{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/dropout/papers/227","list_of":"/method/dropout","method":"Dropout","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":227,"pages_in_order":275,"rows_per_page":100,"rows":[22601,22700],"of":27472,"counts":{"archive_papers_tagged":27472,"with_a_code_link":12129,"where_syntology_ran_a_sample":3620,"not_listed_spam_title":0,"listed":27472,"listed_where_code_ran":3620,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":3044,"every_run_a_failure_of_syntologys_instrument":576,"listed_with_a_run_with_no_instrument_failure":3044,"listed_every_run_a_failure_of_syntologys_instrument":576,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/dropout","prev":"/method/dropout/papers/226","next":"/method/dropout/papers/228","papers":[{"paper":"/paper/vulnerability-analysis-of-face-morphing","slug":"vulnerability-analysis-of-face-morphing","title":"Vulnerability Analysis of Face Morphing Attacks from Landmarks and Generative Adversarial Networks","date":"2020-12-09","arxiv_id":"2012.05344","n_code_links":1,"syntology":null},{"paper":null,"slug":"discourse-parsing-of-contentious-non","title":"Discourse Parsing of Contentious, Non-Convergent Online Discussions","date":"2020-12-08","arxiv_id":"2012.04585","n_code_links":0,"syntology":null},{"paper":null,"slug":"do-mass-media-shape-public-opinion-toward","title":"Large-scale Quantitative Evidence of Media Impact on Public Opinion toward China","date":"2020-12-08","arxiv_id":"2012.07575","n_code_links":0,"syntology":null},{"paper":null,"slug":"efficient-estimation-of-influence-of-a","title":"Efficient Estimation of Influence of a Training Instance","date":"2020-12-08","arxiv_id":"2012.04207","n_code_links":0,"syntology":null},{"paper":"/paper/extractive-opinion-summarization-in-quantized","slug":"extractive-opinion-summarization-in-quantized","title":"Extractive Opinion Summarization in Quantized Transformer Spaces","date":"2020-12-08","arxiv_id":"2012.04443","n_code_links":2,"syntology":null},{"paper":"/paper/from-bag-of-sentences-to-document-distantly","slug":"from-bag-of-sentences-to-document-distantly","title":"From Bag of Sentences to Document: Distantly Supervised Relation Extraction via Machine Reading Comprehension","date":"2020-12-08","arxiv_id":"2012.04334","n_code_links":1,"syntology":null},{"paper":null,"slug":"parameter-efficient-multimodal-transformers-1","title":"Parameter Efficient Multimodal Transformers for Video Representation Learning","date":"2020-12-08","arxiv_id":"2012.04124","n_code_links":0,"syntology":null},{"paper":"/paper/tado-time-varying-attention-with-dual","slug":"tado-time-varying-attention-with-dual","title":"TADO: Time-varying Attention with Dual-Optimizer Model","date":"2020-12-08","arxiv_id":"2012.04558","n_code_links":1,"syntology":null},{"paper":null,"slug":"texture-transform-attention-for-realistic","title":"Texture Transform Attention for Realistic Image Inpainting","date":"2020-12-08","arxiv_id":"2012.04242","n_code_links":0,"syntology":null},{"paper":null,"slug":"an-empirical-survey-of-unsupervised-text","title":"An Empirical Survey of Unsupervised Text Representation Methods on Twitter Data","date":"2020-12-07","arxiv_id":"2012.03468","n_code_links":0,"syntology":null},{"paper":"/paper/cx-db8-a-queryable-extractive-summarizer-and","slug":"cx-db8-a-queryable-extractive-summarizer-and","title":"CX DB8: A queryable extractive summarizer and semantic search engine","date":"2020-12-07","arxiv_id":"2012.03942","n_code_links":2,"syntology":null},{"paper":null,"slug":"dartmouth-cs-at-wnut-2020-task-2-informative","title":"Dartmouth CS at WNUT-2020 Task 2: Informative COVID-19 Tweet Classification Using BERT","date":"2020-12-07","arxiv_id":"2012.04539","n_code_links":0,"syntology":null},{"paper":null,"slug":"deep-policy-networks-for-npc-behaviors-that","title":"Deep Policy Networks for NPC Behaviors that Adapt to Changing Design Parameters in Roguelike Games","date":"2020-12-07","arxiv_id":"2012.03532","n_code_links":0,"syntology":null},{"paper":"/paper/detecting-insincere-questions-from-text-a","slug":"detecting-insincere-questions-from-text-a","title":"Detecting Insincere Questions from Text: A Transfer Learning Approach","date":"2020-12-07","arxiv_id":"2012.07587","n_code_links":1,"syntology":null},{"paper":"/paper/diffprune-neural-network-pruning-with","slug":"diffprune-neural-network-pruning-with","title":"DiffPrune: Neural Network Pruning with Deterministic Approximate Binary Gates and $L_0$ Regularization","date":"2020-12-07","arxiv_id":"2012.03653","n_code_links":1,"syntology":null},{"paper":null,"slug":"document-graph-for-neural-machine-translation","title":"Document Graph for Neural Machine Translation","date":"2020-12-07","arxiv_id":"2012.03477","n_code_links":0,"syntology":null},{"paper":null,"slug":"kgplm-knowledge-guided-language-model-pre","title":"KgPLM: Knowledge-guided Language Model Pre-training via Generative and Discriminative Learning","date":"2020-12-07","arxiv_id":"2012.03551","n_code_links":0,"syntology":null},{"paper":"/paper/ubar-towards-fully-end-to-end-task-oriented","slug":"ubar-towards-fully-end-to-end-task-oriented","title":"UBAR: Towards Fully End-to-End Task-Oriented Dialog Systems with GPT-2","date":"2020-12-07","arxiv_id":"2012.03539","n_code_links":1,"syntology":null},{"paper":null,"slug":"using-previous-acoustic-context-to-improve","title":"Using previous acoustic context to improve Text-to-Speech synthesis","date":"2020-12-07","arxiv_id":"2012.03763","n_code_links":0,"syntology":null},{"paper":"/paper/re-satellite-image-time-series-classification","slug":"re-satellite-image-time-series-classification","title":"[Re] Satellite Image Time Series Classification with Pixel-Set Encoders and Temporal Self-Attention","date":"2020-12-06","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"data-efficient-methods-for-dialogue-systems","title":"Data-Efficient Methods for Dialogue Systems","date":"2020-12-05","arxiv_id":"2012.02929","n_code_links":0,"syntology":null},{"paper":null,"slug":"enhanced-offensive-language-detection-through","title":"Enhanced Offensive Language Detection Through Data Augmentation","date":"2020-12-05","arxiv_id":"2012.02954","n_code_links":0,"syntology":null},{"paper":"/paper/parallel-blockwise-knowledge-distillation-for","slug":"parallel-blockwise-knowledge-distillation-for","title":"Parallel Blockwise Knowledge Distillation for Deep Neural Network Compression","date":"2020-12-05","arxiv_id":"2012.03096","n_code_links":1,"syntology":null},{"paper":null,"slug":"paranet-deep-regular-representation-for-3d","title":"ParaNet: Deep Regular Representation for 3D Point Clouds","date":"2020-12-05","arxiv_id":"2012.03028","n_code_links":0,"syntology":null},{"paper":"/paper/pre-training-protein-language-models-with","slug":"pre-training-protein-language-models-with","title":"Pre-training Protein Language Models with Label-Agnostic Binding Pairs Enhances Performance in Downstream Tasks","date":"2020-12-05","arxiv_id":"2012.03084","n_code_links":1,"syntology":null},{"paper":null,"slug":"automated-detection-of-cyberbullying-against","title":"Automated Detection of Cyberbullying Against Women and Immigrants and Cross-domain Adaptability","date":"2020-12-04","arxiv_id":"2012.02565","n_code_links":0,"syntology":null},{"paper":null,"slug":"bayesian-active-learning-for-wearable-stress","title":"Bayesian Active Learning for Wearable Stress and Affect Detection","date":"2020-12-04","arxiv_id":"2012.02702","n_code_links":0,"syntology":null},{"paper":null,"slug":"cued-speech-at-trec-2020-podcast","title":"CUED_speech at TREC 2020 Podcast Summarisation Track","date":"2020-12-04","arxiv_id":"2012.02535","n_code_links":0,"syntology":null},{"paper":"/paper/echobert-a-transformer-based-approach-for","slug":"echobert-a-transformer-based-approach-for","title":"EchoBERT: A Transformer-Based Approach for Behavior Detection in Echograms","date":"2020-12-04","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"fine-tuning-bert-for-low-resource-natural","title":"Fine-tuning BERT for Low-Resource Natural Language Understanding via Active Learning","date":"2020-12-04","arxiv_id":"2012.02462","n_code_links":0,"syntology":null},{"paper":null,"slug":"modelling-general-properties-of-nouns-by","title":"Modelling General Properties of Nouns by Selectively Averaging Contextualised Embeddings","date":"2020-12-04","arxiv_id":"2012.07580","n_code_links":0,"syntology":null},{"paper":null,"slug":"playing-text-based-games-with-common-sense","title":"Playing Text-Based Games with Common Sense","date":"2020-12-04","arxiv_id":"2012.02757","n_code_links":0,"syntology":null},{"paper":null,"slug":"pre-trained-language-models-as-knowledge","title":"Pre-trained language models as knowledge bases for Automotive Complaint Analysis","date":"2020-12-04","arxiv_id":"2012.02558","n_code_links":0,"syntology":null},{"paper":null,"slug":"relational-pretrained-transformers-towards","title":"RPT: Relational Pre-trained Transformer Is Almost All You Need towards Democratizing Data Preparation","date":"2020-12-04","arxiv_id":"2012.02469","n_code_links":0,"syntology":null},{"paper":null,"slug":"spread-mechanism-and-influence-measurement-of","title":"Spread Mechanism and Influence Measurement of Online Rumors in China During the COVID-19 Pandemic","date":"2020-12-04","arxiv_id":"2012.02446","n_code_links":0,"syntology":null},{"paper":"/paper/bert-hlstms-bert-and-hierarchical-lstms-for","slug":"bert-hlstms-bert-and-hierarchical-lstms-for","title":"BERT-hLSTMs: BERT and Hierarchical LSTMs for Visual Storytelling","date":"2020-12-03","arxiv_id":"2012.02128","n_code_links":0,"syntology":null},{"paper":null,"slug":"circles-are-like-ellipses-or-ellipses-are","title":"Circles are like Ellipses, or Ellipses are like Circles? Measuring the Degree of Asymmetry of Static and Contextual Embeddings and the Implications to Representation Learning","date":"2020-12-03","arxiv_id":"2012.01631","n_code_links":0,"syntology":null},{"paper":"/paper/dialogbert-discourse-aware-response","slug":"dialogbert-discourse-aware-response","title":"DialogBERT: Discourse-Aware Response Generation via Learning to Recover and Rank Utterances","date":"2020-12-03","arxiv_id":"2012.01775","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":0,"n_instrument":2,"unverified":0,"pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":null}},{"paper":null,"slug":"evolving-character-level-densenet","title":"Evolving Character-Level DenseNet Architectures using Genetic Programming","date":"2020-12-03","arxiv_id":"2012.02327","n_code_links":0,"syntology":null},{"paper":null,"slug":"federated-learning-with-diversified","title":"Federated Learning for Personalized Humor Recognition","date":"2020-12-03","arxiv_id":"2012.01675","n_code_links":0,"syntology":null},{"paper":null,"slug":"gottbert-a-pure-german-language-model","title":"GottBERT: a pure German Language Model","date":"2020-12-03","arxiv_id":"2012.02110","n_code_links":0,"syntology":null},{"paper":null,"slug":"multiple-networks-are-more-efficient-than-one","title":"Wisdom of Committees: An Overlooked Approach To Faster and More Accurate Models","date":"2020-12-03","arxiv_id":"2012.01988","n_code_links":0,"syntology":null},{"paper":"/paper/sentiment-analysis-in-bengali-via-transfer","slug":"sentiment-analysis-in-bengali-via-transfer","title":"Sentiment analysis in Bengali via transfer learning using multi-lingual BERT","date":"2020-12-03","arxiv_id":"2012.07538","n_code_links":1,"syntology":null},{"paper":null,"slug":"trace-early-detection-of-chronic-kidney","title":"TRACE: Early Detection of Chronic Kidney Disease Onset with Transformer-Enhanced Feature Embedding","date":"2020-12-03","arxiv_id":"2012.03729","n_code_links":0,"syntology":null},{"paper":"/paper/triplet-entropy-loss-improving-the","slug":"triplet-entropy-loss-improving-the","title":"Triplet Entropy Loss: Improving The Generalisation of Short Speech Language Identification Systems","date":"2020-12-03","arxiv_id":"2012.03775","n_code_links":1,"syntology":null},{"paper":null,"slug":"what-makes-a-star-teacher-a-hierarchical-bert","title":"What Makes a Star Teacher? A Hierarchical BERT Model for Evaluating Teacher's Performance in Online Education","date":"2020-12-03","arxiv_id":"2012.01633","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-framework-and-dataset-for-abstract-art","title":"A Framework and Dataset for Abstract Art Generation via CalligraphyGAN","date":"2020-12-02","arxiv_id":"2012.00744","n_code_links":0,"syntology":null},{"paper":"/paper/contour-transformer-network-for-one-shot","slug":"contour-transformer-network-for-one-shot","title":"Contour Transformer Network for One-shot Segmentation of Anatomical Structures","date":"2020-12-02","arxiv_id":"2012.01480","n_code_links":1,"syntology":null},{"paper":null,"slug":"data-driven-analysis-of-turbulent-flame","title":"Data-driven Analysis of Turbulent Flame Images","date":"2020-12-02","arxiv_id":"2012.01485","n_code_links":0,"syntology":null},{"paper":"/paper/exploiting-bert-to-improve-aspect-based","slug":"exploiting-bert-to-improve-aspect-based","title":"Exploiting BERT to improve aspect-based sentiment analysis performance on Persian language","date":"2020-12-02","arxiv_id":"2012.07510","n_code_links":1,"syntology":null},{"paper":"/paper/how-can-we-know-when-language-models-know","slug":"how-can-we-know-when-language-models-know","title":"How Can We Know When Language Models Know? On the Calibration of Language Models for Question Answering","date":"2020-12-02","arxiv_id":"2012.00955","n_code_links":1,"syntology":null},{"paper":"/paper/two-stage-single-image-reflection-removal","slug":"two-stage-single-image-reflection-removal","title":"Two-Stage Single Image Reflection Removal with Reflection-Aware Guidance","date":"2020-12-02","arxiv_id":"2012.00945","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-bayesian-nonparametrics-view-into-deep","title":"A Bayesian Nonparametrics View into Deep Representations","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"a-deep-generative-approach-to-native-language","title":"A Deep Generative Approach to Native Language Identification","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"a-large-scale-corpus-of-e-mail-conversations","title":"A Large-Scale Corpus of E-mail Conversations with Standard and Two-Level Dialogue Act Annotations","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"a-neural-local-coherence-analysis-model-for","title":"A Neural Local Coherence Analysis Model for Clarity Text Scoring","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/adversarial-sparse-transformer-for-time","slug":"adversarial-sparse-transformer-for-time","title":"Adversarial Sparse Transformer for Time Series Forecasting","date":"2020-12-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/affective-and-contextual-embedding-for","slug":"affective-and-contextual-embedding-for","title":"Affective and Contextual Embedding for Sarcasm Detection","date":"2020-12-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"alexu-aux-bert-at-semeval-2020-task-3","title":"AlexU-AUX-BERT at SemEval-2020 Task 3: Improving BERT Contextual Similarity Using Multiple Auxiliary Contexts","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"alexu-backtranslation-tl-at-semeval-2020-task","title":"AlexU-BackTranslation-TL at SemEval-2020 Task 12: Improving Offensive Language Detection Using Data Augmentation and Transfer Learning","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"alt-at-semeval-2020-task-12-arabic-and","title":"ALT at SemEval-2020 Task 12: Arabic and English Offensive Language Identification in Social Media","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/analogy-models-for-neural-word-inflection","slug":"analogy-models-for-neural-word-inflection","title":"Analogy Models for Neural Word Inflection","date":"2020-12-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"arabizi-language-models-for-sentiment","title":"Arabizi Language Models for Sentiment Analysis","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"assessing-polyseme-sense-similarity-through","title":"Assessing Polyseme Sense Similarity through Co-predication Acceptability and Contextualised Embedding Distance","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"asymptotic-convergence-rate-of-dropout-on","title":"Asymptotic convergence rate of Dropout on shallow linear neural networks","date":"2020-12-01","arxiv_id":"2012.01978","n_code_links":0,"syntology":null},{"paper":"/paper/attentively-embracing-noise-for-robust-latent","slug":"attentively-embracing-noise-for-robust-latent","title":"Attentively Embracing Noise for Robust Latent Representation in BERT","date":"2020-12-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/autoencoders-that-don-t-overfit-towards-the","slug":"autoencoders-that-don-t-overfit-towards-the","title":"Autoencoders that don't overfit towards the Identity","date":"2020-12-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"bert-at-semeval-2020-task-8-using-bert-to","title":"BERT at SemEval-2020 Task 8: Using BERT to Analyse Meme Emotions","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/bert-based-cohesion-analysis-of-japanese","slug":"bert-based-cohesion-analysis-of-japanese","title":"BERT-based Cohesion Analysis of Japanese Texts","date":"2020-12-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"bert-based-neural-collaborative-filtering-and","title":"BERT-Based Neural Collaborative Filtering and Fixed-Length Contiguous Tokens Explanation","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"bertatde-at-semeval-2020-task-6-extracting","title":"BERTatDE at SemEval-2020 Task 6: Extracting Term-definition Pairs in Free Text Using Pre-trained Model","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"bilingual-subword-segmentation-for-neural","title":"Bilingual Subword Segmentation for Neural Machine Translation","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"blcu-nlp-at-semeval-2020-task-5-data","title":"BLCU-NLP at SemEval-2020 Task 5: Data Augmentation for Efficient Counterfactual Detecting","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"byteam-at-semeval-2020-task-5-detecting","title":"BYteam at SemEval-2020 Task 5: Detecting Counterfactual Statements with BERT and Ensembles","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"cardiff-university-at-semeval-2020-task-6","title":"Cardiff University at SemEval-2020 Task 6: Fine-tuning BERT for Domain-Specific Definition Classification","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"citiusnlp-at-semeval-2020-task-3-comparing","title":"CitiusNLP at SemEval-2020 Task 3: Comparing Two Approaches for Word Vector Contextualization","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/classifier-probes-may-just-learn-from-linear","slug":"classifier-probes-may-just-learn-from-linear","title":"Classifier Probes May Just Learn from Linear Context Features","date":"2020-12-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"climatext-a-dataset-for-climate-change-topic","title":"ClimaText: A Dataset for Climate Change Topic Detection","date":"2020-12-01","arxiv_id":"2012.00483","n_code_links":0,"syntology":null},{"paper":null,"slug":"cn-hit-mi-t-at-semeval-2020-task-8-memotion","title":"CN-HIT-MI.T at SemEval-2020 Task 8: Memotion Analysis Based on BERT","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/cogltx-applying-bert-to-long-texts","slug":"cogltx-applying-bert-to-long-texts","title":"CogLTX: Applying BERT to Long Texts","date":"2020-12-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"comparing-probabilistic-distributional-and","title":"Comparing Probabilistic, Distributional and Transformer-Based Models on Logical Metonymy Interpretation","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/constituency-lattice-encoding-for-aspect-term","slug":"constituency-lattice-encoding-for-aspect-term","title":"Constituency Lattice Encoding for Aspect Term Extraction","date":"2020-12-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"contextualized-embeddings-for-enriching","title":"Contextualized Embeddings for Enriching Linguistic Analyses on Politeness","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/cpm-a-large-scale-generative-chinese-pre","slug":"cpm-a-large-scale-generative-chinese-pre","title":"CPM: A Large-scale Generative Chinese Pre-trained Language Model","date":"2020-12-01","arxiv_id":"2012.00413","n_code_links":10,"syntology":null},{"paper":null,"slug":"data-selection-for-bilingual-lexicon","title":"Data Selection for Bilingual Lexicon Induction from Specialized Comparable Corpora","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"deep-learning-brasil-nlp-at-semeval-2020-task-1","title":"Deep Learning Brasil - NLP at SemEval-2020 Task 9: Sentiment Analysis of Code-Mixed Tweets Using Ensemble of Language Models","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"deepyang-at-semeval-2020-task-4-using-the","title":"DEEPYANG at SemEval-2020 Task 4: Using the Hidden Layer State of BERT Model for Differentiating Common Sense","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"denoising-pre-training-and-data-augmentation","title":"Denoising Pre-Training and Data Augmentation Strategies for Enhanced RDF Verbalization with Transformers","date":"2020-12-01","arxiv_id":"2012.00571","n_code_links":0,"syntology":null},{"paper":"/paper/detecting-urgency-status-of-crisis-tweets-a","slug":"detecting-urgency-status-of-crisis-tweets-a","title":"Detecting Urgency Status of Crisis Tweets: A Transfer Learning Approach for Low Resource Languages","date":"2020-12-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"discovery-team-at-semeval-2020-task-1-context","title":"Discovery Team at SemEval-2020 Task 1: Context-sensitive Embeddings Not Always Better than Static for Semantic Change Detection","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"diverse-dialogue-generation-with-context","title":"Diverse dialogue generation with context dependent dynamic loss function","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"document-level-neural-machine-translation-3","title":"Document-Level Neural Machine Translation Using BERT as Context Encoder","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"domain-transfer-based-data-augmentation-for","title":"Domain Transfer based Data Augmentation for Neural Query Translation","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"don-t-invite-bert-to-drink-a-bottle-modeling","title":"Don't Invite BERT to Drink a Bottle: Modeling the Interpretation of Metonymies Using BERT and Distributional Representations","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"emotion-classification-by-jointly-learning-to","title":"Emotion Classification by Jointly Learning to Lexiconize and Classify","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"evaluating-pretrained-transformer-based","title":"Evaluating Pretrained Transformer-based Models on the Task of Fine-Grained Named Entity Recognition","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"evaluating-unsupervised-representation","title":"Evaluating Unsupervised Representation Learning for Detecting Stances of Fake News","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"exploring-statistical-and-neural-models-for","title":"Exploring Statistical and Neural Models for Noun Ellipsis Detection and Resolution in English","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"fastmatch-accelerating-the-inference-of-bert","title":"FASTMATCH: Accelerating the Inference of BERT-based Text Matching","date":"2020-12-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/fbk-dh-at-semeval-2020-task-12-using-multi","slug":"fbk-dh-at-semeval-2020-task-12-using-multi","title":"FBK-DH at SemEval-2020 Task 12: Using Multi-channel BERT for Multilingual Offensive Language Detection","date":"2020-12-01","arxiv_id":null,"n_code_links":1,"syntology":null}],"record_sha256":"4508e4d179970995f5204936df6fbdd49ea5d48d8052a3db595352a7d4499c68","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}