{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/linear-warmup-with-linear-decay/papers/57","list_of":"/method/linear-warmup-with-linear-decay","method":"Linear Warmup With Linear Decay","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":57,"pages_in_order":71,"rows_per_page":100,"rows":[5601,5700],"of":7076,"counts":{"archive_papers_tagged":7076,"with_a_code_link":2913,"where_syntology_ran_a_sample":650,"not_listed_spam_title":0,"listed":7076,"listed_where_code_ran":650,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":531,"every_run_a_failure_of_syntologys_instrument":119,"listed_with_a_run_with_no_instrument_failure":531,"listed_every_run_a_failure_of_syntologys_instrument":119,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/linear-warmup-with-linear-decay","prev":"/method/linear-warmup-with-linear-decay/papers/56","next":"/method/linear-warmup-with-linear-decay/papers/58","papers":[{"paper":"/paper/accelerating-training-of-transformer-based","slug":"accelerating-training-of-transformer-based","title":"Accelerating Training of Transformer-Based Language Models with Progressive Layer Dropping","date":"2020-10-26","arxiv_id":"2010.13369","n_code_links":1,"syntology":null},{"paper":"/paper/fine-grained-information-status-1","slug":"fine-grained-information-status-1","title":"Fine-grained Information Status Classification Using Discourse Context-Aware BERT","date":"2020-10-26","arxiv_id":"2010.14759","n_code_links":1,"syntology":null},{"paper":"/paper/semi-supervised-spoken-language-understanding","slug":"semi-supervised-spoken-language-understanding","title":"Semi-Supervised Spoken Language Understanding via Self-Supervised Speech and Language Model Pretraining","date":"2020-10-26","arxiv_id":"2010.13826","n_code_links":1,"syntology":null},{"paper":null,"slug":"upb-at-semeval-2020-task-12-multilingual","title":"UPB at SemEval-2020 Task 12: Multilingual Offensive Language Detection on Social Media by Fine-tuning a Variety of BERT-based Models","date":"2020-10-26","arxiv_id":"2010.13609","n_code_links":0,"syntology":null},{"paper":null,"slug":"commonsense-knowledge-adversarial-dataset","title":"Commonsense knowledge adversarial dataset that challenges ELECTRA","date":"2020-10-25","arxiv_id":"2010.13049","n_code_links":0,"syntology":null},{"paper":null,"slug":"contextualized-word-embeddings-encode-aspects","title":"Contextualized Word Embeddings Encode Aspects of Human-Like Word Sense Knowledge","date":"2020-10-25","arxiv_id":"2010.13057","n_code_links":0,"syntology":null},{"paper":null,"slug":"crab-class-representation-attentive-bert-for","title":"CRAB: Class Representation Attentive BERT for Hate Speech Identification in Social Media","date":"2020-10-25","arxiv_id":"2010.13028","n_code_links":0,"syntology":null},{"paper":"/paper/two-stage-textual-knowledge-distillation-to","slug":"two-stage-textual-knowledge-distillation-to","title":"Two-stage Textual Knowledge Distillation for End-to-End Spoken Language Understanding","date":"2020-10-25","arxiv_id":"2010.13105","n_code_links":1,"syntology":null},{"paper":null,"slug":"char2subword-extending-the-subword-embedding","title":"Char2Subword: Extending the Subword Embedding Space Using Robust Character Compositionality","date":"2020-10-24","arxiv_id":"2010.12730","n_code_links":0,"syntology":null},{"paper":"/paper/cough-a-challenge-dataset-and-models-for","slug":"cough-a-challenge-dataset-and-models-for","title":"COUGH: A Challenge Dataset and Models for COVID-19 FAQ Retrieval","date":"2020-10-24","arxiv_id":"2010.12800","n_code_links":1,"syntology":null},{"paper":"/paper/multi-domain-dialogue-state-tracking-a-purely","slug":"multi-domain-dialogue-state-tracking-a-purely","title":"Jointly Optimizing State Operation Prediction and Value Generation for Dialogue State Tracking","date":"2020-10-24","arxiv_id":"2010.14061","n_code_links":2,"syntology":null},{"paper":"/paper/pre-trained-summarization-distillation","slug":"pre-trained-summarization-distillation","title":"Pre-trained Summarization Distillation","date":"2020-10-24","arxiv_id":"2010.13002","n_code_links":1,"syntology":null},{"paper":"/paper/barthez-a-skilled-pretrained-french-sequence","slug":"barthez-a-skilled-pretrained-french-sequence","title":"BARThez: a Skilled Pretrained French Sequence-to-Sequence Model","date":"2020-10-23","arxiv_id":"2010.12321","n_code_links":5,"syntology":null},{"paper":"/paper/did-you-ask-a-good-question-a-cross-domain","slug":"did-you-ask-a-good-question-a-cross-domain","title":"Did You Ask a Good Question? A Cross-Domain Question Intention Classification Benchmark for Text-to-SQL","date":"2020-10-23","arxiv_id":"2010.12634","n_code_links":1,"syntology":null},{"paper":"/paper/ernie-gram-pre-training-with-explicitly-n","slug":"ernie-gram-pre-training-with-explicitly-n","title":"ERNIE-Gram: Pre-Training with Explicitly N-Gram Masked Language Modeling for Natural Language Understanding","date":"2020-10-23","arxiv_id":"2010.12148","n_code_links":2,"syntology":null},{"paper":null,"slug":"gibert-introducing-linguistic-knowledge-into","title":"GiBERT: Introducing Linguistic Knowledge into BERT through a Lightweight Gated Injection Method","date":"2020-10-23","arxiv_id":"2010.12532","n_code_links":0,"syntology":null},{"paper":"/paper/hatebert-retraining-bert-for-abusive-language","slug":"hatebert-retraining-bert-for-abusive-language","title":"HateBERT: Retraining BERT for Abusive Language Detection in English","date":"2020-10-23","arxiv_id":"2010.12472","n_code_links":1,"syntology":null},{"paper":"/paper/lightseq-a-high-performance-inference-library","slug":"lightseq-a-high-performance-inference-library","title":"LightSeq: A High Performance Inference Library for Transformers","date":"2020-10-23","arxiv_id":"2010.13887","n_code_links":1,"syntology":null},{"paper":"/paper/long-document-ranking-with-query-directed","slug":"long-document-ranking-with-query-directed","title":"Long Document Ranking with Query-Directed Sparse Transformer","date":"2020-10-23","arxiv_id":"2010.12683","n_code_links":1,"syntology":null},{"paper":null,"slug":"on-the-transformer-growth-for-progressive","title":"On the Transformer Growth for Progressive BERT Training","date":"2020-10-23","arxiv_id":"2010.12562","n_code_links":0,"syntology":null},{"paper":"/paper/posterior-differential-regularization-with-f","slug":"posterior-differential-regularization-with-f","title":"Posterior Differential Regularization with f-divergence for Improving Model Robustness","date":"2020-10-23","arxiv_id":"2010.12638","n_code_links":2,"syntology":null},{"paper":null,"slug":"pre-trained-model-for-chinese-word","title":"Pre-training with Meta Learning for Chinese Word Segmentation","date":"2020-10-23","arxiv_id":"2010.12272","n_code_links":0,"syntology":null},{"paper":null,"slug":"st-bert-cross-modal-language-model-pre","title":"ST-BERT: Cross-modal Language Model Pre-training For End-to-end Spoken Language Understanding","date":"2020-10-23","arxiv_id":"2010.12283","n_code_links":0,"syntology":null},{"paper":null,"slug":"topic-modeling-with-contextualized-word","title":"Topic Modeling with Contextualized Word Representation Clusters","date":"2020-10-23","arxiv_id":"2010.12626","n_code_links":0,"syntology":null},{"paper":"/paper/distilling-dense-representations-for-ranking","slug":"distilling-dense-representations-for-ranking","title":"Distilling Dense Representations for Ranking using Tightly-Coupled Teachers","date":"2020-10-22","arxiv_id":"2010.11386","n_code_links":2,"syntology":null},{"paper":"/paper/improving-bert-performance-for-aspect-based","slug":"improving-bert-performance-for-aspect-based","title":"Improving BERT Performance for Aspect-Based Sentiment Analysis","date":"2020-10-22","arxiv_id":"2010.11731","n_code_links":2,"syntology":{"ran":7,"of":11,"n_ran_checked":4,"n_instrument":3,"unverified":4,"pointer_only":2,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 2 honoured, 0 violated, 2 with no contract checked; 3 where Syntology's instrument failed) · 4 unverified","official":{"repos":["IMPLabUniPr/BERT-for-ABSA"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":"/paper/investigating-the-true-performance-of","slug":"investigating-the-true-performance-of","title":"Exploiting News Article Structure for Automatic Corpus Generation of Entailment Datasets","date":"2020-10-22","arxiv_id":"2010.11574","n_code_links":1,"syntology":null},{"paper":"/paper/knowledge-distillation-for-bert-unsupervised","slug":"knowledge-distillation-for-bert-unsupervised","title":"Knowledge Distillation for BERT Unsupervised Domain Adaptation","date":"2020-10-22","arxiv_id":"2010.11478","n_code_links":1,"syntology":null},{"paper":"/paper/language-models-are-open-knowledge-graphs-1","slug":"language-models-are-open-knowledge-graphs-1","title":"Language Models are Open Knowledge Graphs","date":"2020-10-22","arxiv_id":"2010.11967","n_code_links":2,"syntology":{"ran":7,"of":8,"n_ran_checked":6,"n_instrument":1,"unverified":1,"pointer_only":1,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 1 violated, 5 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":null}},{"paper":"/paper/self-alignment-pre-training-for-biomedical","slug":"self-alignment-pre-training-for-biomedical","title":"Self-Alignment Pretraining for Biomedical Entity Representations","date":"2020-10-22","arxiv_id":"2010.11784","n_code_links":1,"syntology":null},{"paper":null,"slug":"towards-fully-bilingual-deep-language","title":"Towards Fully Bilingual Deep Language Modeling","date":"2020-10-22","arxiv_id":"2010.11639","n_code_links":0,"syntology":null},{"paper":null,"slug":"unicase-rethinking-casing-in-language-models","title":"UniCase -- Rethinking Casing in Language Models","date":"2020-10-22","arxiv_id":"2010.11936","n_code_links":0,"syntology":null},{"paper":null,"slug":"detection-of-covid-19-informative-tweets","title":"Detection of COVID-19 informative tweets using RoBERTa","date":"2020-10-21","arxiv_id":"2010.11238","n_code_links":0,"syntology":null},{"paper":"/paper/generalized-conditioned-dialogue-generation","slug":"generalized-conditioned-dialogue-generation","title":"A Simple and Efficient Multi-Task Learning Approach for Conditioned Dialogue Generation","date":"2020-10-21","arxiv_id":"2010.11140","n_code_links":1,"syntology":null},{"paper":"/paper/german-s-next-language-model","slug":"german-s-next-language-model","title":"German's Next Language Model","date":"2020-10-21","arxiv_id":"2010.10906","n_code_links":1,"syntology":null},{"paper":null,"slug":"latte-mix-measuring-sentence-semantic","title":"Latte-Mix: Measuring Sentence Semantic Similarity with Latent Categorical Mixtures","date":"2020-10-21","arxiv_id":"2010.11351","n_code_links":0,"syntology":null},{"paper":"/paper/automets-the-autocomplete-for-medical-text","slug":"automets-the-autocomplete-for-medical-text","title":"AutoMeTS: The Autocomplete for Medical Text Simplification","date":"2020-10-20","arxiv_id":"2010.10573","n_code_links":1,"syntology":null},{"paper":null,"slug":"bert2dnn-bert-distillation-with-massive","title":"BERT2DNN: BERT Distillation with Massive Unlabeled Data for Online E-Commerce Search","date":"2020-10-20","arxiv_id":"2010.10442","n_code_links":0,"syntology":null},{"paper":"/paper/characterbert-reconciling-elmo-and-bert-for","slug":"characterbert-reconciling-elmo-and-bert-for","title":"CharacterBERT: Reconciling ELMo and BERT for Word-Level Open-Vocabulary Representations From Characters","date":"2020-10-20","arxiv_id":"2010.10392","n_code_links":2,"syntology":{"ran":6,"of":8,"n_ran_checked":4,"n_instrument":2,"unverified":2,"pointer_only":4,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 2 honoured, 0 violated, 2 with no contract checked; 2 where Syntology's instrument failed) · 2 unverified","official":{"repos":["helboukkouri/character-bert"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/conjnli-natural-language-inference-over","slug":"conjnli-natural-language-inference-over","title":"ConjNLI: Natural Language Inference Over Conjunctive Sentences","date":"2020-10-20","arxiv_id":"2010.10418","n_code_links":1,"syntology":{"ran":5,"of":7,"n_ran_checked":4,"n_instrument":1,"unverified":2,"pointer_only":1,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","official":{"repos":["swarnaHub/ConjNLI"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/cort-complementary-rankings-from-transformers","slug":"cort-complementary-rankings-from-transformers","title":"CoRT: Complementary Rankings from Transformers","date":"2020-10-20","arxiv_id":"2010.10252","n_code_links":1,"syntology":null},{"paper":"/paper/language-representation-in-multilingual","slug":"language-representation-in-multilingual","title":"Looking for Clues of Language in Multilingual BERT to Improve Cross-lingual Generalization","date":"2020-10-20","arxiv_id":"2010.10041","n_code_links":1,"syntology":null},{"paper":"/paper/optimal-subarchitecture-extraction-for-bert","slug":"optimal-subarchitecture-extraction-for-bert","title":"Optimal Subarchitecture Extraction For BERT","date":"2020-10-20","arxiv_id":"2010.10499","n_code_links":3,"syntology":null},{"paper":null,"slug":"performance-of-transfer-learning-model-vs","title":"Performance of Transfer Learning Model vs. Traditional Neural Network in Low System Resource Environment","date":"2020-10-20","arxiv_id":"2011.07962","n_code_links":0,"syntology":null},{"paper":"/paper/prop-pre-training-with-representative-words","slug":"prop-pre-training-with-representative-words","title":"PROP: Pre-training with Representative Words Prediction for Ad-hoc Retrieval","date":"2020-10-20","arxiv_id":"2010.10137","n_code_links":1,"syntology":null},{"paper":null,"slug":"text-classification-of-covid-19-press","title":"Text Classification of Manifestos and COVID-19 Press Briefings using BERT and Convolutional Neural Networks","date":"2020-10-20","arxiv_id":"2010.10267","n_code_links":0,"syntology":null},{"paper":null,"slug":"what-makes-multilingual-bert-multilingual","title":"What makes multilingual BERT multilingual?","date":"2020-10-20","arxiv_id":"2010.10938","n_code_links":0,"syntology":null},{"paper":"/paper/bertnesia-investigating-the-capture-and","slug":"bertnesia-investigating-the-capture-and","title":"BERTnesia: Investigating the capture and forgetting of knowledge in BERT","date":"2020-10-19","arxiv_id":"2010.09313","n_code_links":1,"syntology":null},{"paper":null,"slug":"better-distractions-transformer-based","title":"Better Distractions: Transformer-based Distractor Generation and Multiple Choice Question Filtering","date":"2020-10-19","arxiv_id":"2010.09598","n_code_links":0,"syntology":null},{"paper":"/paper/cold-start-active-learning-through-self","slug":"cold-start-active-learning-through-self","title":"Cold-start Active Learning through Self-supervised Language Modeling","date":"2020-10-19","arxiv_id":"2010.09535","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["forest-snow/alps"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/colloql-robust-cross-domain-text-to-sql-over","slug":"colloql-robust-cross-domain-text-to-sql-over","title":"ColloQL: Robust Cross-Domain Text-to-SQL Over Search Queries","date":"2020-10-19","arxiv_id":"2010.09927","n_code_links":1,"syntology":null},{"paper":"/paper/cross-lingual-transfer-in-zero-shot-cross","slug":"cross-lingual-transfer-in-zero-shot-cross","title":"Cross-Lingual Transfer in Zero-Shot Cross-Language Entity Linking","date":"2020-10-19","arxiv_id":"2010.09828","n_code_links":1,"syntology":null},{"paper":"/paper/drug-repurposing-for-covid-19-via-knowledge","slug":"drug-repurposing-for-covid-19-via-knowledge","title":"Drug Repurposing for COVID-19 via Knowledge Graph Completion","date":"2020-10-19","arxiv_id":"2010.09600","n_code_links":1,"syntology":null},{"paper":"/paper/the-relx-dataset-and-matching-the","slug":"the-relx-dataset-and-matching-the","title":"The RELX Dataset and Matching the Multilingual Blanks for Cross-Lingual Relation Classification","date":"2020-10-19","arxiv_id":"2010.09381","n_code_links":1,"syntology":{"ran":0,"of":1,"n_ran_checked":0,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"0 ran · 1 unverified","official":{"repos":["boun-tabi/RELX"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":[]}}},{"paper":null,"slug":"explaining-and-improving-model-behavior-with","title":"Explaining and Improving Model Behavior with k Nearest Neighbor Representations","date":"2020-10-18","arxiv_id":"2010.09030","n_code_links":0,"syntology":null},{"paper":"/paper/towards-interpreting-bert-for-reading","slug":"towards-interpreting-bert-for-reading","title":"Towards Interpreting BERT for Reading Comprehension Based QA","date":"2020-10-18","arxiv_id":"2010.08983","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":0,"n_instrument":3,"unverified":0,"pointer_only":2,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 3 where Syntology's instrument failed) · 0 unverified","official":{"repos":["iitmnlp/BERT-Analysis-RCQA"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official","unlocated"]}}},{"paper":null,"slug":"answer-checking-in-context-a-multi-modal","title":"Answer-checking in Context: A Multi-modal FullyAttention Network for Visual Question Answering","date":"2020-10-17","arxiv_id":"2010.08708","n_code_links":0,"syntology":null},{"paper":null,"slug":"habertor-an-efficient-and-effective-deep","title":"HABERTOR: An Efficient and Effective Deep Hatespeech Detector","date":"2020-10-17","arxiv_id":"2010.08865","n_code_links":0,"syntology":null},{"paper":null,"slug":"hierarchical-multitask-learning-approach-for","title":"Hierarchical Multitask Learning Approach for BERT","date":"2020-10-17","arxiv_id":"2011.04451","n_code_links":0,"syntology":null},{"paper":null,"slug":"question-answering-over-knowledge-base-using-1","title":"Question Answering over Knowledge Base using Language Model Embeddings","date":"2020-10-17","arxiv_id":"2010.08883","n_code_links":0,"syntology":null},{"paper":"/paper/tweetbert-a-pretrained-language","slug":"tweetbert-a-pretrained-language","title":"TweetBERT: A Pretrained Language Representation Model for Twitter Text Analysis","date":"2020-10-17","arxiv_id":"2010.11091","n_code_links":1,"syntology":null},{"paper":"/paper/coarse-to-fine-pre-training-for-named-entity","slug":"coarse-to-fine-pre-training-for-named-entity","title":"Coarse-to-Fine Pre-training for Named Entity Recognition","date":"2020-10-16","arxiv_id":"2010.08210","n_code_links":1,"syntology":{"ran":4,"of":13,"n_ran_checked":4,"n_instrument":0,"unverified":9,"pointer_only":13,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 1 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 9 unverified","official":{"repos":["strawberryx/CoFEE"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":9,"ran_from_kinds":["official","unlocated"]}}},{"paper":"/paper/delaying-interaction-layers-in-transformer","slug":"delaying-interaction-layers-in-transformer","title":"Delaying Interaction Layers in Transformer-based Encoders for Efficient Open Domain Question Answering","date":"2020-10-16","arxiv_id":"2010.08422","n_code_links":1,"syntology":null},{"paper":"/paper/it-s-not-greek-to-mbert-inducing-word-level","slug":"it-s-not-greek-to-mbert-inducing-word-level","title":"It's not Greek to mBERT: Inducing Word-Level Translations from Multilingual BERT","date":"2020-10-16","arxiv_id":"2010.08275","n_code_links":1,"syntology":{"ran":5,"of":6,"n_ran_checked":2,"n_instrument":3,"unverified":1,"pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 2 honoured, 0 violated, 0 with no contract checked; 3 where Syntology's instrument failed) · 1 unverified","official":{"repos":["gonenhila/mbert"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"linguistically-informed-transformations-lit-a","title":"Linguistically-Informed Transformations (LIT): A Method for Automatically Generating Contrast Sets","date":"2020-10-16","arxiv_id":"2010.08580","n_code_links":0,"syntology":null},{"paper":"/paper/context-guided-bert-for-targeted-aspect-based","slug":"context-guided-bert-for-targeted-aspect-based","title":"Context-Guided BERT for Targeted Aspect-Based Sentiment Analysis","date":"2020-10-15","arxiv_id":"2010.07523","n_code_links":1,"syntology":null},{"paper":"/paper/does-chinese-bert-encode-word-structure","slug":"does-chinese-bert-encode-word-structure","title":"Does Chinese BERT Encode Word Structure?","date":"2020-10-15","arxiv_id":"2010.07711","n_code_links":1,"syntology":null},{"paper":"/paper/neural-deepfake-detection-with-factual","slug":"neural-deepfake-detection-with-factual","title":"Neural Deepfake Detection with Factual Structure of Text","date":"2020-10-15","arxiv_id":"2010.07475","n_code_links":1,"syntology":null},{"paper":null,"slug":"nuig-shubhanker-dravidian-codemix-fire2020","title":"NUIG-Shubhanker@Dravidian-CodeMix-FIRE2020: Sentiment Analysis of Code-Mixed Dravidian text using XLNet","date":"2020-10-15","arxiv_id":"2010.07773","n_code_links":0,"syntology":null},{"paper":null,"slug":"response-selection-for-multi-party","title":"Response Selection for Multi-Party Conversations withDynamic Topic Tracking","date":"2020-10-15","arxiv_id":"2010.07785","n_code_links":0,"syntology":null},{"paper":null,"slug":"unsupervised-bitext-mining-and-translation","title":"Unsupervised Bitext Mining and Translation via Self-trained Contextual Embeddings","date":"2020-10-15","arxiv_id":"2010.07761","n_code_links":0,"syntology":null},{"paper":null,"slug":"an-investigation-on-different-underlying","title":"An Investigation on Different Underlying Quantization Schemes for Pre-trained Language Models","date":"2020-10-14","arxiv_id":"2010.07109","n_code_links":0,"syntology":null},{"paper":null,"slug":"da-transformer-distance-aware-transformer","title":"DA-Transformer: Distance-aware Transformer","date":"2020-10-14","arxiv_id":"2010.06925","n_code_links":0,"syntology":null},{"paper":null,"slug":"geometry-matters-exploring-language-examples-1","title":"Geometry matters: Exploring language examples at the decision boundary","date":"2020-10-14","arxiv_id":"2010.07212","n_code_links":0,"syntology":null},{"paper":"/paper/no-rumours-please-a-multi-indic-lingual","slug":"no-rumours-please-a-multi-indic-lingual","title":"No Rumours Please! A Multi-Indic-Lingual Approach for COVID Fake-Tweet Detection","date":"2020-10-14","arxiv_id":"2010.06906","n_code_links":1,"syntology":null},{"paper":"/paper/aspect-based-document-similarity-for-research","slug":"aspect-based-document-similarity-for-research","title":"Aspect-based Document Similarity for Research Papers","date":"2020-10-13","arxiv_id":"2010.06395","n_code_links":1,"syntology":{"ran":2,"of":3,"n_ran_checked":2,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["malteos/aspect-document-similarity"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/bert-emd-many-to-many-layer-mapping-for-bert","slug":"bert-emd-many-to-many-layer-mapping-for-bert","title":"BERT-EMD: Many-to-Many Layer Mapping for BERT Compression with Earth Mover's Distance","date":"2020-10-13","arxiv_id":"2010.06133","n_code_links":1,"syntology":null},{"paper":null,"slug":"capt-contrastive-pre-training-for","title":"CAPT: Contrastive Pre-Training for Learning Denoised Sequence Representations","date":"2020-10-13","arxiv_id":"2010.06351","n_code_links":0,"syntology":null},{"paper":null,"slug":"improving-text-generation-evaluation-with","title":"Improving Text Generation Evaluation with Batch Centering and Tempered Word Mover Distance","date":"2020-10-13","arxiv_id":"2010.06150","n_code_links":0,"syntology":null},{"paper":"/paper/incorporating-bert-into-parallel-sequence","slug":"incorporating-bert-into-parallel-sequence","title":"Incorporating BERT into Parallel Sequence Decoding with Adapters","date":"2020-10-13","arxiv_id":"2010.06138","n_code_links":1,"syntology":{"ran":8,"of":13,"n_ran_checked":5,"n_instrument":3,"unverified":5,"pointer_only":3,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 1 honoured, 0 violated, 4 with no contract checked; 3 where Syntology's instrument failed) · 5 unverified","official":{"repos":["lemmonation/abnet"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":5,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"interpreting-attention-models-with-human","title":"Interpreting Attention Models with Human Visual Attention in Machine Reading Comprehension","date":"2020-10-13","arxiv_id":"2010.06396","n_code_links":0,"syntology":null},{"paper":null,"slug":"multilingual-argument-mining-datasets-and","title":"Multilingual Argument Mining: Datasets and Analysis","date":"2020-10-13","arxiv_id":"2010.06432","n_code_links":0,"syntology":null},{"paper":"/paper/pretrained-transformers-for-text-ranking-bert","slug":"pretrained-transformers-for-text-ranking-bert","title":"Pretrained Transformers for Text Ranking: BERT and Beyond","date":"2020-10-13","arxiv_id":"2010.06467","n_code_links":1,"syntology":null},{"paper":"/paper/probing-for-multilingual-numerical","slug":"probing-for-multilingual-numerical","title":"Probing for Multilingual Numerical Understanding in Transformer-Based Language Models","date":"2020-10-13","arxiv_id":"2010.06666","n_code_links":1,"syntology":null},{"paper":null,"slug":"chatbot-interaction-with-artificial","title":"Chatbot Interaction with Artificial Intelligence: Human Data Augmentation with T5 and Language Transformer Ensemble for Text Classification","date":"2020-10-12","arxiv_id":"2010.05990","n_code_links":0,"syntology":null},{"paper":"/paper/counterfactual-variable-control-for-robust","slug":"counterfactual-variable-control-for-robust","title":"Counterfactual Variable Control for Robust and Interpretable Question Answering","date":"2020-10-12","arxiv_id":"2010.05581","n_code_links":1,"syntology":null},{"paper":"/paper/cross-modal-bert-for-text-audio-sentiment","slug":"cross-modal-bert-for-text-audio-sentiment","title":"Cross-Modal BERT for Text-Audio Sentiment Analysis","date":"2020-10-12","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"efsg-evolutionary-fooling-sentences-generator","title":"EFSG: Evolutionary Fooling Sentences Generator","date":"2020-10-12","arxiv_id":"2010.05736","n_code_links":0,"syntology":null},{"paper":"/paper/from-hero-to-zeroe-a-benchmark-of-low-level","slug":"from-hero-to-zeroe-a-benchmark-of-low-level","title":"From Hero to Zéroe: A Benchmark of Low-Level Adversarial Attacks","date":"2020-10-12","arxiv_id":"2010.05648","n_code_links":1,"syntology":{"ran":2,"of":3,"n_ran_checked":0,"n_instrument":2,"unverified":1,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 1 unverified","official":{"repos":["yannikbenz/zeroe"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/huji-ku-at-mrp-2020-two-transition-based","slug":"huji-ku-at-mrp-2020-two-transition-based","title":"HUJI-KU at MRP~2020: Two Transition-based Neural Parsers","date":"2020-10-12","arxiv_id":"2010.05710","n_code_links":0,"syntology":null},{"paper":"/paper/improving-compositional-generalization-in","slug":"improving-compositional-generalization-in","title":"Improving Compositional Generalization in Semantic Parsing","date":"2020-10-12","arxiv_id":"2010.05647","n_code_links":1,"syntology":{"ran":2,"of":3,"n_ran_checked":2,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["inbaroren/improving-compgen-in-semparse"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"layer-wise-guided-training-for-bert-learning","title":"Layer-wise Guided Training for BERT: Learning Incrementally Refined Document Representations","date":"2020-10-12","arxiv_id":"2010.05763","n_code_links":0,"syntology":null},{"paper":"/paper/load-what-you-need-smaller-versions-of","slug":"load-what-you-need-smaller-versions-of","title":"Load What You Need: Smaller Versions of Multilingual BERT","date":"2020-10-12","arxiv_id":"2010.05609","n_code_links":2,"syntology":null},{"paper":null,"slug":"probing-pretrained-language-models-for","title":"Probing Pretrained Language Models for Lexical Semantics","date":"2020-10-12","arxiv_id":"2010.05731","n_code_links":0,"syntology":null},{"paper":"/paper/zero-shot-entity-linking-with-efficient-long","slug":"zero-shot-entity-linking-with-efficient-long","title":"Zero-shot Entity Linking with Efficient Long Range Sequence Modeling","date":"2020-10-12","arxiv_id":"2010.06065","n_code_links":1,"syntology":null},{"paper":null,"slug":"connecting-the-dots-between-fact-verification","title":"Connecting the Dots Between Fact Verification and Fake News Detection","date":"2020-10-11","arxiv_id":"2010.05202","n_code_links":0,"syntology":null},{"paper":"/paper/data-agnostic-roberta-based-natural-language","slug":"data-agnostic-roberta-based-natural-language","title":"Data Agnostic RoBERTa-based Natural Language to SQL Query Generation","date":"2020-10-11","arxiv_id":"2010.05243","n_code_links":1,"syntology":null},{"paper":null,"slug":"detecting-foodborne-illness-complaints-in","title":"Detecting Foodborne Illness Complaints in Multiple Languages Using English Annotations Only","date":"2020-10-11","arxiv_id":"2010.05194","n_code_links":0,"syntology":null},{"paper":"/paper/incremental-processing-in-the-age-of-non","slug":"incremental-processing-in-the-age-of-non","title":"Incremental Processing in the Age of Non-Incremental Encoders: An Empirical Assessment of Bidirectional Models for Incremental NLU","date":"2020-10-11","arxiv_id":"2010.05330","n_code_links":1,"syntology":null},{"paper":"/paper/learning-which-features-matter-roberta","slug":"learning-which-features-matter-roberta","title":"Learning Which Features Matter: RoBERTa Acquires a Preference for Linguistic Generalizations (Eventually)","date":"2020-10-11","arxiv_id":"2010.05358","n_code_links":1,"syntology":null}],"record_sha256":"6c8249425476a5e5be894a792d8f831b24a7c32ec22004b23d37b908a0387063","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}