{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/weight-decay/papers/101","list_of":"/method/weight-decay","method":"Weight Decay","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":101,"pages_in_order":108,"rows_per_page":100,"rows":[10001,10100],"of":10713,"counts":{"archive_papers_tagged":10713,"with_a_code_link":4533,"where_syntology_ran_a_sample":1291,"not_listed_spam_title":0,"listed":10713,"listed_where_code_ran":1291,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":1064,"every_run_a_failure_of_syntologys_instrument":227,"listed_with_a_run_with_no_instrument_failure":1064,"listed_every_run_a_failure_of_syntologys_instrument":227,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/weight-decay","prev":"/method/weight-decay/papers/100","next":"/method/weight-decay/papers/102","papers":[{"paper":"/paper/bert-is-not-a-knowledge-base-yet-factual","slug":"bert-is-not-a-knowledge-base-yet-factual","title":"E-BERT: Efficient-Yet-Effective Entity Embeddings for BERT","date":"2019-11-09","arxiv_id":"1911.03681","n_code_links":1,"syntology":null},{"paper":"/paper/convert-efficient-and-accurate-conversational","slug":"convert-efficient-and-accurate-conversational","title":"ConveRT: Efficient and Accurate Conversational Representations from Transformers","date":"2019-11-09","arxiv_id":"1911.03688","n_code_links":5,"syntology":{"ran":8,"of":9,"n_ran_checked":8,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 0 violated, 8 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":null}},{"paper":null,"slug":"multi-perspective-inferrer-reasoning","title":"Multi-Perspective Inferrer: Reasoning Sentences Relationship from Holistic Perspective","date":"2019-11-09","arxiv_id":"1911.03668","n_code_links":0,"syntology":null},{"paper":null,"slug":"the-dialogue-dodecathlon-open-domain","title":"The Dialogue Dodecathlon: Open-Domain Knowledge and Image Grounded Conversational Agents","date":"2019-11-09","arxiv_id":"1911.03768","n_code_links":0,"syntology":null},{"paper":null,"slug":"zero-shot-paraphrase-generation-with","title":"Zero-Shot Paraphrase Generation with Multilingual Language Models","date":"2019-11-09","arxiv_id":"1911.03597","n_code_links":0,"syntology":null},{"paper":null,"slug":"cross-lingual-relevance-transfer-for-document","title":"Cross-Lingual Relevance Transfer for Document Retrieval","date":"2019-11-08","arxiv_id":"1911.02989","n_code_links":0,"syntology":null},{"paper":"/paper/graph-to-graph-transformer-for-transition","slug":"graph-to-graph-transformer-for-transition","title":"Graph-to-Graph Transformer for Transition-based Dependency Parsing","date":"2019-11-08","arxiv_id":"1911.03561","n_code_links":1,"syntology":null},{"paper":"/paper/how-language-neutral-is-multilingual-bert","slug":"how-language-neutral-is-multilingual-bert","title":"How Language-Neutral is Multilingual BERT?","date":"2019-11-08","arxiv_id":"1911.03310","n_code_links":1,"syntology":null},{"paper":null,"slug":"pretrained-language-models-for-document-level","title":"Pretrained Language Models for Document-Level Neural Machine Translation","date":"2019-11-08","arxiv_id":"1911.03110","n_code_links":0,"syntology":null},{"paper":null,"slug":"resurrecting-submodularity-in-neural","title":"Resurrecting Submodularity for Neural Text Generation","date":"2019-11-08","arxiv_id":"1911.03014","n_code_links":0,"syntology":null},{"paper":null,"slug":"sept-improving-scientific-named-entity","title":"SEPT: Improving Scientific Named Entity Recognition with Span Representation","date":"2019-11-08","arxiv_id":"1911.03353","n_code_links":0,"syntology":null},{"paper":"/paper/towards-hierarchical-importance-attribution-1","slug":"towards-hierarchical-importance-attribution-1","title":"Towards Hierarchical Importance Attribution: Explaining Compositional Semantics for Neural Sequence Models","date":"2019-11-08","arxiv_id":"1911.06194","n_code_links":3,"syntology":null},{"paper":null,"slug":"transforming-wikipedia-into-augmented-data","title":"Transforming Wikipedia into Augmented Data for Query-Focused Summarization","date":"2019-11-08","arxiv_id":"1911.03324","n_code_links":0,"syntology":null},{"paper":null,"slug":"what-would-elsa-do-freezing-layers-during","title":"What Would Elsa Do? Freezing Layers During Transformer Fine-Tuning","date":"2019-11-08","arxiv_id":"1911.03090","n_code_links":0,"syntology":null},{"paper":"/paper/berts-of-a-feather-do-not-generalize-together","slug":"berts-of-a-feather-do-not-generalize-together","title":"BERTs of a feather do not generalize together: Large variability in generalization across models with similar test set performance","date":"2019-11-07","arxiv_id":"1911.02969","n_code_links":1,"syntology":null},{"paper":"/paper/blockwise-self-attention-for-long-document","slug":"blockwise-self-attention-for-long-document","title":"Blockwise Self-Attention for Long Document Understanding","date":"2019-11-07","arxiv_id":"1911.02972","n_code_links":1,"syntology":null},{"paper":"/paper/conversation-generation-with-concept-flow","slug":"conversation-generation-with-concept-flow","title":"Grounded Conversation Generation as Guided Traverses in Commonsense Knowledge Graphs","date":"2019-11-07","arxiv_id":"1911.02707","n_code_links":2,"syntology":{"ran":6,"of":6,"n_ran_checked":6,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 1 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["thunlp/ConceptFlow"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"explicit-pairwise-word-interaction-modeling","title":"Explicit Pairwise Word Interaction Modeling Improves Pretrained Transformers for English Semantic Similarity Tasks","date":"2019-11-07","arxiv_id":"1911.02847","n_code_links":0,"syntology":null},{"paper":null,"slug":"the-lig-system-for-the-english-czech-text","title":"The LIG system for the English-Czech Text Translation Task of IWSLT 2019","date":"2019-11-07","arxiv_id":"1911.02898","n_code_links":0,"syntology":null},{"paper":null,"slug":"using-dynamic-embeddings-to-improve-static","title":"How Can BERT Help Lexical Semantics Tasks?","date":"2019-11-07","arxiv_id":"1911.02929","n_code_links":0,"syntology":null},{"paper":null,"slug":"learning-to-answer-by-learning-to-ask-getting","title":"Learning to Answer by Learning to Ask: Getting the Best of GPT-2 and BERT Worlds","date":"2019-11-06","arxiv_id":"1911.02365","n_code_links":0,"syntology":null},{"paper":null,"slug":"unsupervised-domain-adaptation-of-contextual","title":"Unsupervised Domain Adaptation of Contextual Embeddings for Low-Resource Duplicate Question Detection","date":"2019-11-06","arxiv_id":"1911.02645","n_code_links":0,"syntology":null},{"paper":null,"slug":"deepening-hidden-representations-from-pre","title":"Deepening Hidden Representations from Pre-trained Language Models","date":"2019-11-05","arxiv_id":"1911.01940","n_code_links":0,"syntology":null},{"paper":"/paper/improving-slot-filling-by-utilizing","slug":"improving-slot-filling-by-utilizing","title":"Improving Slot Filling by Utilizing Contextual Information","date":"2019-11-05","arxiv_id":"1911.01680","n_code_links":0,"syntology":null},{"paper":null,"slug":"incremental-sense-weight-training-for-the","title":"Incremental Sense Weight Training for the Interpretation of Contextualized Word Embeddings","date":"2019-11-05","arxiv_id":"1911.01623","n_code_links":0,"syntology":null},{"paper":"/paper/mml-maximal-multiverse-learning-for-robust","slug":"mml-maximal-multiverse-learning-for-robust","title":"MML: Maximal Multiverse Learning for Robust Fine-Tuning of Language Models","date":"2019-11-05","arxiv_id":"1911.06182","n_code_links":1,"syntology":null},{"paper":"/paper/unsupervised-cross-lingual-representation-1","slug":"unsupervised-cross-lingual-representation-1","title":"Unsupervised Cross-lingual Representation Learning at Scale","date":"2019-11-05","arxiv_id":"1911.02116","n_code_links":35,"syntology":{"ran":40,"of":59,"n_ran_checked":35,"n_instrument":5,"unverified":19,"pointer_only":52,"phrase":"40 ran (of which 13 constructed an object rather than computing a result; 35 with no instrument failure: 4 honoured, 1 violated, 30 with no contract checked; 5 where Syntology's instrument failed) · 19 unverified","official":{"repos":["facebookresearch/cc_net","facebookresearch/XLM"],"state":"official (archive's flag): 12 ran","n_ran":12,"n_constructed":0,"n_ran_no_instrument_failure":12,"n_unverified":3,"ran_from_kinds":["listed","official","unlocated"]}}},{"paper":"/paper/assessing-social-and-intersectional-biases-in","slug":"assessing-social-and-intersectional-biases-in","title":"Assessing Social and Intersectional Biases in Contextualized Word Representations","date":"2019-11-04","arxiv_id":"1911.01485","n_code_links":1,"syntology":null},{"paper":null,"slug":"bas-an-answer-selection-method-using-bert","title":"BAS: An Answer Selection Method Using BERT Language Model","date":"2019-11-04","arxiv_id":"1911.01528","n_code_links":0,"syntology":null},{"paper":null,"slug":"sentence-level-bert-and-multi-task-learning","title":"Sentence-Level BERT and Multi-Task Learning of Age and Gender in Social Media","date":"2019-11-02","arxiv_id":"1911.00637","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-deep-learning-based-system-for-pharmaconer","title":"A Deep Learning-Based System for PharmaCoNER","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"a-multi-task-learning-framework-for-2","title":"A Multi-Task Learning Framework for Extracting Bacteria Biotope Information","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/a-recurrent-bert-based-model-for-question","slug":"a-recurrent-bert-based-model-for-question","title":"A Recurrent BERT-based Model for Question Generation","date":"2019-11-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"aggregating-bidirectional-encoder","title":"Aggregating Bidirectional Encoder Representations Using MatchLSTM for Sequence Matching","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"applying-bert-to-document-retrieval-with","title":"Applying BERT to Document Retrieval with Birch","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"bert-goes-to-law-school-quantifying-the","title":"BERT Goes to Law School: Quantifying the Competitive Advantage of Access to Large Legal Corpora in Contract Understanding","date":"2019-11-01","arxiv_id":"1911.00473","n_code_links":0,"syntology":null},{"paper":"/paper/bert-is-not-an-interlingua-and-the-bias-of","slug":"bert-is-not-an-interlingua-and-the-bias-of","title":"BERT is Not an Interlingua and the Bias of Tokenization","date":"2019-11-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/biomedical-named-entity-recognition-with","slug":"biomedical-named-entity-recognition-with","title":"Biomedical Named Entity Recognition with Multilingual BERT","date":"2019-11-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"blcu-nlp-at-coin-shared-task1-stagewise-fine","title":"BLCU-NLP at COIN-Shared Task1: Stagewise Fine-tuning BERT for Commonsense Inference in Everyday Narrations","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"caunlp-at-nlp4if-2019-shared-task-context","title":"CAUnLP at NLP4IF 2019 Shared Task: Context-Dependent BERT for Sentence-Level Propaganda Detection","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"cler-cross-task-learning-with-expert","title":"CLER: Cross-task Learning with Expert Representation to Generalize Reading and Understanding","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"combining-unsupervised-pre-training-and","title":"Combining Unsupervised Pre-training and Annotator Rationales to Improve Low-shot Text Classification","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"contextualized-cross-lingual-event-trigger","title":"Contextualized Cross-Lingual Event Trigger Extraction with Minimal Resources","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"cost-sensitive-bert-for-generalisable","title":"Cost-Sensitive BERT for Generalisable Sentence Classification on Imbalanced Data","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"cross-domain-modeling-of-sentence-level","title":"Cross-Domain Modeling of Sentence-Level Evidence for Document Retrieval","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"deep-bidirectional-transformers-for-relation-1","title":"Deep Bidirectional Transformers for Relation Extraction without Supervision","date":"2019-11-01","arxiv_id":"1911.00313","n_code_links":0,"syntology":null},{"paper":null,"slug":"detection-of-propaganda-using-logistic","title":"Detection of Propaganda Using Logistic Regression","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"divisive-language-and-propaganda-detection","title":"Divisive Language and Propaganda Detection using Multi-head Attention Transformers with Deep Learning BERT-based Language Models for Binary Classification","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"domain-adaptation-with-bert-based-domain","title":"Domain Adaptation with BERT-based Domain Classification and Data Selection","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"enhancing-bert-for-lexical-normalization","title":"Enhancing BERT for Lexical Normalization","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"evaluating-bert-for-natural-language","title":"Evaluating BERT for natural language inference: A case study on the CommitmentBank","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"event-causality-recognition-exploiting","title":"Event Causality Recognition Exploiting Multiple Annotators' Judgments and Background Knowledge","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"extract-and-aggregate-a-novel-domain","title":"Extract and Aggregate: A Novel Domain-Independent Approach to Factual Data Verification","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"extractive-narrativeqa-with-heuristic-pre","title":"Extractive NarrativeQA with Heuristic Pre-Training","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/faspell-a-fast-adaptable-simple-powerful","slug":"faspell-a-fast-adaptable-simple-powerful","title":"FASPell: A Fast, Adaptable, Simple, Powerful Chinese Spell Checker Based On DAE-Decoder Paradigm","date":"2019-11-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"fine-grained-propaganda-detection-with-fine","title":"Fine-Grained Propaganda Detection with Fine-Tuned BERT","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"fine-tune-bert-with-sparse-self-attention","title":"Fine-tune BERT with Sparse Self-Attention Mechanism","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"from-monolingual-to-multilingual-faq","title":"From Monolingual to Multilingual FAQ Assistant using Multilingual Co-training","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"fully-unsupervised-crosslingual-semantic","title":"Fully Unsupervised Crosslingual Semantic Textual Similarity Metric Based on BERT for Identifying Parallel Data","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"gem-generative-enhanced-model-for-adversarial","title":"GEM: Generative Enhanced Model for adversarial attacks","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/hit-scir-at-mrp-2019-a-unified-pipeline-for","slug":"hit-scir-at-mrp-2019-a-unified-pipeline-for","title":"HIT-SCIR at MRP 2019: A Unified Pipeline for Meaning Representation Parsing via Efficient Training and Effective Encoding","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"how-well-do-nli-models-capture-verb","title":"How well do NLI models capture verb veridicality?","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"improving-pre-trained-multilingual-model-with","title":"Improving Pre-Trained Multilingual Model with Vocabulary Expansion","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"inspecting-unification-of-encoding-and","title":"Inspecting Unification of Encoding and Matching with Transformer: A Case Study of Machine Reading Comprehension","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"jbnu-at-mrp-2019-multi-level-biaffine","title":"JBNU at MRP 2019: Multi-level Biaffine Attention for Semantic Dependency Parsing","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"jeff-da-at-coin-shared-task","title":"Jeff Da at COIN - Shared Task: BIG MOOD: Relating Transformers to Explicit Commonsense Knowledge","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"justdeep-at-nlp4if-2019-task-1-propaganda","title":"JUSTDeep at NLP4IF 2019 Task 1: Propaganda Detection using Ensemble Deep Learning Models","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"lexicalat-lexical-based-adversarial","title":"LexicalAT: Lexical-Based Adversarial Reinforcement Training for Robust Sentiment Classification","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"multi-view-domain-adapted-sentence-embeddings","title":"Multi-View Domain Adapted Sentence Embeddings for Low-Resource Unsupervised Duplicate Question Detection","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"named-entity-recognition-is-there-a-glass-1","title":"Named Entity Recognition - Is There a Glass Ceiling?","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/natural-language-generation-for-effective","slug":"natural-language-generation-for-effective","title":"Natural Language Generation for Effective Knowledge Distillation","date":"2019-11-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"no-youre-not-alone-a-better-way-to-find","title":"No, you're not alone: A better way to find people with similar experiences on Reddit","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"nsitnlp4if-2019-propaganda-detection-from","title":"NSIT@NLP4IF-2019: Propaganda Detection from News Articles using Transfer Learning","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/pingan-smart-health-and-sjtu-at-coin-shared","slug":"pingan-smart-health-and-sjtu-at-coin-shared","title":"Pingan Smart Health and SJTU at COIN - Shared Task: utilizing Pre-trained Language Models and Common-sense Knowledge in Machine Reading Tasks","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"policy-preference-detection-in-parliamentary","title":"Policy Preference Detection in Parliamentary Debate Motions","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"pre-training-bert-on-domain-resources-for","title":"Pre-Training BERT on Domain Resources for Short Answer Grading","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"question-answering-using-hierarchical","title":"Question Answering Using Hierarchical Attention on Top of BERT Features","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"recycling-a-pre-trained-bert-encoder-for","title":"Recycling a Pre-trained BERT Encoder for Neural Machine Translation","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"reevaluating-argument-component-extraction-in","title":"Reevaluating Argument Component Extraction in Low Resource Settings","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"relation-module-for-non-answerable","title":"Relation Module for Non-Answerable Predictions on Reading Comprehension","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"selecting-planning-and-rewriting-a-modular","title":"Selecting, Planning, and Rewriting: A Modular Approach for Data-to-Document Generation and Translation","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"sentence-level-propaganda-detection-in-news","title":"Sentence-Level Propaganda Detection in News Articles with Transfer Learning and BERT-BiLSTM-Capsule Model","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"suda-alibaba-at-mrp-2019-graph-based-models","title":"SUDA-Alibaba at MRP 2019: Graph-Based Models with BERT","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"sum-qe-a-bert-based-summary-quality","title":"SUM-QE: a BERT-based Summary Quality Estimation Model","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"team-domlin-exploiting-evidence-enhancement","title":"Team DOMLIN: Exploiting Evidence Enhancement for the FEVER Shared Task","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"transfer-learning-in-biomedical-named-entity","title":"Transfer Learning in Biomedical Named Entity Recognition: An Evaluation of BERT in the PharmaCoNER task","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/tupa-at-mrp-2019-a-multi-task-baseline-system","slug":"tupa-at-mrp-2019-a-multi-task-baseline-system","title":"TUPA at MRP 2019: A Multi-Task Baseline System","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"unsupervised-labeled-parsing-with-deep-inside","title":"Unsupervised Labeled Parsing with Deep Inside-Outside Recursive Autoencoders","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"what-does-this-word-mean-explaining","title":"What Does This Word Mean? Explaining Contextualized Embeddings with Natural Language Definition","date":"2019-11-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"when-choosing-plausible-alternatives-clever-1","title":"When Choosing Plausible Alternatives, Clever Hans can be Clever","date":"2019-11-01","arxiv_id":"1911.00225","n_code_links":0,"syntology":null},{"paper":null,"slug":"dianet-bert-and-hierarchical-attention-multi","title":"DiaNet: BERT and Hierarchical Attention Multi-Task Learning of Fine-Grained Dialect","date":"2019-10-31","arxiv_id":"1910.14243","n_code_links":0,"syntology":null},{"paper":"/paper/do-multi-hop-readers-dream-of-reasoning","slug":"do-multi-hop-readers-dream-of-reasoning","title":"Do Multi-hop Readers Dream of Reasoning Chains?","date":"2019-10-31","arxiv_id":"1910.14520","n_code_links":1,"syntology":{"ran":9,"of":11,"n_ran_checked":6,"n_instrument":3,"unverified":2,"pointer_only":2,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 2 honoured, 1 violated, 3 with no contract checked; 3 where Syntology's instrument failed) · 2 unverified","official":{"repos":["helloeve/bert-co-matching"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"human-centric-metric-for-accelerating","title":"Human-centric Metric for Accelerating Pathology Reports Annotation","date":"2019-10-31","arxiv_id":"1911.01226","n_code_links":0,"syntology":null},{"paper":null,"slug":"limit-bert-linguistic-informed-multi-task","title":"LIMIT-BERT : Linguistic Informed Multi-Task BERT","date":"2019-10-31","arxiv_id":"1910.14296","n_code_links":0,"syntology":null},{"paper":"/paper/multi-stage-document-ranking-with-bert","slug":"multi-stage-document-ranking-with-bert","title":"Multi-Stage Document Ranking with BERT","date":"2019-10-31","arxiv_id":"1910.14424","n_code_links":3,"syntology":{"ran":4,"of":5,"n_ran_checked":4,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":null}},{"paper":null,"slug":"positional-attention-based-frame","title":"Positional Attention-based Frame Identification with BERT: A Deep Learning Approach to Target Disambiguation and Semantic Frame Selection","date":"2019-10-31","arxiv_id":"1910.14549","n_code_links":0,"syntology":null},{"paper":"/paper/pseudolikelihood-reranking-with-masked","slug":"pseudolikelihood-reranking-with-masked","title":"Masked Language Model Scoring","date":"2019-10-31","arxiv_id":"1910.14659","n_code_links":6,"syntology":{"ran":10,"of":10,"n_ran_checked":7,"n_instrument":3,"unverified":0,"pointer_only":3,"phrase":"10 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 2 honoured, 0 violated, 5 with no contract checked; 3 where Syntology's instrument failed) · 0 unverified","official":{"repos":["awslabs/mlm-scoring"],"state":"official: no sample here; runs from other or unrecorded repositories","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed","unlocated"]}}},{"paper":null,"slug":"transfer-learning-from-transformers-to-fake","title":"Transfer Learning from Transformers to Fake News Challenge Stance Detection (FNC-1) Task","date":"2019-10-31","arxiv_id":"1910.14353","n_code_links":0,"syntology":null},{"paper":"/paper/discourse-aware-neural-extractive-model-for","slug":"discourse-aware-neural-extractive-model-for","title":"Discourse-Aware Neural Extractive Text Summarization","date":"2019-10-30","arxiv_id":"1910.14142","n_code_links":1,"syntology":{"ran":10,"of":11,"n_ran_checked":10,"n_instrument":0,"unverified":1,"pointer_only":1,"phrase":"10 ran (of which 0 constructed an object rather than computing a result; 10 with no instrument failure: 0 honoured, 0 violated, 10 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["jiacheng-xu/DiscoBERT"],"state":"official (archive's flag): 10 ran","n_ran":10,"n_constructed":0,"n_ran_no_instrument_failure":10,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"lsh-sampling-breaks-the-computation-chicken","title":"Lsh-sampling Breaks the Computation Chicken-and-egg Loop in Adaptive Stochastic Gradient Estimation","date":"2019-10-30","arxiv_id":"1910.14162","n_code_links":0,"syntology":null}],"record_sha256":"95cea870933463437fa31da25fccf19c98b4bb35e3fe5026f210ea7a58fec2f4","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}