{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/bert/papers/57","list_of":"/method/bert","method":"BERT","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":57,"pages_in_order":70,"rows_per_page":100,"rows":[5601,5700],"of":6938,"counts":{"archive_papers_tagged":6938,"with_a_code_link":2862,"where_syntology_ran_a_sample":640,"not_listed_spam_title":0,"listed":6938,"listed_where_code_ran":640,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":520,"every_run_a_failure_of_syntologys_instrument":120,"listed_with_a_run_with_no_instrument_failure":520,"listed_every_run_a_failure_of_syntologys_instrument":120,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/bert","prev":"/method/bert/papers/56","next":"/method/bert/papers/58","papers":[{"paper":"/paper/where-are-the-facts-searching-for-fact","slug":"where-are-the-facts-searching-for-fact","title":"Where Are the Facts? Searching for Fact-checked Information to Alleviate the Spread of Fake News","date":"2020-10-07","arxiv_id":"2010.03159","n_code_links":2,"syntology":{"ran":5,"of":5,"n_ran_checked":5,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["nguyenvo09/EMNLP2020"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/why-do-you-think-that-exploring-faithful","slug":"why-do-you-think-that-exploring-faithful","title":"Why do you think that? Exploring Faithful Sentence-Level Rationales Without Supervision","date":"2020-10-07","arxiv_id":"2010.03384","n_code_links":1,"syntology":null},{"paper":"/paper/analyzing-individual-neurons-in-pre-trained","slug":"analyzing-individual-neurons-in-pre-trained","title":"Analyzing Individual Neurons in Pre-trained Language Models","date":"2020-10-06","arxiv_id":"2010.02695","n_code_links":1,"syntology":null},{"paper":"/paper/bert-knows-punta-cana-is-not-just-beautiful","slug":"bert-knows-punta-cana-is-not-just-beautiful","title":"BERT Knows Punta Cana is not just beautiful, it's gorgeous: Ranking Scalar Adjectives with Contextualised Representations","date":"2020-10-06","arxiv_id":"2010.02686","n_code_links":1,"syntology":null},{"paper":"/paper/cross-lingual-text-classification-with","slug":"cross-lingual-text-classification-with","title":"Cross-Lingual Text Classification with Minimal Resources by Transferring a Sparse Teacher","date":"2020-10-06","arxiv_id":"2010.02562","n_code_links":1,"syntology":null},{"paper":"/paper/do-explicit-alignments-robustly-improve","slug":"do-explicit-alignments-robustly-improve","title":"Do Explicit Alignments Robustly Improve Multilingual Encoders?","date":"2020-10-06","arxiv_id":"2010.02537","n_code_links":1,"syntology":null},{"paper":"/paper/exploring-bert-s-sensitivity-to-lexical-cues","slug":"exploring-bert-s-sensitivity-to-lexical-cues","title":"Exploring BERT's Sensitivity to Lexical Cues using Tests from Semantic Priming","date":"2020-10-06","arxiv_id":"2010.03010","n_code_links":1,"syntology":{"ran":2,"of":3,"n_ran_checked":2,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["kanishkamisra/emnlp-bert-priming"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/improving-efficient-neural-ranking-models","slug":"improving-efficient-neural-ranking-models","title":"Improving Efficient Neural Ranking Models with Cross-Architecture Knowledge Distillation","date":"2020-10-06","arxiv_id":"2010.02666","n_code_links":1,"syntology":null},{"paper":null,"slug":"incorporating-behavioral-hypotheses-for-query","title":"Incorporating Behavioral Hypotheses for Query Generation","date":"2020-10-06","arxiv_id":"2010.02667","n_code_links":0,"syntology":null},{"paper":"/paper/intrinsic-probing-through-dimension-selection","slug":"intrinsic-probing-through-dimension-selection","title":"Intrinsic Probing through Dimension Selection","date":"2020-10-06","arxiv_id":"2010.02812","n_code_links":1,"syntology":{"ran":2,"of":3,"n_ran_checked":2,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["rycolab/intrinsic-probing"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"legal-bert-the-muppets-straight-out-of-law","title":"LEGAL-BERT: The Muppets straight out of Law School","date":"2020-10-06","arxiv_id":"2010.02559","n_code_links":0,"syntology":null},{"paper":"/paper/neural-mask-generator-learning-to-generate","slug":"neural-mask-generator-learning-to-generate","title":"Neural Mask Generator: Learning to Generate Adaptive Word Maskings for Language Model Adaptation","date":"2020-10-06","arxiv_id":"2010.02705","n_code_links":1,"syntology":null},{"paper":null,"slug":"on-the-interplay-between-fine-tuning-and","title":"On the Interplay Between Fine-tuning and Sentence-level Probing for Linguistic Knowledge in Pre-trained Transformers","date":"2020-10-06","arxiv_id":"2010.02616","n_code_links":0,"syntology":null},{"paper":"/paper/poison-attacks-against-text-datasets-with","slug":"poison-attacks-against-text-datasets-with","title":"Poison Attacks against Text Datasets with Conditional Adversarially Regularized Autoencoder","date":"2020-10-06","arxiv_id":"2010.02684","n_code_links":2,"syntology":{"ran":2,"of":2,"n_ran_checked":1,"n_instrument":1,"unverified":0,"pointer_only":2,"phrase":"2 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["alvinchangw/CARA_EMNLP2020"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/the-multilingual-amazon-reviews-corpus","slug":"the-multilingual-amazon-reviews-corpus","title":"The Multilingual Amazon Reviews Corpus","date":"2020-10-06","arxiv_id":"2010.02573","n_code_links":1,"syntology":null},{"paper":null,"slug":"how-effective-is-task-agnostic-data","title":"How Effective is Task-Agnostic Data Augmentation for Pretrained Transformers?","date":"2020-10-05","arxiv_id":"2010.01764","n_code_links":0,"syntology":null},{"paper":"/paper/improving-amr-parsing-with-sequence-to","slug":"improving-amr-parsing-with-sequence-to","title":"Improving AMR Parsing with Sequence-to-Sequence Pre-training","date":"2020-10-05","arxiv_id":"2010.01771","n_code_links":1,"syntology":null},{"paper":"/paper/infobert-improving-robustness-of-language-1","slug":"infobert-improving-robustness-of-language-1","title":"InfoBERT: Improving Robustness of Language Models from An Information Theoretic Perspective","date":"2020-10-05","arxiv_id":"2010.02329","n_code_links":2,"syntology":{"ran":4,"of":4,"n_ran_checked":3,"n_instrument":1,"unverified":0,"pointer_only":4,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 2 honoured, 0 violated, 1 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["AI-secure/InfoBERT"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"linguistic-profiling-of-a-neural-language","title":"Linguistic Profiling of a Neural Language Model","date":"2020-10-05","arxiv_id":"2010.01869","n_code_links":0,"syntology":null},{"paper":null,"slug":"mixup-transfomer-dynamic-data-augmentation","title":"Mixup-Transformer: Dynamic Data Augmentation for NLP Tasks","date":"2020-10-05","arxiv_id":"2010.02394","n_code_links":0,"syntology":null},{"paper":null,"slug":"pair-planning-and-iterative-refinement-in-pre","title":"PAIR: Planning and Iterative Refinement in Pre-trained Transformers for Long Text Generation","date":"2020-10-05","arxiv_id":"2010.02301","n_code_links":0,"syntology":null},{"paper":"/paper/pareto-probing-trading-off-accuracy-for","slug":"pareto-probing-trading-off-accuracy-for","title":"Pareto Probing: Trading Off Accuracy for Complexity","date":"2020-10-05","arxiv_id":"2010.02180","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":1,"n_instrument":2,"unverified":0,"pointer_only":3,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":{"repos":["rycolab/pareto-probing"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/pmi-masking-principled-masking-of-correlated-1","slug":"pmi-masking-principled-masking-of-correlated-1","title":"PMI-Masking: Principled masking of correlated spans","date":"2020-10-05","arxiv_id":"2010.01825","n_code_links":1,"syntology":null},{"paper":"/paper/pruning-redundant-mappings-in-transformer","slug":"pruning-redundant-mappings-in-transformer","title":"Pruning Redundant Mappings in Transformer Models via Spectral-Normalized Identity Prior","date":"2020-10-05","arxiv_id":"2010.01791","n_code_links":1,"syntology":null},{"paper":null,"slug":"pum-at-semeval-2020-task-12-aggregation-of","title":"PUM at SemEval-2020 Task 12: Aggregation of Transformer-based models' features for offensive language recognition","date":"2020-10-05","arxiv_id":"2010.01897","n_code_links":0,"syntology":null},{"paper":"/paper/self-training-improves-pre-training-for","slug":"self-training-improves-pre-training-for","title":"Self-training Improves Pre-training for Natural Language Understanding","date":"2020-10-05","arxiv_id":"2010.02194","n_code_links":1,"syntology":null},{"paper":"/paper/unsupervised-reference-free-summary-quality","slug":"unsupervised-reference-free-summary-quality","title":"Unsupervised Reference-Free Summary Quality Evaluation via Contrastive Learning","date":"2020-10-05","arxiv_id":"2010.01781","n_code_links":1,"syntology":null},{"paper":"/paper/x-srl-a-parallel-cross-lingual-semantic-role","slug":"x-srl-a-parallel-cross-lingual-semantic-role","title":"X-SRL: A Parallel Cross-Lingual Semantic Role Labeling Dataset","date":"2020-10-05","arxiv_id":"2010.01998","n_code_links":1,"syntology":null},{"paper":"/paper/an-empirical-study-on-large-scale-multi-label","slug":"an-empirical-study-on-large-scale-multi-label","title":"An Empirical Study on Large-Scale Multi-Label Text Classification Including Few and Zero-Shot Labels","date":"2020-10-04","arxiv_id":"2010.01653","n_code_links":1,"syntology":{"ran":7,"of":8,"n_ran_checked":6,"n_instrument":1,"unverified":1,"pointer_only":1,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":{"repos":["iliaschalkidis/lmtc-eurlex57k"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/on-losses-for-modern-language-models","slug":"on-losses-for-modern-language-models","title":"On Losses for Modern Language Models","date":"2020-10-04","arxiv_id":"2010.01694","n_code_links":1,"syntology":null},{"paper":"/paper/mining-knowledge-for-natural-language","slug":"mining-knowledge-for-natural-language","title":"Mining Knowledge for Natural Language Inference from Wikipedia Categories","date":"2020-10-03","arxiv_id":"2010.01239","n_code_links":1,"syntology":{"ran":1,"of":2,"n_ran_checked":1,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["ZeweiChu/WikiNLI"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"personality-trait-detection-using-bagged-svm","title":"Personality Trait Detection Using Bagged SVM over BERT Word Embedding Ensembles","date":"2020-10-03","arxiv_id":"2010.01309","n_code_links":0,"syntology":null},{"paper":"/paper/cost-effective-selection-of-pretraining-data","slug":"cost-effective-selection-of-pretraining-data","title":"Cost-effective Selection of Pretraining Data: A Case Study of Pretraining BERT on Social Media","date":"2020-10-02","arxiv_id":"2010.01150","n_code_links":0,"syntology":null},{"paper":null,"slug":"long-tail-zero-and-few-shot-learning-via","title":"Data-Efficient Pretraining via Contrastive Self-Supervision","date":"2020-10-02","arxiv_id":"2010.01061","n_code_links":0,"syntology":null},{"paper":"/paper/luke-deep-contextualized-entity","slug":"luke-deep-contextualized-entity","title":"LUKE: Deep Contextualized Entity Representations with Entity-aware Self-attention","date":"2020-10-02","arxiv_id":"2010.01057","n_code_links":9,"syntology":{"ran":3,"of":10,"n_ran_checked":3,"n_instrument":0,"unverified":7,"pointer_only":0,"phrase":"3 ran (of which 2 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 7 unverified","official":{"repos":["studio-ousia/luke"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":2,"n_ran_no_instrument_failure":3,"n_unverified":5,"ran_from_kinds":["official"]}}},{"paper":"/paper/multicqa-zero-shot-transfer-of-self","slug":"multicqa-zero-shot-transfer-of-self","title":"MultiCQA: Zero-Shot Transfer of Self-Supervised Text Matching Models on a Massive Scale","date":"2020-10-02","arxiv_id":"2010.00980","n_code_links":1,"syntology":null},{"paper":"/paper/stil-simultaneous-slot-filling-translation","slug":"stil-simultaneous-slot-filling-translation","title":"STIL -- Simultaneous Slot Filling, Translation, Intent Classification, and Language Identification: Initial Results using mBART on MultiATIS++","date":"2020-10-02","arxiv_id":"2010.00760","n_code_links":1,"syntology":null},{"paper":null,"slug":"beyond-the-text-analysis-of-privacy","title":"Beyond The Text: Analysis of Privacy Statements through Syntactic and Semantic Role Labeling","date":"2020-10-01","arxiv_id":"2010.00678","n_code_links":0,"syntology":null},{"paper":"/paper/colake-contextualized-language-and-knowledge","slug":"colake-contextualized-language-and-knowledge","title":"CoLAKE: Contextualized Language and Knowledge Embedding","date":"2020-10-01","arxiv_id":"2010.00309","n_code_links":1,"syntology":{"ran":6,"of":7,"n_ran_checked":6,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["txsun1997/CoLAKE"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"detecting-white-supremacist-hate-speech-using","title":"Detecting White Supremacist Hate Speech using Domain Specific Word Embedding with Deep Learning and BERT","date":"2020-10-01","arxiv_id":"2010.00357","n_code_links":0,"syntology":null},{"paper":null,"slug":"evaluating-multilingual-bert-for-estonian","title":"Evaluating Multilingual BERT for Estonian","date":"2020-10-01","arxiv_id":"2010.00454","n_code_links":0,"syntology":null},{"paper":"/paper/refvos-a-closer-look-at-referring-expressions","slug":"refvos-a-closer-look-at-referring-expressions","title":"RefVOS: A Closer Look at Referring Expressions for Video Object Segmentation","date":"2020-10-01","arxiv_id":"2010.00263","n_code_links":2,"syntology":null},{"paper":null,"slug":"rrf102-meeting-the-trec-covid-challenge-with","title":"RRF102: Meeting the TREC-COVID Challenge with a 100+ Runs Ensemble","date":"2020-10-01","arxiv_id":"2010.00200","n_code_links":0,"syntology":null},{"paper":"/paper/understanding-tables-with-intermediate-pre","slug":"understanding-tables-with-intermediate-pre","title":"Understanding tables with intermediate pre-training","date":"2020-10-01","arxiv_id":"2010.00571","n_code_links":1,"syntology":null},{"paper":"/paper/a-tale-of-two-linkings-dynamically-gating","slug":"a-tale-of-two-linkings-dynamically-gating","title":"A Tale of Two Linkings: Dynamically Gating between Schema Linking and Structural Linking for Text-to-SQL Parsing","date":"2020-09-30","arxiv_id":"2009.14809","n_code_links":1,"syntology":null},{"paper":"/paper/a-vietnamese-dataset-for-evaluating-machine","slug":"a-vietnamese-dataset-for-evaluating-machine","title":"A Vietnamese Dataset for Evaluating Machine Reading Comprehension","date":"2020-09-30","arxiv_id":"2009.14725","n_code_links":0,"syntology":null},{"paper":null,"slug":"auber-automated-bert-regularization","title":"AUBER: Automated BERT Regularization","date":"2020-09-30","arxiv_id":"2009.14409","n_code_links":0,"syntology":null},{"paper":"/paper/bert-for-monolingual-and-cross-lingual","slug":"bert-for-monolingual-and-cross-lingual","title":"BERT for Monolingual and Cross-Lingual Reverse Dictionary","date":"2020-09-30","arxiv_id":"2009.14790","n_code_links":1,"syntology":null},{"paper":null,"slug":"pea-kd-parameter-efficient-and-accurate","title":"Pea-KD: Parameter-efficient and Accurate Knowledge Distillation on BERT","date":"2020-09-30","arxiv_id":"2009.14822","n_code_links":0,"syntology":null},{"paper":"/paper/contrastive-distillation-on-intermediate","slug":"contrastive-distillation-on-intermediate","title":"Contrastive Distillation on Intermediate Representations for Language Model Compression","date":"2020-09-29","arxiv_id":"2009.14167","n_code_links":1,"syntology":{"ran":2,"of":4,"n_ran_checked":2,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["intersun/CoDIR"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"cross-lingual-alignment-methods-for","title":"Cross-lingual Alignment Methods for Multilingual BERT: A Comparative Study","date":"2020-09-29","arxiv_id":"2009.14304","n_code_links":0,"syntology":null},{"paper":null,"slug":"from-twitter-to-traffic-predictor-next-day","title":"From Twitter to Traffic Predictor: Next-Day Morning Traffic Prediction Using Social Media Data","date":"2020-09-29","arxiv_id":"2009.13794","n_code_links":0,"syntology":null},{"paper":null,"slug":"gender-prediction-using-limited-twitter-data","title":"Gender prediction using limited Twitter Data","date":"2020-09-29","arxiv_id":"2010.02005","n_code_links":0,"syntology":null},{"paper":"/paper/hint3-raising-the-bar-for-intent-detection-in","slug":"hint3-raising-the-bar-for-intent-detection-in","title":"HINT3: Raising the bar for Intent Detection in the Wild","date":"2020-09-29","arxiv_id":"2009.13833","n_code_links":1,"syntology":null},{"paper":null,"slug":"map-a-matrix-based-prediction-approach-to","title":"MaP: A Matrix-based Prediction Approach to Improve Span Extraction in Machine Reading Comprehension","date":"2020-09-29","arxiv_id":"2009.14348","n_code_links":0,"syntology":null},{"paper":null,"slug":"neural-retrieval-for-question-answering-with","title":"Neural Retrieval for Question Answering with Cross-Attention Supervised Data Augmentation","date":"2020-09-29","arxiv_id":"2009.13815","n_code_links":0,"syntology":null},{"paper":null,"slug":"test-positive-at-w-nut-2020-shared-task-3","title":"TEST_POSITIVE at W-NUT 2020 Shared Task-3: Joint Event Multi-task Learning for Slot Filling in Noisy Text","date":"2020-09-29","arxiv_id":"2009.14262","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-simple-and-efficient-ensemble-classifier","title":"A Simple and Efficient Ensemble Classifier Combining Multiple Neural Network Models on Social Media Datasets in Vietnamese","date":"2020-09-28","arxiv_id":"2009.13060","n_code_links":0,"syntology":null},{"paper":null,"slug":"accelerating-multi-model-inference-by-merging","title":"Accelerating Multi-Model Inference by Merging DNNs of Different Weights","date":"2020-09-28","arxiv_id":"2009.13062","n_code_links":0,"syntology":null},{"paper":"/paper/dialoglue-a-natural-language-understanding","slug":"dialoglue-a-natural-language-understanding","title":"DialoGLUE: A Natural Language Understanding Benchmark for Task-Oriented Dialogue","date":"2020-09-28","arxiv_id":"2009.13570","n_code_links":1,"syntology":null},{"paper":null,"slug":"fancy-man-lauches-zippo-at-wnut-2020-shared","title":"Fancy Man Lauches Zippo at WNUT 2020 Shared Task-1: A Bert Case Model for Wet Lab Entity Extraction","date":"2020-09-28","arxiv_id":"2009.12997","n_code_links":0,"syntology":null},{"paper":null,"slug":"knowledge-aware-procedural-text-understanding","title":"Knowledge-Aware Procedural Text Understanding with Multi-Stage Training","date":"2020-09-28","arxiv_id":"2009.13199","n_code_links":0,"syntology":null},{"paper":null,"slug":"pin-a-novel-parallel-interactive-network-for","title":"PIN: A Novel Parallel Interactive Network for Spoken Language Understanding","date":"2020-09-28","arxiv_id":"2009.13431","n_code_links":0,"syntology":null},{"paper":"/paper/ternarybert-distillation-aware-ultra-low-bit","slug":"ternarybert-distillation-aware-ultra-low-bit","title":"TernaryBERT: Distillation-aware Ultra-low Bit BERT","date":"2020-09-27","arxiv_id":"2009.12812","n_code_links":5,"syntology":null},{"paper":null,"slug":"metaphor-detection-using-deep-contextualized","title":"Metaphor Detection using Deep Contextualized Word Embeddings","date":"2020-09-26","arxiv_id":"2009.12565","n_code_links":0,"syntology":null},{"paper":null,"slug":"techniques-to-improve-q-a-accuracy-with","title":"Techniques to Improve Q&A Accuracy with Transformer-based models on Large Complex Documents","date":"2020-09-26","arxiv_id":"2009.12695","n_code_links":0,"syntology":null},{"paper":"/paper/a-little-goes-a-long-way-improving-toxic","slug":"a-little-goes-a-long-way-improving-toxic","title":"A little goes a long way: Improving toxic language classification despite data scarcity","date":"2020-09-25","arxiv_id":"2009.12344","n_code_links":1,"syntology":null},{"paper":"/paper/an-unsupervised-sentence-embedding-method","slug":"an-unsupervised-sentence-embedding-method","title":"An Unsupervised Sentence Embedding Method by Mutual Information Maximization","date":"2020-09-25","arxiv_id":"2009.12061","n_code_links":1,"syntology":null},{"paper":"/paper/bet-a-backtranslation-approach-for-easy-data","slug":"bet-a-backtranslation-approach-for-easy-data","title":"BET: A Backtranslation Approach for Easy Data Augmentation in Transformer-based Paraphrase Identification Context","date":"2020-09-25","arxiv_id":"2009.12452","n_code_links":1,"syntology":null},{"paper":"/paper/hetseq-distributed-gpu-training-on","slug":"hetseq-distributed-gpu-training-on","title":"HetSeq: Distributed GPU Training on Heterogeneous Infrastructure","date":"2020-09-25","arxiv_id":"2009.14783","n_code_links":1,"syntology":null},{"paper":"/paper/a-comparative-study-of-feature-types-for-age","slug":"a-comparative-study-of-feature-types-for-age","title":"A Comparative Study of Feature Types for Age-Based Text Classification","date":"2020-09-24","arxiv_id":"2009.11898","n_code_links":1,"syntology":null},{"paper":"/paper/adapting-bert-for-word-sense-disambiguation","slug":"adapting-bert-for-word-sense-disambiguation","title":"Adapting BERT for Word Sense Disambiguation with Gloss Selection Objective and Example Sentences","date":"2020-09-24","arxiv_id":"2009.11795","n_code_links":1,"syntology":{"ran":1,"of":2,"n_ran_checked":0,"n_instrument":1,"unverified":1,"pointer_only":2,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":{"repos":["BPYap/BERT-WSD"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/anchibert-a-pre-trained-model-for-ancient","slug":"anchibert-a-pre-trained-model-for-ancient","title":"AnchiBERT: A Pre-Trained Model for Ancient ChineseLanguage Understanding and Generation","date":"2020-09-24","arxiv_id":"2009.11473","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-token-wise-cnn-based-method-for-sentence","title":"A Token-wise CNN-based Method for Sentence Compression","date":"2020-09-23","arxiv_id":"2009.11260","n_code_links":0,"syntology":null},{"paper":null,"slug":"autorc-improving-bert-based-relation","title":"AutoRC: Improving BERT Based Relation Classification Models via Architecture Search","date":"2020-09-22","arxiv_id":"2009.10680","n_code_links":0,"syntology":null},{"paper":"/paper/constructing-interval-variables-via-faceted","slug":"constructing-interval-variables-via-faceted","title":"Constructing interval variables via faceted Rasch measurement and multitask deep learning: a hate speech application","date":"2020-09-22","arxiv_id":"2009.10277","n_code_links":2,"syntology":{"ran":3,"of":3,"n_ran_checked":3,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["ck37/coral-ordinal"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/grace-gradient-harmonized-and-cascaded","slug":"grace-gradient-harmonized-and-cascaded","title":"GRACE: Gradient Harmonized and Cascaded Labeling for Aspect-based Sentiment Analysis","date":"2020-09-22","arxiv_id":"2009.10557","n_code_links":1,"syntology":null},{"paper":null,"slug":"on-data-augmentation-for-extreme-multi-label","title":"On Data Augmentation for Extreme Multi-label Classification","date":"2020-09-22","arxiv_id":"2009.10778","n_code_links":0,"syntology":null},{"paper":"/paper/latin-bert-a-contextual-language-model-for","slug":"latin-bert-a-contextual-language-model-for","title":"Latin BERT: A Contextual Language Model for Classical Philology","date":"2020-09-21","arxiv_id":"2009.10053","n_code_links":1,"syntology":null},{"paper":"/paper/profile-consistency-identification-for-open","slug":"profile-consistency-identification-for-open","title":"Profile Consistency Identification for Open-domain Dialogue Agents","date":"2020-09-21","arxiv_id":"2009.09680","n_code_links":1,"syntology":{"ran":1,"of":3,"n_ran_checked":1,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"1 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified; the one sample that ran constructed an object rather than computing a result","official":{"repos":["songhaoyu/KvPI"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/ted-triple-supervision-decouples-end-to-end","slug":"ted-triple-supervision-decouples-end-to-end","title":"\"Listen, Understand and Translate\": Triple Supervision Decouples End-to-end Speech-to-text Translation","date":"2020-09-21","arxiv_id":"2009.09704","n_code_links":1,"syntology":null},{"paper":null,"slug":"when-they-say-weed-causes-depression-but-it-s","title":"\"When they say weed causes depression, but it's your fav antidepressant\": Knowledge-aware Attention Framework for Relationship Extraction","date":"2020-09-21","arxiv_id":"2009.10155","n_code_links":0,"syntology":null},{"paper":"/paper/dual-path-cnn-with-max-gated-block-for-text","slug":"dual-path-cnn-with-max-gated-block-for-text","title":"Dual-path CNN with Max Gated block for Text-Based Person Re-identification","date":"2020-09-20","arxiv_id":"2009.09343","n_code_links":1,"syntology":null},{"paper":"/paper/longformer-for-ms-marco-document-re-ranking","slug":"longformer-for-ms-marco-document-re-ranking","title":"Longformer for MS MARCO Document Re-ranking Task","date":"2020-09-20","arxiv_id":"2009.09392","n_code_links":1,"syntology":null},{"paper":"/paper/persian-ezafe-recognition-using-transformers","slug":"persian-ezafe-recognition-using-transformers","title":"Persian Ezafe Recognition Using Transformers and Its Role in Part-Of-Speech Tagging","date":"2020-09-20","arxiv_id":"2009.09474","n_code_links":1,"syntology":null},{"paper":null,"slug":"vicomtech-at-ehealth-kd-challenge-2020-deep","title":"Vicomtech at eHealth-KD Challenge 2020: Deep End-to-End Model for Entity and Relation Extraction in Medical Text","date":"2020-09-20","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"virtualflow-decoupling-deep-learning-model","title":"VirtualFlow: Decoupling Deep Learning Models from the Underlying Hardware","date":"2020-09-20","arxiv_id":"2009.09523","n_code_links":0,"syntology":null},{"paper":"/paper/conditionally-adaptive-multi-task-learning","slug":"conditionally-adaptive-multi-task-learning","title":"Conditionally Adaptive Multi-Task Learning: Improving Transfer Learning in NLP Using Fewer Parameters & Less Data","date":"2020-09-19","arxiv_id":"2009.09139","n_code_links":1,"syntology":null},{"paper":null,"slug":"nominal-compound-chain-extraction-a-new-task","title":"Nominal Compound Chain Extraction: A New Task for Semantic-enriched Lexical Chain","date":"2020-09-19","arxiv_id":"2009.09173","n_code_links":0,"syntology":null},{"paper":null,"slug":"prior-art-search-and-reranking-for-generated","title":"Prior Art Search and Reranking for Generated Patent Text","date":"2020-09-19","arxiv_id":"2009.09132","n_code_links":0,"syntology":null},{"paper":"/paper/farstail-a-persian-natural-language-inference","slug":"farstail-a-persian-natural-language-inference","title":"FarsTail: A Persian Natural Language Inference Dataset","date":"2020-09-18","arxiv_id":"2009.08820","n_code_links":1,"syntology":null},{"paper":"/paper/fasthan-a-bert-based-joint-many-task-toolkit","slug":"fasthan-a-bert-based-joint-many-task-toolkit","title":"fastHan: A BERT-based Multi-Task Toolkit for Chinese NLP","date":"2020-09-18","arxiv_id":"2009.08633","n_code_links":1,"syntology":null},{"paper":null,"slug":"neu-at-wnut-2020-task-2-data-augmentation-to","title":"NEU at WNUT-2020 Task 2: Data Augmentation To Tell BERT That Death Is Not Necessarily Informative","date":"2020-09-18","arxiv_id":"2009.08590","n_code_links":0,"syntology":null},{"paper":"/paper/the-birth-of-romanian-bert","slug":"the-birth-of-romanian-bert","title":"The birth of Romanian BERT","date":"2020-09-18","arxiv_id":"2009.08712","n_code_links":1,"syntology":null},{"paper":"/paper/will-it-unblend","slug":"will-it-unblend","title":"Will it Unblend?","date":"2020-09-18","arxiv_id":"2009.09123","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-multimodal-memes-classification-a-survey","title":"A Multimodal Memes Classification: A Survey and Open Research Issues","date":"2020-09-17","arxiv_id":"2009.08395","n_code_links":0,"syntology":null},{"paper":null,"slug":"compositional-and-lexical-semantics-in","title":"Compositional and Lexical Semantics in RoBERTa, BERT and DistilBERT: A Case Study on CoQA","date":"2020-09-17","arxiv_id":"2009.08257","n_code_links":0,"syntology":null},{"paper":null,"slug":"cross-modal-alignment-with-mixture-experts","title":"Cross-Modal Alignment with Mixture Experts Neural Network for Intral-City Retail Recommendation","date":"2020-09-17","arxiv_id":"2009.09926","n_code_links":0,"syntology":null},{"paper":"/paper/dsc-iit-ism-at-semeval-2020-task-6-boosting","slug":"dsc-iit-ism-at-semeval-2020-task-6-boosting","title":"DSC IIT-ISM at SemEval-2020 Task 6: Boosting BERT with Dependencies for Definition Extraction","date":"2020-09-17","arxiv_id":"2009.08180","n_code_links":1,"syntology":null},{"paper":null,"slug":"efficient-transformer-based-large-scale","title":"Efficient Transformer-based Large Scale Language Representations using Hardware-friendly Block Structured Pruning","date":"2020-09-17","arxiv_id":"2009.08065","n_code_links":0,"syntology":null}],"record_sha256":"7a1d9a7a096ce59350aa33febb2af286aac57595aa88e848e1a02603fedfbfe9","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}