{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/roberta/papers/6","list_of":"/method/roberta","method":"RoBERTa","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":6,"pages_in_order":10,"rows_per_page":100,"rows":[501,600],"of":913,"counts":{"archive_papers_tagged":913,"with_a_code_link":399,"where_syntology_ran_a_sample":87,"not_listed_spam_title":0,"listed":913,"listed_where_code_ran":87,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":66,"every_run_a_failure_of_syntologys_instrument":21,"listed_with_a_run_with_no_instrument_failure":66,"listed_every_run_a_failure_of_syntologys_instrument":21,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/roberta","prev":"/method/roberta/papers/5","next":"/method/roberta/papers/7","papers":[{"paper":null,"slug":"do-transformers-use-variable-binding","title":"Do Transformers know symbolic rules, and would we know if they did?","date":"2022-02-19","arxiv_id":"2203.00162","n_code_links":0,"syntology":null},{"paper":"/paper/automatic-issue-classifier-a-transfer","slug":"automatic-issue-classifier-a-transfer","title":"Automatic Issue Classifier: A Transfer Learning Framework for Classifying Issue Reports","date":"2022-02-12","arxiv_id":"2202.06149","n_code_links":1,"syntology":null},{"paper":"/paper/pnlp-mixer-an-efficient-all-mlp-architecture","slug":"pnlp-mixer-an-efficient-all-mlp-architecture","title":"pNLP-Mixer: an Efficient all-MLP Architecture for Language","date":"2022-02-09","arxiv_id":"2202.04350","n_code_links":1,"syntology":null},{"paper":null,"slug":"do-language-models-learn-position-role","title":"Do Language Models Learn Position-Role Mappings?","date":"2022-02-08","arxiv_id":"2202.03611","n_code_links":0,"syntology":null},{"paper":null,"slug":"logical-reasoning-for-task-oriented-dialogue","title":"Logical Reasoning for Task Oriented Dialogue Systems","date":"2022-02-08","arxiv_id":"2202.04161","n_code_links":0,"syntology":null},{"paper":"/paper/memory-efficient-backpropagation-through","slug":"memory-efficient-backpropagation-through","title":"Memory-Efficient Backpropagation through Large Linear Layers","date":"2022-01-31","arxiv_id":"2201.13195","n_code_links":2,"syntology":null},{"paper":"/paper/black-box-prompt-learning-for-pre-trained","slug":"black-box-prompt-learning-for-pre-trained","title":"Black-box Prompt Learning for Pre-trained Language Models","date":"2022-01-21","arxiv_id":"2201.08531","n_code_links":1,"syntology":{"ran":3,"of":4,"n_ran_checked":3,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["shizhediao/black-box-prompt-learning"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/dual-contrastive-learning-text-classification","slug":"dual-contrastive-learning-text-classification","title":"Dual Contrastive Learning: Text Classification via Label-Aware Data Augmentation","date":"2022-01-21","arxiv_id":"2201.08702","n_code_links":2,"syntology":null},{"paper":null,"slug":"applying-softtriple-loss-for-supervised-1","title":"Applying SoftTriple Loss for Supervised Language Model Fine Tuning","date":"2022-01-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"bridge-the-gap-between-cv-and-nlp-a-gradient-1","title":"Bridge the Gap Between CV and NLP! A Gradient-based Textual Adversarial Attack Framework","date":"2022-01-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"deck-behavioral-tests-to-improve","title":"DECK: Behavioral Tests to Improve Interpretability and Generalizability of BERT Models Detecting Depression from Text","date":"2022-01-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"impli-investigatng-nli-models-performance-on","title":"IMPLI: Investigatng NLI Models' Performance on Figurative Language","date":"2022-01-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"template-temprel-classification-model-trained","title":"TEMPLATE: TempRel Classification Model Trained with Embedded Temporal Relation Knowledge","date":"2022-01-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"uncovering-surprising-event-boundaries-in","title":"Uncovering Surprising Event Boundaries in Narratives","date":"2022-01-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"understand-before-answer-improve-temporal","title":"Understand before Answer: Improve Temporal Reading Comprehension via Precise Question Understanding","date":"2022-01-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/what-do-tokens-know-about-their-characters","slug":"what-do-tokens-know-about-their-characters","title":"What do tokens know about their characters and how do they know it?","date":"2022-01-16","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/knowledge-graph-augmented-network-towards","slug":"knowledge-graph-augmented-network-towards","title":"Knowledge Graph Augmented Network Towards Multiview Representation Learning for Aspect-based Sentiment Analysis","date":"2022-01-13","arxiv_id":"2201.04831","n_code_links":1,"syntology":null},{"paper":null,"slug":"towards-automated-error-analysis-learning-to","title":"Towards Automated Error Analysis: Learning to Characterize Errors","date":"2022-01-13","arxiv_id":"2201.05017","n_code_links":0,"syntology":null},{"paper":"/paper/promptbert-improving-bert-sentence-embeddings-1","slug":"promptbert-improving-bert-sentence-embeddings-1","title":"PromptBERT: Improving BERT Sentence Embeddings with Prompts","date":"2022-01-12","arxiv_id":"2201.04337","n_code_links":1,"syntology":{"ran":1,"of":5,"n_ran_checked":1,"n_instrument":0,"unverified":4,"pointer_only":5,"phrase":"1 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified; the one sample that ran constructed an object rather than computing a result","official":{"repos":["kongds/prompt-bert"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"explaining-prediction-uncertainty-of-pre","title":"Explaining Predictive Uncertainty by Looking Back at Model Explanations","date":"2022-01-11","arxiv_id":"2201.03742","n_code_links":0,"syntology":null},{"paper":"/paper/black-box-tuning-for-language-model-as-a","slug":"black-box-tuning-for-language-model-as-a","title":"Black-Box Tuning for Language-Model-as-a-Service","date":"2022-01-10","arxiv_id":"2201.03514","n_code_links":2,"syntology":{"ran":3,"of":4,"n_ran_checked":3,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["txsun1997/black-box-tuning"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/dense-to-sparse-gate-for-mixture-of-experts-1","slug":"dense-to-sparse-gate-for-mixture-of-experts-1","title":"EvoMoE: An Evolutional Mixture-of-Experts Training Framework via Dense-To-Sparse Gate","date":"2021-12-29","arxiv_id":"2112.14397","n_code_links":2,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["codecaution/evomoe"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"contextual-sentence-analysis-for-the","title":"Contextual Sentence Analysis for the Sentiment Prediction on Financial Data","date":"2021-12-27","arxiv_id":"2112.13790","n_code_links":0,"syntology":null},{"paper":null,"slug":"secondary-use-of-clinical-problem-list","title":"Secondary Use of Clinical Problem List Entries for Neural Network-Based Disease Code Assignment","date":"2021-12-27","arxiv_id":"2112.13756","n_code_links":0,"syntology":null},{"paper":null,"slug":"understanding-roberta-s-mood-the-role-of","title":"Evaluating Contextual Embeddings and their Extraction Layers for Depression Assessment","date":"2021-12-27","arxiv_id":"2112.13795","n_code_links":0,"syntology":null},{"paper":null,"slug":"challenging-america-modeling-language-in","title":"Challenging America: Modeling language in longer time scales","date":"2021-12-17","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/commonsense-knowledge-augmented-pretrained-1","slug":"commonsense-knowledge-augmented-pretrained-1","title":"Knowledge-Augmented Language Models for Cause-Effect Relation Classification","date":"2021-12-16","arxiv_id":"2112.08615","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":0,"n_instrument":2,"unverified":0,"pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":{"repos":["phosseini/causal-reasoning"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"applying-softtriple-loss-for-supervised","title":"Applying SoftTriple Loss for Supervised Language Model Fine Tuning","date":"2021-12-15","arxiv_id":"2112.08462","n_code_links":0,"syntology":null},{"paper":"/paper/wechsel-effective-initialization-of-subword-1","slug":"wechsel-effective-initialization-of-subword-1","title":"WECHSEL: Effective initialization of subword embeddings for cross-lingual transfer of monolingual language models","date":"2021-12-13","arxiv_id":"2112.06598","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["cpjku/wechsel"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/racebert-a-transformer-based-model-for","slug":"racebert-a-transformer-based-model-for","title":"raceBERT -- A Transformer-based Model for Predicting Race and Ethnicity from Names","date":"2021-12-07","arxiv_id":"2112.03807","n_code_links":1,"syntology":null},{"paper":"/paper/bridging-pre-trained-models-and-downstream","slug":"bridging-pre-trained-models-and-downstream","title":"Bridging Pre-trained Models and Downstream Tasks for Source Code Understanding","date":"2021-12-04","arxiv_id":"2112.02268","n_code_links":1,"syntology":{"ran":12,"of":13,"n_ran_checked":12,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"12 ran (of which 0 constructed an object rather than computing a result; 12 with no instrument failure: 0 honoured, 0 violated, 12 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["wangdeze18/DACL"],"state":"official (archive's flag): 12 ran","n_ran":12,"n_constructed":0,"n_ran_no_instrument_failure":12,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/sentiment-analysis-and-effect-of-covid-19","slug":"sentiment-analysis-and-effect-of-covid-19","title":"Sentiment Analysis and Effect of COVID-19 Pandemic using College SubReddit Data","date":"2021-11-30","arxiv_id":"2112.04351","n_code_links":1,"syntology":null},{"paper":null,"slug":"customer-sentiment-analysis-using-weak","title":"Customer Sentiment Analysis using Weak Supervision for Customer-Agent Chat","date":"2021-11-29","arxiv_id":"2111.14282","n_code_links":0,"syntology":null},{"paper":null,"slug":"anna-enhanced-language-representation-for","title":"ANNA: Enhanced Language Representation for Question Answering","date":"2021-11-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"assessing-the-coherence-modeling-capabilities","title":"Assessing the Coherence Modeling Capabilities of Pretrained Transformer-based Language Models","date":"2021-11-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"bert-got-a-date-introducing-transformers-to-1","title":"BERT got a Date: Introducing Transformers to Temporal Tagging","date":"2021-11-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"impact-of-tokenization-on-language-models-an","title":"Impact of Tokenization on Language Models: An Analysis for Turkish","date":"2021-11-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/improved-grammatical-error-correction-by","slug":"improved-grammatical-error-correction-by","title":"Improved grammatical error correction by ranking elementary edits","date":"2021-11-16","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"interpreting-the-robustness-of-neural-nlp","title":"Interpreting the Robustness of Neural NLP Models to Textual Perturbations","date":"2021-11-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"learning-to-ignore-adversarial-attacks","title":"Learning to Ignore Adversarial Attacks","date":"2021-11-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"nsp-bert-a-prompt-based-zero-shot-learner-1","title":"NSP-BERT: A Prompt-based Zero-Shot Learner Through an Original Pre-training Task —— Next Sentence Prediction","date":"2021-11-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"on-the-robustness-of-reading-comprehension-1","title":"On the Robustness of Reading Comprehension Models to Entity Renaming","date":"2021-11-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"perturbations-in-the-wild-leveraging-human","title":"Perturbations in the Wild: Leveraging Human-Written Text Perturbations for Realistic Adversarial Attack and Defense","date":"2021-11-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"promptbert-improving-bert-sentence-embeddings","title":"PromptBERT: Improving BERT Sentence Embeddings with Prompts","date":"2021-11-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"unsupervised-multiple-choice-question","title":"Unsupervised multiple-choice question generation for out-of-domain Q\\&A fine-tuning","date":"2021-11-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"assessing-gender-bias-in-medical-and","title":"Assessing gender bias in medical and scientific masked language models with StereoSet","date":"2021-11-15","arxiv_id":"2111.08088","n_code_links":0,"syntology":null},{"paper":"/paper/character-level-hypernetworks-for-hate-speech","slug":"character-level-hypernetworks-for-hate-speech","title":"Character-level HyperNetworks for Hate Speech Detection","date":"2021-11-11","arxiv_id":"2111.06336","n_code_links":1,"syntology":null},{"paper":null,"slug":"improving-large-scale-language-models-and","title":"Improving Large-scale Language Models and Resources for Filipino","date":"2021-11-11","arxiv_id":"2111.06053","n_code_links":0,"syntology":null},{"paper":null,"slug":"amazon-sagemaker-model-parallelism-a-general","title":"Amazon SageMaker Model Parallelism: A General and Flexible Framework for Large Model Training","date":"2021-11-10","arxiv_id":"2111.05972","n_code_links":0,"syntology":null},{"paper":"/paper/bagbert-bert-based-bagging-stacking-for-multi","slug":"bagbert-bert-based-bagging-stacking-for-multi","title":"BagBERT: BERT-based bagging-stacking for multi-topic classification","date":"2021-11-10","arxiv_id":"2111.05808","n_code_links":1,"syntology":null},{"paper":"/paper/tacl-improving-bert-pre-training-with-token","slug":"tacl-improving-bert-pre-training-with-token","title":"TaCL: Improving BERT Pre-training with Token-aware Contrastive Learning","date":"2021-11-07","arxiv_id":"2111.04198","n_code_links":2,"syntology":null},{"paper":null,"slug":"ibert-idiom-cloze-style-reading-comprehension","title":"IBERT: Idiom Cloze-style reading comprehension with Attention","date":"2021-11-05","arxiv_id":"2112.02994","n_code_links":0,"syntology":null},{"paper":"/paper/an-empirical-study-of-the-effectiveness-of-an","slug":"an-empirical-study-of-the-effectiveness-of-an","title":"An Empirical Study of the Effectiveness of an Ensemble of Stand-alone Sentiment Detection Tools for Software Engineering Datasets","date":"2021-11-04","arxiv_id":"2111.03196","n_code_links":1,"syntology":null},{"paper":"/paper/an-empirical-study-of-training-end-to-end","slug":"an-empirical-study-of-training-end-to-end","title":"An Empirical Study of Training End-to-End Vision-and-Language Transformers","date":"2021-11-03","arxiv_id":"2111.02387","n_code_links":3,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["zdou0830/meter"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/dsee-dually-sparsity-embedded-efficient-1","slug":"dsee-dually-sparsity-embedded-efficient-1","title":"DSEE: Dually Sparsity-embedded Efficient Tuning of Pre-trained Language Models","date":"2021-10-30","arxiv_id":"2111.00160","n_code_links":1,"syntology":{"ran":10,"of":13,"n_ran_checked":8,"n_instrument":2,"unverified":3,"pointer_only":10,"phrase":"10 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 0 violated, 8 with no contract checked; 2 where Syntology's instrument failed) · 3 unverified","official":{"repos":["vita-group/dsee"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":3,"ran_from_kinds":["official","unlocated"]}}},{"paper":"/paper/bridge-the-gap-between-cv-and-nlp-a-gradient","slug":"bridge-the-gap-between-cv-and-nlp-a-gradient","title":"Bridge the Gap Between CV and NLP! A Gradient-based Textual Adversarial Attack Framework","date":"2021-10-28","arxiv_id":"2110.15317","n_code_links":1,"syntology":null},{"paper":null,"slug":"hate-and-offensive-speech-detection-in-hindi","title":"Hate and Offensive Speech Detection in Hindi and Marathi","date":"2021-10-23","arxiv_id":"2110.12200","n_code_links":0,"syntology":null},{"paper":null,"slug":"impli-investigating-nli-models-performance-on","title":"IMPLI: Investigating NLI Models' Performance on Figurative Language","date":"2021-10-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"models-in-a-spelling-bee-language-models-1","title":"Models In a Spelling Bee: Language Models Implicitly Learn the Character Composition of Tokens","date":"2021-10-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/on-the-robustness-of-reading-comprehension","slug":"on-the-robustness-of-reading-comprehension","title":"On the Robustness of Reading Comprehension Models to Entity Renaming","date":"2021-10-16","arxiv_id":"2110.08555","n_code_links":1,"syntology":null},{"paper":null,"slug":"wechsel-effective-initialization-of-subword","title":"WECHSEL: Effective initialization of subword embeddings for cross-lingual transfer of monolingual language models","date":"2021-10-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/evaluating-the-faithfulness-of-importance","slug":"evaluating-the-faithfulness-of-importance","title":"Evaluating the Faithfulness of Importance Measures in NLP by Recursively Masking Allegedly Important Tokens and Retraining","date":"2021-10-15","arxiv_id":"2110.08412","n_code_links":1,"syntology":{"ran":4,"of":4,"n_ran_checked":4,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["AndreasMadsen/nlp-roar-interpretability"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"causally-estimating-the-sensitivity-of-neural-1","title":"Interpreting the Robustness of Neural NLP Models to Textual Perturbations","date":"2021-10-14","arxiv_id":"2110.07159","n_code_links":0,"syntology":null},{"paper":"/paper/p-adapters-robustly-extracting-factual-1","slug":"p-adapters-robustly-extracting-factual-1","title":"P-Adapters: Robustly Extracting Factual Information from Language Models with Diverse Prompts","date":"2021-10-14","arxiv_id":"2110.07280","n_code_links":1,"syntology":null},{"paper":null,"slug":"extracting-feelings-of-people-regarding-covid","title":"Extracting Feelings of People Regarding COVID-19 by Social Network Mining","date":"2021-10-12","arxiv_id":"2110.06151","n_code_links":0,"syntology":null},{"paper":"/paper/leveraging-recent-advances-in-pre-trained","slug":"leveraging-recent-advances-in-pre-trained","title":"Leveraging recent advances in Pre-Trained Language Models forEye-Tracking Prediction","date":"2021-10-09","arxiv_id":"2110.04475","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-comparative-study-of-transformer-based","title":"A Comparative Study of Transformer-Based Language Models on Extractive Question Answering","date":"2021-10-07","arxiv_id":"2110.03142","n_code_links":0,"syntology":null},{"paper":"/paper/8-bit-optimizers-via-block-wise-quantization","slug":"8-bit-optimizers-via-block-wise-quantization","title":"8-bit Optimizers via Block-wise Quantization","date":"2021-10-06","arxiv_id":"2110.02861","n_code_links":3,"syntology":null},{"paper":null,"slug":"analyzing-the-impact-of-covid-19-on-economy","title":"Analyzing the Impact of COVID-19 on Economy from the Perspective of Users Reviews","date":"2021-10-05","arxiv_id":"2110.02198","n_code_links":0,"syntology":null},{"paper":"/paper/foodchem-a-food-chemical-relation-extraction","slug":"foodchem-a-food-chemical-relation-extraction","title":"FoodChem: A food-chemical relation extraction model","date":"2021-10-05","arxiv_id":"2110.02019","n_code_links":1,"syntology":null},{"paper":null,"slug":"exploiting-pre-trained-asr-models-for","title":"Exploiting Pre-Trained ASR Models for Alzheimer's Disease Recognition Through Spontaneous Speech","date":"2021-10-04","arxiv_id":"2110.01493","n_code_links":0,"syntology":null},{"paper":"/paper/bert-got-a-date-introducing-transformers-to","slug":"bert-got-a-date-introducing-transformers-to","title":"BERT got a Date: Introducing Transformers to Temporal Tagging","date":"2021-09-30","arxiv_id":"2109.14927","n_code_links":1,"syntology":null},{"paper":null,"slug":"scala-speeding-up-fine-tuning-of-pre-trained","title":"ScaLA: Speeding-Up Fine-tuning of Pre-trained Transformer Networks via Efficient and Scalable Adversarial Perturbation","date":"2021-09-29","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/effective-use-of-graph-convolution-network","slug":"effective-use-of-graph-convolution-network","title":"Effective Use of Graph Convolution Network and Contextual Sub-Tree forCommodity News Event Extraction","date":"2021-09-27","arxiv_id":"2109.12781","n_code_links":1,"syntology":null},{"paper":null,"slug":"finetuning-transformer-models-to-build-asag","title":"Finetuning Transformer Models to Build ASAG System","date":"2021-09-25","arxiv_id":"2109.12300","n_code_links":0,"syntology":null},{"paper":null,"slug":"bertweetfr-domain-adaptation-of-pre-trained","title":"BERTweetFR : Domain Adaptation of Pre-Trained Language Models for French Tweets","date":"2021-09-21","arxiv_id":"2109.10234","n_code_links":0,"syntology":null},{"paper":"/paper/mirrorwic-on-eliciting-word-in-context","slug":"mirrorwic-on-eliciting-word-in-context","title":"MirrorWiC: On Eliciting Word-in-Context Representations from Pretrained Language Models","date":"2021-09-19","arxiv_id":"2109.09237","n_code_links":1,"syntology":null},{"paper":null,"slug":"fine-tuned-transformers-show-clusters-of","title":"Fine-Tuned Transformers Show Clusters of Similar Representations Across Layers","date":"2021-09-17","arxiv_id":"2109.08406","n_code_links":0,"syntology":null},{"paper":"/paper/grounding-natural-language-instructions-can","slug":"grounding-natural-language-instructions-can","title":"Grounding Natural Language Instructions: Can Large Language Models Capture Spatial Information?","date":"2021-09-17","arxiv_id":"2109.08634","n_code_links":1,"syntology":null},{"paper":null,"slug":"the-futility-of-stilts-for-the-classification","title":"The futility of STILTs for the classification of lexical borrowings in Spanish","date":"2021-09-17","arxiv_id":"2109.08607","n_code_links":0,"syntology":null},{"paper":null,"slug":"efficient-domain-adaptation-of-language","title":"Efficient Domain Adaptation of Language Models via Adaptive Tokenization","date":"2021-09-15","arxiv_id":"2109.07460","n_code_links":0,"syntology":null},{"paper":null,"slug":"legal-transformer-models-may-not-always-help","title":"Legal Transformer Models May Not Always Help","date":"2021-09-14","arxiv_id":"2109.06862","n_code_links":0,"syntology":null},{"paper":"/paper/question-answering-over-electronic-devices-a","slug":"question-answering-over-electronic-devices-a","title":"Question Answering over Electronic Devices: A New Benchmark Dataset and a Multi-Task Learning based QA Framework","date":"2021-09-13","arxiv_id":"2109.05897","n_code_links":1,"syntology":{"ran":9,"of":9,"n_ran_checked":8,"n_instrument":1,"unverified":0,"pointer_only":0,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 0 violated, 8 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["abhi1nandy2/emnlp-2021-findings"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/teasel-a-transformer-based-speech-prefixed","slug":"teasel-a-transformer-based-speech-prefixed","title":"TEASEL: A Transformer-Based Speech-Prefixed Language Model","date":"2021-09-12","arxiv_id":"2109.05522","n_code_links":1,"syntology":{"ran":2,"of":4,"n_ran_checked":2,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":null}},{"paper":"/paper/d-rex-dialogue-relation-extraction-with","slug":"d-rex-dialogue-relation-extraction-with","title":"D-REX: Dialogue Relation Extraction with Explanations","date":"2021-09-10","arxiv_id":"2109.05126","n_code_links":1,"syntology":null},{"paper":null,"slug":"mining-points-of-interest-via-address","title":"Mining Points of Interest via Address Embeddings: An Unsupervised Approach","date":"2021-09-09","arxiv_id":"2109.04467","n_code_links":0,"syntology":null},{"paper":"/paper/multi-granularity-textual-adversarial-attack","slug":"multi-granularity-textual-adversarial-attack","title":"Multi-granularity Textual Adversarial Attack with Behavior Cloning","date":"2021-09-09","arxiv_id":"2109.04367","n_code_links":1,"syntology":null},{"paper":"/paper/word-level-coreference-resolution","slug":"word-level-coreference-resolution","title":"Word-Level Coreference Resolution","date":"2021-09-09","arxiv_id":"2109.04127","n_code_links":1,"syntology":null},{"paper":"/paper/nsp-bert-a-prompt-based-zero-shot-learner","slug":"nsp-bert-a-prompt-based-zero-shot-learner","title":"NSP-BERT: A Prompt-based Few-Shot Learner Through an Original Pre-training Task--Next Sentence Prediction","date":"2021-09-08","arxiv_id":"2109.03564","n_code_links":1,"syntology":null},{"paper":null,"slug":"empathetic-dialogue-generation-with-pre","title":"Empathetic Dialogue Generation with Pre-trained RoBERTa-GPT2 and External Knowledge","date":"2021-09-07","arxiv_id":"2109.03004","n_code_links":0,"syntology":null},{"paper":null,"slug":"how-much-pretraining-data-do-language-models","title":"How much pretraining data do language models need to learn syntax?","date":"2021-09-07","arxiv_id":"2109.03160","n_code_links":0,"syntology":null},{"paper":null,"slug":"error-detection-in-large-scale-natural","title":"Error Detection in Large-Scale Natural Language Understanding Systems Using Transformer Models","date":"2021-09-04","arxiv_id":"2109.01754","n_code_links":0,"syntology":null},{"paper":null,"slug":"conqx-semantic-expansion-of-spoken-queries","title":"ConQX: Semantic Expansion of Spoken Queries for Intent Detection based on Conditioned Text Generation","date":"2021-09-02","arxiv_id":"2109.00729","n_code_links":0,"syntology":null},{"paper":null,"slug":"so-cloze-yet-so-far-n400-amplitude-is-better","title":"So Cloze yet so Far: N400 Amplitude is Better Predicted by Distributional Information than Human Predictability Judgements","date":"2021-09-02","arxiv_id":"2109.01226","n_code_links":0,"syntology":null},{"paper":null,"slug":"fight-fire-with-fire-fine-tuning-hate","title":"Fight Fire with Fire: Fine-tuning Hate Detectors using Large Samples of Generated Hate Speech","date":"2021-09-01","arxiv_id":"2109.00591","n_code_links":0,"syntology":null},{"paper":"/paper/towards-improving-adversarial-training-of-nlp","slug":"towards-improving-adversarial-training-of-nlp","title":"Towards Improving Adversarial Training of NLP Models","date":"2021-09-01","arxiv_id":"2109.00544","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["QData/TextAttack-A2T"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/evaluating-the-robustness-of-neural-language","slug":"evaluating-the-robustness-of-neural-language","title":"Evaluating the Robustness of Neural Language Models to Input Perturbations","date":"2021-08-27","arxiv_id":"2108.12237","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["mmoradi-iut/nlp-perturbation"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/a-statutory-article-retrieval-dataset-in","slug":"a-statutory-article-retrieval-dataset-in","title":"A Statutory Article Retrieval Dataset in French","date":"2021-08-26","arxiv_id":"2108.11792","n_code_links":1,"syntology":null},{"paper":"/paper/emoberta-speaker-aware-emotion-recognition-in","slug":"emoberta-speaker-aware-emotion-recognition-in","title":"EmoBERTa: Speaker-Aware Emotion Recognition in Conversation with RoBERTa","date":"2021-08-26","arxiv_id":"2108.12009","n_code_links":1,"syntology":null},{"paper":"/paper/rethinking-why-intermediate-task-fine-tuning","slug":"rethinking-why-intermediate-task-fine-tuning","title":"Rethinking Why Intermediate-Task Fine-Tuning Works","date":"2021-08-26","arxiv_id":"2108.11696","n_code_links":1,"syntology":null}],"record_sha256":"51f164b3b9f03e807df48775d1fea1fe2518fcc4a84ebb3f975605c0a2b48948","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}