{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/bert/papers/50","list_of":"/method/bert","method":"BERT","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":50,"pages_in_order":70,"rows_per_page":100,"rows":[4901,5000],"of":6938,"counts":{"archive_papers_tagged":6938,"with_a_code_link":2862,"where_syntology_ran_a_sample":640,"not_listed_spam_title":0,"listed":6938,"listed_where_code_ran":640,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":520,"every_run_a_failure_of_syntologys_instrument":120,"listed_with_a_run_with_no_instrument_failure":520,"listed_every_run_a_failure_of_syntologys_instrument":120,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/bert","prev":"/method/bert/papers/49","next":"/method/bert/papers/51","papers":[{"paper":"/paper/keep-learning-self-supervised-meta-learning","slug":"keep-learning-self-supervised-meta-learning","title":"Keep Learning: Self-supervised Meta-learning for Learning from Inference","date":"2021-04-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"maximal-multiverse-learning-for-promoting","title":"Maximal Multiverse Learning for Promoting Cross-Task Generalization of Fine-Tuned Language Models","date":"2021-04-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/multilingual-entity-and-relation-extraction","slug":"multilingual-entity-and-relation-extraction","title":"Multilingual Entity and Relation Extraction Dataset and Model","date":"2021-04-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"neural-driven-search-based-paraphrase","title":"Neural-Driven Search-Based Paraphrase Generation","date":"2021-04-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/nlquad-a-non-factoid-long-question-answering","slug":"nlquad-a-non-factoid-long-question-answering","title":"NLQuAD: A Non-Factoid Long Question Answering Data Set","date":"2021-04-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"on-the-in-effectiveness-of-images-for-text","title":"On the (In)Effectiveness of Images for Text Classification","date":"2021-04-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/probing-for-idiomaticity-in-vector-space","slug":"probing-for-idiomaticity-in-vector-space","title":"Probing for idiomaticity in vector space models","date":"2021-04-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"retrieval-re-ranking-and-multi-task-learning","title":"Retrieval, Re-ranking and Multi-task Learning for Knowledge-Base Question Answering","date":"2021-04-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/an-in-depth-analysis-of-passage-level-label","slug":"an-in-depth-analysis-of-passage-level-label","title":"An In-depth Analysis of Passage-Level Label Transfer for Contextual Document Ranking","date":"2021-03-30","arxiv_id":"2103.16669","n_code_links":1,"syntology":null},{"paper":null,"slug":"automatic-graph-partitioning-for-very-large","title":"Automatic Graph Partitioning for Very Large-scale Deep Learning","date":"2021-03-30","arxiv_id":"2103.16063","n_code_links":0,"syntology":null},{"paper":"/paper/grounding-dialogue-systems-via-knowledge","slug":"grounding-dialogue-systems-via-knowledge","title":"Grounding Dialogue Systems via Knowledge Graph Aware Decoding with Pre-trained Transformers","date":"2021-03-30","arxiv_id":"2103.16289","n_code_links":1,"syntology":null},{"paper":"/paper/kaleido-bert-vision-language-pre-training-on","slug":"kaleido-bert-vision-language-pre-training-on","title":"Kaleido-BERT: Vision-Language Pre-training on Fashion Domain","date":"2021-03-30","arxiv_id":"2103.16110","n_code_links":1,"syntology":{"ran":3,"of":5,"n_ran_checked":3,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 2 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["mczhuge/Kaleido-BERT"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"contextual-text-embeddings-for-twi","title":"Contextual Text Embeddings for Twi","date":"2021-03-29","arxiv_id":"2103.15963","n_code_links":0,"syntology":null},{"paper":null,"slug":"retraining-distilbert-for-a-voice-shopping","title":"Retraining DistilBERT for a Voice Shopping Assistant by Using Universal Dependencies","date":"2021-03-29","arxiv_id":"2103.15737","n_code_links":0,"syntology":null},{"paper":"/paper/whitening-sentence-representations-for-better","slug":"whitening-sentence-representations-for-better","title":"Whitening Sentence Representations for Better Semantics and Faster Retrieval","date":"2021-03-29","arxiv_id":"2103.15316","n_code_links":3,"syntology":{"ran":2,"of":5,"n_ran_checked":2,"n_instrument":0,"unverified":3,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 1 honoured, 1 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":{"repos":["bojone/BERT-whitening"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":null,"slug":"png-bert-augmented-bert-on-phonemes-and","title":"PnG BERT: Augmented BERT on Phonemes and Graphemes for Neural TTS","date":"2021-03-28","arxiv_id":"2103.15060","n_code_links":0,"syntology":null},{"paper":null,"slug":"machine-learning-meets-natural-language","title":"Machine Learning Meets Natural Language Processing -- The story so far","date":"2021-03-27","arxiv_id":"2104.10213","n_code_links":0,"syntology":null},{"paper":null,"slug":"unsupervised-self-training-for-sentiment","title":"Unsupervised Self-Training for Sentiment Analysis of Code-Switched Data","date":"2021-03-27","arxiv_id":"2103.14797","n_code_links":0,"syntology":null},{"paper":null,"slug":"bert4so-neural-sentence-ordering-by-fine","title":"BERT4SO: Neural Sentence Ordering by Fine-tuning BERT","date":"2021-03-25","arxiv_id":"2103.13584","n_code_links":0,"syntology":null},{"paper":null,"slug":"bertinho-galician-bert-representations","title":"Bertinho: Galician BERT Representations","date":"2021-03-25","arxiv_id":"2103.13799","n_code_links":0,"syntology":null},{"paper":null,"slug":"k-xlnet-a-general-method-for-combining","title":"K-XLNet: A General Method for Combining Explicit Knowledge with Language Model Pretraining","date":"2021-03-25","arxiv_id":"2104.10649","n_code_links":0,"syntology":null},{"paper":"/paper/predicting-directionality-in-causal-relations","slug":"predicting-directionality-in-causal-relations","title":"Predicting Directionality in Causal Relations in Text","date":"2021-03-25","arxiv_id":"2103.13606","n_code_links":2,"syntology":null},{"paper":null,"slug":"visual-grounding-strategies-for-text-only","title":"Visual Grounding Strategies for Text-Only Natural Language Processing","date":"2021-03-25","arxiv_id":"2103.13942","n_code_links":0,"syntology":null},{"paper":"/paper/czert-czech-bert-like-model-for-language","slug":"czert-czech-bert-like-model-for-language","title":"Czert -- Czech BERT-like Model for Language Representation","date":"2021-03-24","arxiv_id":"2103.13031","n_code_links":1,"syntology":null},{"paper":"/paper/are-neural-language-models-good-plagiarists-a","slug":"are-neural-language-models-good-plagiarists-a","title":"Are Neural Language Models Good Plagiarists? A Benchmark for Neural Paraphrase Detection","date":"2021-03-23","arxiv_id":"2103.12450","n_code_links":0,"syntology":null},{"paper":null,"slug":"repairing-pronouns-in-translation-with-bert","title":"Repairing Pronouns in Translation with BERT-Based Post-Editing","date":"2021-03-23","arxiv_id":"2103.12838","n_code_links":0,"syntology":null},{"paper":null,"slug":"the-nlp-cookbook-modern-recipes-for","title":"The NLP Cookbook: Modern Recipes for Transformer based Deep Learning Architectures","date":"2021-03-23","arxiv_id":"2104.10640","n_code_links":0,"syntology":null},{"paper":null,"slug":"tmr-evaluating-ner-recall-on-tough-mentions","title":"TMR: Evaluating NER Recall on Tough Mentions","date":"2021-03-23","arxiv_id":"2103.12312","n_code_links":0,"syntology":null},{"paper":null,"slug":"variable-name-recovery-in-decompiled-binary","title":"Variable Name Recovery in Decompiled Binary Code using Constrained Masked Language Modeling","date":"2021-03-23","arxiv_id":"2103.12801","n_code_links":0,"syntology":null},{"paper":null,"slug":"bert-a-review-of-applications-in-natural","title":"BERT: A Review of Applications in Natural Language Processing and Understanding","date":"2021-03-22","arxiv_id":"2103.11943","n_code_links":0,"syntology":null},{"paper":null,"slug":"bridging-the-gap-between-supervised","title":"Bridging the gap between supervised classification and unsupervised topic modelling for social-media assisted crisis management","date":"2021-03-22","arxiv_id":"2103.11835","n_code_links":0,"syntology":null},{"paper":"/paper/hybrid-model-for-patent-classification-using","slug":"hybrid-model-for-patent-classification-using","title":"PatentSBERTa: A Deep NLP based Hybrid Model for Patent Distance and Classification using Augmented SBERT","date":"2021-03-22","arxiv_id":"2103.11933","n_code_links":2,"syntology":null},{"paper":"/paper/open-domain-question-answering-over-tables","slug":"open-domain-question-answering-over-tables","title":"Open Domain Question Answering over Tables via Dense Retrieval","date":"2021-03-22","arxiv_id":"2103.12011","n_code_links":1,"syntology":null},{"paper":null,"slug":"namerec-highly-accurate-and-fine-grained","title":"NameRec*: Highly Accurate and Fine-grained Person Name Recognition","date":"2021-03-21","arxiv_id":"2103.11360","n_code_links":0,"syntology":null},{"paper":"/paper/rosita-refined-bert-compression-with","slug":"rosita-refined-bert-compression-with","title":"ROSITA: Refined BERT cOmpreSsion with InTegrAted techniques","date":"2021-03-21","arxiv_id":"2103.11367","n_code_links":1,"syntology":{"ran":0,"of":1,"n_ran_checked":0,"n_instrument":0,"unverified":1,"pointer_only":1,"phrase":"0 ran · 1 unverified","official":{"repos":["llyx97/Rosita"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":[]}}},{"paper":null,"slug":"cost-effective-deployment-of-bert-models-in","title":"Cost-effective Deployment of BERT Models in Serverless Environment","date":"2021-03-19","arxiv_id":"2103.10673","n_code_links":0,"syntology":null},{"paper":"/paper/muril-multilingual-representations-for-indian","slug":"muril-multilingual-representations-for-indian","title":"MuRIL: Multilingual Representations for Indian Languages","date":"2021-03-19","arxiv_id":"2103.10730","n_code_links":1,"syntology":null},{"paper":"/paper/all-nlp-tasks-are-generation-tasks-a-general","slug":"all-nlp-tasks-are-generation-tasks-a-general","title":"GLM: General Language Model Pretraining with Autoregressive Blank Infilling","date":"2021-03-18","arxiv_id":"2103.10360","n_code_links":8,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["THUDM/GLM"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":null,"slug":"contextual-biasing-of-language-models-for","title":"Contextual Biasing of Language Models for Speech Recognition in Goal-Oriented Conversational Agents","date":"2021-03-18","arxiv_id":"2103.10325","n_code_links":0,"syntology":null},{"paper":"/paper/model-extraction-and-adversarial","slug":"model-extraction-and-adversarial","title":"Model Extraction and Adversarial Transferability, Your BERT is Vulnerable!","date":"2021-03-18","arxiv_id":"2103.10013","n_code_links":1,"syntology":null},{"paper":null,"slug":"code-word-detection-in-fraud-investigations","title":"Code Word Detection in Fraud Investigations using a Deep-Learning Approach","date":"2021-03-17","arxiv_id":"2103.09606","n_code_links":0,"syntology":null},{"paper":"/paper/on-the-role-of-images-for-analyzing-claims-in","slug":"on-the-role-of-images-for-analyzing-claims-in","title":"On the Role of Images for Analyzing Claims in Social Media","date":"2021-03-17","arxiv_id":"2103.09602","n_code_links":1,"syntology":null},{"paper":"/paper/uniparma-semeval-2021-task-5-toxic-spans","slug":"uniparma-semeval-2021-task-5-toxic-spans","title":"UniParma at SemEval-2021 Task 5: Toxic Spans Detection Using CharacterBERT and Bag-of-Words Model","date":"2021-03-17","arxiv_id":"2103.09645","n_code_links":1,"syntology":null},{"paper":null,"slug":"kgsynnet-a-novel-entity-synonyms-discovery","title":"KGSynNet: A Novel Entity Synonyms Discovery Framework with Knowledge Graph","date":"2021-03-16","arxiv_id":"2103.08893","n_code_links":0,"syntology":null},{"paper":null,"slug":"text-mining-of-stocktwits-data-for-predicting","title":"Text Mining of Stocktwits Data for Predicting Stock Prices","date":"2021-03-13","arxiv_id":"2103.16388","n_code_links":0,"syntology":null},{"paper":null,"slug":"comparing-the-performance-of-nlp-toolkits-and","title":"Comparing the Performance of NLP Toolkits and Evaluation measures in Legal Tech","date":"2021-03-12","arxiv_id":"2103.11792","n_code_links":0,"syntology":null},{"paper":null,"slug":"explaining-and-improving-bert-performance-on","title":"Explaining and Improving BERT Performance on Lexical Semantic Change Detection","date":"2021-03-12","arxiv_id":"2103.07259","n_code_links":0,"syntology":null},{"paper":null,"slug":"is-bert-a-cross-disciplinary-knowledge","title":"Is BERT a Cross-Disciplinary Knowledge Learner? A Surprising Finding of Pre-trained Models' Transferability","date":"2021-03-12","arxiv_id":"2103.07162","n_code_links":0,"syntology":null},{"paper":null,"slug":"composite-re-ranking-for-efficient-document","title":"Composite Re-Ranking for Efficient Document Search with BERT","date":"2021-03-11","arxiv_id":"2103.06499","n_code_links":0,"syntology":null},{"paper":null,"slug":"evaluation-of-morphological-embeddings-for-1","title":"Evaluation of Morphological Embeddings for the Russian Language","date":"2021-03-11","arxiv_id":"2103.06628","n_code_links":0,"syntology":null},{"paper":null,"slug":"fairfil-contrastive-neural-debiasing-method-1","title":"FairFil: Contrastive Neural Debiasing Method for Pretrained Text Encoders","date":"2021-03-11","arxiv_id":"2103.06413","n_code_links":0,"syntology":null},{"paper":"/paper/improving-bi-encoder-document-ranking-models","slug":"improving-bi-encoder-document-ranking-models","title":"Improving Bi-encoder Document Ranking Models with Two Rankers and Multi-teacher Distillation","date":"2021-03-11","arxiv_id":"2103.06523","n_code_links":1,"syntology":null},{"paper":null,"slug":"lightmbert-a-simple-yet-effective-method-for","title":"LightMBERT: A Simple Yet Effective Method for Multilingual BERT Distillation","date":"2021-03-11","arxiv_id":"2103.06418","n_code_links":0,"syntology":null},{"paper":null,"slug":"self-supervised-text-to-sql-learning-with","title":"Self-supervised Text-to-SQL Learning with Header Alignment Training","date":"2021-03-11","arxiv_id":"2103.06402","n_code_links":0,"syntology":null},{"paper":"/paper/towards-multi-sense-cross-lingual-alignment-1","slug":"towards-multi-sense-cross-lingual-alignment-1","title":"Towards Multi-Sense Cross-Lingual Alignment of Contextual Embeddings","date":"2021-03-11","arxiv_id":"2103.06459","n_code_links":1,"syntology":null},{"paper":null,"slug":"ceqe-contextualized-embeddings-for-query","title":"CEQE: Contextualized Embeddings for Query Expansion","date":"2021-03-09","arxiv_id":"2103.05256","n_code_links":0,"syntology":null},{"paper":"/paper/language-models-have-a-moral-dimension","slug":"language-models-have-a-moral-dimension","title":"Large Pre-trained Language Models Contain Human-like Biases of What is Right and Wrong to Do","date":"2021-03-08","arxiv_id":"2103.11790","n_code_links":1,"syntology":{"ran":6,"of":6,"n_ran_checked":6,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["ml-research/MoRT_NMI"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"skillbert-skilling-the-bert-to-classify-1","title":"SKILLBERT: “SKILLING” THE BERT TO CLASSIFY SKILLS!","date":"2021-03-08","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/syntax-bert-improving-pre-trained","slug":"syntax-bert-improving-pre-trained","title":"Syntax-BERT: Improving Pre-trained Transformers with Syntax Trees","date":"2021-03-07","arxiv_id":"2103.04350","n_code_links":1,"syntology":null},{"paper":null,"slug":"fine-tuning-pretrained-multilingual-bert","title":"Fine-tuning Pretrained Multilingual BERT Model for Indonesian Aspect-based Sentiment Analysis","date":"2021-03-05","arxiv_id":"2103.03732","n_code_links":0,"syntology":null},{"paper":null,"slug":"malbert-using-transformers-for-cybersecurity","title":"MalBERT: Using Transformers for Cybersecurity and Malicious Software Detection","date":"2021-03-05","arxiv_id":"2103.03806","n_code_links":0,"syntology":null},{"paper":null,"slug":"non-invasive-self-attention-for-side","title":"Non-invasive Self-attention for Side Information Fusion in Sequential Recommendation","date":"2021-03-05","arxiv_id":"2103.03578","n_code_links":0,"syntology":null},{"paper":null,"slug":"hardware-acceleration-of-fully-quantized-bert","title":"Hardware Acceleration of Fully Quantized BERT for Efficient Natural Language Processing","date":"2021-03-04","arxiv_id":"2103.02800","n_code_links":0,"syntology":null},{"paper":null,"slug":"few-shot-learning-for-slot-tagging-with","title":"Few-shot Learning for Slot Tagging with Attentive Relational Network","date":"2021-03-03","arxiv_id":"2103.02333","n_code_links":0,"syntology":null},{"paper":null,"slug":"natural-language-understanding-for","title":"Natural Language Understanding for Argumentative Dialogue Systems in the Opinion Building Domain","date":"2021-03-03","arxiv_id":"2103.02691","n_code_links":0,"syntology":null},{"paper":null,"slug":"hate-towards-the-political-opponent-a-twitter","title":"Hate Towards the Political Opponent: A Twitter Corpus Study of the 2020 US Elections on the Basis of Offensive Speech and Stance Detection","date":"2021-03-02","arxiv_id":"2103.01664","n_code_links":0,"syntology":null},{"paper":null,"slug":"bert-based-knowledge-extraction-method-of","title":"BERT-based knowledge extraction method of unstructured domain text","date":"2021-03-01","arxiv_id":"2103.00728","n_code_links":0,"syntology":null},{"paper":null,"slug":"bert-based-patent-novelty-search-by-training","title":"BERT based patent novelty search by training claims to their own description","date":"2021-03-01","arxiv_id":"2103.01126","n_code_links":0,"syntology":null},{"paper":null,"slug":"combat-covid-19-infodemic-using-explainable","title":"Combat COVID-19 Infodemic Using Explainable Natural Language Processing Models","date":"2021-03-01","arxiv_id":"2103.00747","n_code_links":0,"syntology":null},{"paper":"/paper/nlp-cuet-dravidianlangtech-eacl2021-offensive","slug":"nlp-cuet-dravidianlangtech-eacl2021-offensive","title":"NLP-CUET@DravidianLangTech-EACL2021: Offensive Language Detection from Multilingual Code-Mixed Text using Transformers","date":"2021-02-28","arxiv_id":"2103.00455","n_code_links":1,"syntology":null},{"paper":"/paper/nlp-cuet-lt-edi-eacl2021-multilingual-code","slug":"nlp-cuet-lt-edi-eacl2021-multilingual-code","title":"NLP-CUET@LT-EDI-EACL2021: Multilingual Code-Mixed Hope Speech Detection using Cross-lingual Representation Learner","date":"2021-02-28","arxiv_id":"2103.00464","n_code_links":1,"syntology":null},{"paper":"/paper/covid-19-tweets-analysis-through-transformer","slug":"covid-19-tweets-analysis-through-transformer","title":"COVID-19 Tweets Analysis through Transformer Language Models","date":"2021-02-27","arxiv_id":"2103.00199","n_code_links":1,"syntology":null},{"paper":null,"slug":"transformers-with-competitive-ensembles-of-1","title":"Transformers with Competitive Ensembles of Independent Mechanisms","date":"2021-02-27","arxiv_id":"2103.00336","n_code_links":0,"syntology":null},{"paper":null,"slug":"multi-task-transfer-learning-for-finding","title":"Multi-task transfer learning for finding actionable information from crisis-related messages on social media","date":"2021-02-26","arxiv_id":"2102.13395","n_code_links":0,"syntology":null},{"paper":null,"slug":"bert-based-acronym-disambiguation-with","title":"BERT-based Acronym Disambiguation with Multiple Training Strategies","date":"2021-02-25","arxiv_id":"2103.00488","n_code_links":0,"syntology":null},{"paper":null,"slug":"emotion-aware-emotion-agnostic-or-automatic","title":"Emotion-Aware, Emotion-Agnostic, or Automatic: Corpus Creation Strategies to Obtain Cognitive Event Appraisal Annotations","date":"2021-02-25","arxiv_id":"2102.12858","n_code_links":0,"syntology":null},{"paper":null,"slug":"pharmke-knowledge-extraction-platform-for","title":"PharmKE: Knowledge Extraction Platform for Pharmaceutical Texts using Transfer Learning","date":"2021-02-25","arxiv_id":"2102.13139","n_code_links":0,"syntology":null},{"paper":"/paper/sentiment-analysis-of-persian-english-code","slug":"sentiment-analysis-of-persian-english-code","title":"Sentiment Analysis of Persian-English Code-mixed Texts","date":"2021-02-25","arxiv_id":"2102.12700","n_code_links":1,"syntology":null},{"paper":null,"slug":"from-universal-language-model-to-downstream","title":"From Universal Language Model to Downstream Task: Improving RoBERTa-Based Vietnamese Hate Speech Detection","date":"2021-02-24","arxiv_id":"2102.12162","n_code_links":0,"syntology":null},{"paper":null,"slug":"hopeful-men-lt-edi-eacl2021-hope-speech","title":"Hopeful_Men@LT-EDI-EACL2021: Hope Speech Detection Using Indic Transliteration and Transformers","date":"2021-02-24","arxiv_id":"2102.12082","n_code_links":0,"syntology":null},{"paper":"/paper/lrg-at-semeval-2021-task-4-improving-reading","slug":"lrg-at-semeval-2021-task-4-improving-reading","title":"LRG at SemEval-2021 Task 4: Improving Reading Comprehension with Abstract Words using Augmentation, Linguistic Features and Voting","date":"2021-02-24","arxiv_id":"2102.12255","n_code_links":1,"syntology":null},{"paper":"/paper/nlrg-at-semeval-2021-task-5-toxic-spans","slug":"nlrg-at-semeval-2021-task-5-toxic-spans","title":"NLRG at SemEval-2021 Task 5: Toxic Spans Detection Leveraging BERT-based Token Classification and Span Prediction Techniques","date":"2021-02-24","arxiv_id":"2102.12254","n_code_links":1,"syntology":null},{"paper":null,"slug":"task-specific-pre-training-and-cross-lingual","title":"Task-Specific Pre-Training and Cross Lingual Transfer for Code-Switched Data","date":"2021-02-24","arxiv_id":"2102.12407","n_code_links":0,"syntology":null},{"paper":null,"slug":"minimally-supervised-structure-rich-text","title":"Minimally-Supervised Structure-Rich Text Categorization via Learning on Text-Rich Networks","date":"2021-02-23","arxiv_id":"2102.11479","n_code_links":0,"syntology":null},{"paper":null,"slug":"robust-and-transferable-anomaly-detection-in","title":"Robust and Transferable Anomaly Detection in Log Data using Pre-Trained Language Models","date":"2021-02-23","arxiv_id":"2102.11570","n_code_links":0,"syntology":null},{"paper":"/paper/visualchexbert-addressing-the-discrepancy","slug":"visualchexbert-addressing-the-discrepancy","title":"VisualCheXbert: Addressing the Discrepancy Between Radiology Report Labels and Image Labels","date":"2021-02-23","arxiv_id":"2102.11467","n_code_links":1,"syntology":{"ran":0,"of":2,"n_ran_checked":0,"n_instrument":0,"unverified":2,"pointer_only":2,"phrase":"0 ran · 2 unverified","official":{"repos":["stanfordmlgroup/VisualCheXbert"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":2,"ran_from_kinds":[]}}},{"paper":"/paper/evaluating-contextualized-language-models-for","slug":"evaluating-contextualized-language-models-for","title":"Evaluating Contextualized Language Models for Hungarian","date":"2021-02-22","arxiv_id":"2102.10848","n_code_links":1,"syntology":null},{"paper":null,"slug":"generating-human-readable-transcript-for","title":"Generating Human Readable Transcript for Automatic Speech Recognition with Pre-trained Language Model","date":"2021-02-22","arxiv_id":"2102.11114","n_code_links":0,"syntology":null},{"paper":null,"slug":"mixup-training-leads-to-reduced-overfitting","title":"MixUp Training Leads to Reduced Overfitting and Improved Calibration for the Transformer Architecture","date":"2021-02-22","arxiv_id":"2102.11402","n_code_links":0,"syntology":null},{"paper":"/paper/parallelizing-legendre-memory-unit-training","slug":"parallelizing-legendre-memory-unit-training","title":"Parallelizing Legendre Memory Unit Training","date":"2021-02-22","arxiv_id":"2102.11417","n_code_links":2,"syntology":null},{"paper":null,"slug":"rubert-a-bilingual-roman-urdu-bert-using","title":"RUBERT: A Bilingual Roman Urdu BERT Using Cross Lingual Transfer Learning","date":"2021-02-22","arxiv_id":"2102.11278","n_code_links":0,"syntology":null},{"paper":"/paper/using-prior-knowledge-to-guide-bert-s","slug":"using-prior-knowledge-to-guide-bert-s","title":"Using Prior Knowledge to Guide BERT's Attention in Semantic Textual Matching Tasks","date":"2021-02-22","arxiv_id":"2102.10934","n_code_links":1,"syntology":null},{"paper":null,"slug":"pre-training-bert-on-arabic-tweets-practical","title":"Pre-Training BERT on Arabic Tweets: Practical Considerations","date":"2021-02-21","arxiv_id":"2102.10684","n_code_links":0,"syntology":null},{"paper":null,"slug":"web-based-application-for-detecting","title":"Web-based Application for Detecting Indonesian Clickbait Headlines using IndoBERT","date":"2021-02-21","arxiv_id":"2102.10601","n_code_links":0,"syntology":null},{"paper":null,"slug":"learning-dynamic-bert-via-trainable-gate","title":"Learning Dynamic BERT via Trainable Gate Variables and a Bi-modal Regularizer","date":"2021-02-19","arxiv_id":"2102.09727","n_code_links":0,"syntology":null},{"paper":"/paper/towards-emotion-recognition-in-hindi-english","slug":"towards-emotion-recognition-in-hindi-english","title":"Towards Emotion Recognition in Hindi-English Code-Mixed Data: A Transformer Based Approach","date":"2021-02-19","arxiv_id":"2102.09943","n_code_links":1,"syntology":null},{"paper":"/paper/using-transformer-based-ensemble-learning-to","slug":"using-transformer-based-ensemble-learning-to","title":"Using Transformer based Ensemble Learning to classify Scientific Articles","date":"2021-02-19","arxiv_id":"2102.09991","n_code_links":2,"syntology":null},{"paper":"/paper/analysis-of-contextual-and-non-contextual","slug":"analysis-of-contextual-and-non-contextual","title":"Analysis Of Contextual and Non-Contextual Word Embedding Models For Hindi NER With Web Application For Data Collection","date":"2021-02-18","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/training-microsoft-news-recommenders-with","slug":"training-microsoft-news-recommenders-with","title":"Training Large-Scale News Recommenders with Pretrained Language Models in the Loop","date":"2021-02-18","arxiv_id":"2102.09268","n_code_links":1,"syntology":null},{"paper":null,"slug":"unibuckernel-geolocating-swiss-german-jodels","title":"UnibucKernel: Geolocating Swiss German Jodels Using Ensemble Learning","date":"2021-02-18","arxiv_id":"2102.09379","n_code_links":0,"syntology":null}],"record_sha256":"3bd058b9306b75e56b9b92c28e23d3453ea4fee2289be65a975cf98a49606c49","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}