{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/wordpiece/papers/51","list_of":"/method/wordpiece","method":"WordPiece","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":51,"pages_in_order":71,"rows_per_page":100,"rows":[5001,5100],"of":7063,"counts":{"archive_papers_tagged":7063,"with_a_code_link":2910,"where_syntology_ran_a_sample":650,"not_listed_spam_title":0,"listed":7063,"listed_where_code_ran":650,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":529,"every_run_a_failure_of_syntologys_instrument":121,"listed_with_a_run_with_no_instrument_failure":529,"listed_every_run_a_failure_of_syntologys_instrument":121,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/wordpiece","prev":"/method/wordpiece/papers/50","next":"/method/wordpiece/papers/52","papers":[{"paper":null,"slug":"png-bert-augmented-bert-on-phonemes-and","title":"PnG BERT: Augmented BERT on Phonemes and Graphemes for Neural TTS","date":"2021-03-28","arxiv_id":"2103.15060","n_code_links":0,"syntology":null},{"paper":null,"slug":"machine-learning-meets-natural-language","title":"Machine Learning Meets Natural Language Processing -- The story so far","date":"2021-03-27","arxiv_id":"2104.10213","n_code_links":0,"syntology":null},{"paper":null,"slug":"unsupervised-self-training-for-sentiment","title":"Unsupervised Self-Training for Sentiment Analysis of Code-Switched Data","date":"2021-03-27","arxiv_id":"2103.14797","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-practical-survey-on-faster-and-lighter","title":"A Practical Survey on Faster and Lighter Transformers","date":"2021-03-26","arxiv_id":"2103.14636","n_code_links":0,"syntology":null},{"paper":null,"slug":"bert4so-neural-sentence-ordering-by-fine","title":"BERT4SO: Neural Sentence Ordering by Fine-tuning BERT","date":"2021-03-25","arxiv_id":"2103.13584","n_code_links":0,"syntology":null},{"paper":null,"slug":"bertinho-galician-bert-representations","title":"Bertinho: Galician BERT Representations","date":"2021-03-25","arxiv_id":"2103.13799","n_code_links":0,"syntology":null},{"paper":null,"slug":"k-xlnet-a-general-method-for-combining","title":"K-XLNet: A General Method for Combining Explicit Knowledge with Language Model Pretraining","date":"2021-03-25","arxiv_id":"2104.10649","n_code_links":0,"syntology":null},{"paper":"/paper/predicting-directionality-in-causal-relations","slug":"predicting-directionality-in-causal-relations","title":"Predicting Directionality in Causal Relations in Text","date":"2021-03-25","arxiv_id":"2103.13606","n_code_links":2,"syntology":null},{"paper":null,"slug":"visual-grounding-strategies-for-text-only","title":"Visual Grounding Strategies for Text-Only Natural Language Processing","date":"2021-03-25","arxiv_id":"2103.13942","n_code_links":0,"syntology":null},{"paper":"/paper/czert-czech-bert-like-model-for-language","slug":"czert-czech-bert-like-model-for-language","title":"Czert -- Czech BERT-like Model for Language Representation","date":"2021-03-24","arxiv_id":"2103.13031","n_code_links":1,"syntology":null},{"paper":"/paper/are-neural-language-models-good-plagiarists-a","slug":"are-neural-language-models-good-plagiarists-a","title":"Are Neural Language Models Good Plagiarists? A Benchmark for Neural Paraphrase Detection","date":"2021-03-23","arxiv_id":"2103.12450","n_code_links":0,"syntology":null},{"paper":null,"slug":"repairing-pronouns-in-translation-with-bert","title":"Repairing Pronouns in Translation with BERT-Based Post-Editing","date":"2021-03-23","arxiv_id":"2103.12838","n_code_links":0,"syntology":null},{"paper":null,"slug":"the-nlp-cookbook-modern-recipes-for","title":"The NLP Cookbook: Modern Recipes for Transformer based Deep Learning Architectures","date":"2021-03-23","arxiv_id":"2104.10640","n_code_links":0,"syntology":null},{"paper":null,"slug":"tmr-evaluating-ner-recall-on-tough-mentions","title":"TMR: Evaluating NER Recall on Tough Mentions","date":"2021-03-23","arxiv_id":"2103.12312","n_code_links":0,"syntology":null},{"paper":null,"slug":"variable-name-recovery-in-decompiled-binary","title":"Variable Name Recovery in Decompiled Binary Code using Constrained Masked Language Modeling","date":"2021-03-23","arxiv_id":"2103.12801","n_code_links":0,"syntology":null},{"paper":null,"slug":"bert-a-review-of-applications-in-natural","title":"BERT: A Review of Applications in Natural Language Processing and Understanding","date":"2021-03-22","arxiv_id":"2103.11943","n_code_links":0,"syntology":null},{"paper":null,"slug":"bridging-the-gap-between-supervised","title":"Bridging the gap between supervised classification and unsupervised topic modelling for social-media assisted crisis management","date":"2021-03-22","arxiv_id":"2103.11835","n_code_links":0,"syntology":null},{"paper":"/paper/hybrid-model-for-patent-classification-using","slug":"hybrid-model-for-patent-classification-using","title":"PatentSBERTa: A Deep NLP based Hybrid Model for Patent Distance and Classification using Augmented SBERT","date":"2021-03-22","arxiv_id":"2103.11933","n_code_links":2,"syntology":null},{"paper":"/paper/identifying-machine-paraphrased-plagiarism","slug":"identifying-machine-paraphrased-plagiarism","title":"Identifying Machine-Paraphrased Plagiarism","date":"2021-03-22","arxiv_id":"2103.11909","n_code_links":2,"syntology":null},{"paper":"/paper/open-domain-question-answering-over-tables","slug":"open-domain-question-answering-over-tables","title":"Open Domain Question Answering over Tables via Dense Retrieval","date":"2021-03-22","arxiv_id":"2103.12011","n_code_links":1,"syntology":null},{"paper":null,"slug":"namerec-highly-accurate-and-fine-grained","title":"NameRec*: Highly Accurate and Fine-grained Person Name Recognition","date":"2021-03-21","arxiv_id":"2103.11360","n_code_links":0,"syntology":null},{"paper":"/paper/rosita-refined-bert-compression-with","slug":"rosita-refined-bert-compression-with","title":"ROSITA: Refined BERT cOmpreSsion with InTegrAted techniques","date":"2021-03-21","arxiv_id":"2103.11367","n_code_links":1,"syntology":{"ran":0,"of":1,"n_ran_checked":0,"n_instrument":0,"unverified":1,"pointer_only":1,"phrase":"0 ran · 1 unverified","official":{"repos":["llyx97/Rosita"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":[]}}},{"paper":null,"slug":"cost-effective-deployment-of-bert-models-in","title":"Cost-effective Deployment of BERT Models in Serverless Environment","date":"2021-03-19","arxiv_id":"2103.10673","n_code_links":0,"syntology":null},{"paper":"/paper/let-your-heart-speak-in-its-mother-tongue","slug":"let-your-heart-speak-in-its-mother-tongue","title":"Let Your Heart Speak in its Mother Tongue: Multilingual Captioning of Cardiac Signals","date":"2021-03-19","arxiv_id":"2103.11011","n_code_links":1,"syntology":null},{"paper":"/paper/muril-multilingual-representations-for-indian","slug":"muril-multilingual-representations-for-indian","title":"MuRIL: Multilingual Representations for Indian Languages","date":"2021-03-19","arxiv_id":"2103.10730","n_code_links":1,"syntology":null},{"paper":"/paper/all-nlp-tasks-are-generation-tasks-a-general","slug":"all-nlp-tasks-are-generation-tasks-a-general","title":"GLM: General Language Model Pretraining with Autoregressive Blank Infilling","date":"2021-03-18","arxiv_id":"2103.10360","n_code_links":8,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["THUDM/GLM"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":null,"slug":"contextual-biasing-of-language-models-for","title":"Contextual Biasing of Language Models for Speech Recognition in Goal-Oriented Conversational Agents","date":"2021-03-18","arxiv_id":"2103.10325","n_code_links":0,"syntology":null},{"paper":"/paper/model-extraction-and-adversarial","slug":"model-extraction-and-adversarial","title":"Model Extraction and Adversarial Transferability, Your BERT is Vulnerable!","date":"2021-03-18","arxiv_id":"2103.10013","n_code_links":1,"syntology":null},{"paper":null,"slug":"code-word-detection-in-fraud-investigations","title":"Code Word Detection in Fraud Investigations using a Deep-Learning Approach","date":"2021-03-17","arxiv_id":"2103.09606","n_code_links":0,"syntology":null},{"paper":"/paper/on-the-role-of-images-for-analyzing-claims-in","slug":"on-the-role-of-images-for-analyzing-claims-in","title":"On the Role of Images for Analyzing Claims in Social Media","date":"2021-03-17","arxiv_id":"2103.09602","n_code_links":1,"syntology":null},{"paper":"/paper/uniparma-semeval-2021-task-5-toxic-spans","slug":"uniparma-semeval-2021-task-5-toxic-spans","title":"UniParma at SemEval-2021 Task 5: Toxic Spans Detection Using CharacterBERT and Bag-of-Words Model","date":"2021-03-17","arxiv_id":"2103.09645","n_code_links":1,"syntology":null},{"paper":null,"slug":"kgsynnet-a-novel-entity-synonyms-discovery","title":"KGSynNet: A Novel Entity Synonyms Discovery Framework with Knowledge Graph","date":"2021-03-16","arxiv_id":"2103.08893","n_code_links":0,"syntology":null},{"paper":null,"slug":"robustly-optimized-and-distilled-training-for","title":"Robustly Optimized and Distilled Training for Natural Language Understanding","date":"2021-03-16","arxiv_id":"2103.08809","n_code_links":0,"syntology":null},{"paper":null,"slug":"text-mining-of-stocktwits-data-for-predicting","title":"Text Mining of Stocktwits Data for Predicting Stock Prices","date":"2021-03-13","arxiv_id":"2103.16388","n_code_links":0,"syntology":null},{"paper":null,"slug":"comparing-the-performance-of-nlp-toolkits-and","title":"Comparing the Performance of NLP Toolkits and Evaluation measures in Legal Tech","date":"2021-03-12","arxiv_id":"2103.11792","n_code_links":0,"syntology":null},{"paper":null,"slug":"explaining-and-improving-bert-performance-on","title":"Explaining and Improving BERT Performance on Lexical Semantic Change Detection","date":"2021-03-12","arxiv_id":"2103.07259","n_code_links":0,"syntology":null},{"paper":null,"slug":"is-bert-a-cross-disciplinary-knowledge","title":"Is BERT a Cross-Disciplinary Knowledge Learner? A Surprising Finding of Pre-trained Models' Transferability","date":"2021-03-12","arxiv_id":"2103.07162","n_code_links":0,"syntology":null},{"paper":null,"slug":"composite-re-ranking-for-efficient-document","title":"Composite Re-Ranking for Efficient Document Search with BERT","date":"2021-03-11","arxiv_id":"2103.06499","n_code_links":0,"syntology":null},{"paper":null,"slug":"evaluation-of-morphological-embeddings-for-1","title":"Evaluation of Morphological Embeddings for the Russian Language","date":"2021-03-11","arxiv_id":"2103.06628","n_code_links":0,"syntology":null},{"paper":null,"slug":"fairfil-contrastive-neural-debiasing-method-1","title":"FairFil: Contrastive Neural Debiasing Method for Pretrained Text Encoders","date":"2021-03-11","arxiv_id":"2103.06413","n_code_links":0,"syntology":null},{"paper":"/paper/improving-bi-encoder-document-ranking-models","slug":"improving-bi-encoder-document-ranking-models","title":"Improving Bi-encoder Document Ranking Models with Two Rankers and Multi-teacher Distillation","date":"2021-03-11","arxiv_id":"2103.06523","n_code_links":1,"syntology":null},{"paper":null,"slug":"lightmbert-a-simple-yet-effective-method-for","title":"LightMBERT: A Simple Yet Effective Method for Multilingual BERT Distillation","date":"2021-03-11","arxiv_id":"2103.06418","n_code_links":0,"syntology":null},{"paper":null,"slug":"self-supervised-text-to-sql-learning-with","title":"Self-supervised Text-to-SQL Learning with Header Alignment Training","date":"2021-03-11","arxiv_id":"2103.06402","n_code_links":0,"syntology":null},{"paper":"/paper/towards-multi-sense-cross-lingual-alignment-1","slug":"towards-multi-sense-cross-lingual-alignment-1","title":"Towards Multi-Sense Cross-Lingual Alignment of Contextual Embeddings","date":"2021-03-11","arxiv_id":"2103.06459","n_code_links":1,"syntology":null},{"paper":null,"slug":"ceqe-contextualized-embeddings-for-query","title":"CEQE: Contextualized Embeddings for Query Expansion","date":"2021-03-09","arxiv_id":"2103.05256","n_code_links":0,"syntology":null},{"paper":"/paper/language-models-have-a-moral-dimension","slug":"language-models-have-a-moral-dimension","title":"Large Pre-trained Language Models Contain Human-like Biases of What is Right and Wrong to Do","date":"2021-03-08","arxiv_id":"2103.11790","n_code_links":1,"syntology":{"ran":6,"of":6,"n_ran_checked":6,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["ml-research/MoRT_NMI"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"skillbert-skilling-the-bert-to-classify-1","title":"SKILLBERT: “SKILLING” THE BERT TO CLASSIFY SKILLS!","date":"2021-03-08","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/syntax-bert-improving-pre-trained","slug":"syntax-bert-improving-pre-trained","title":"Syntax-BERT: Improving Pre-trained Transformers with Syntax Trees","date":"2021-03-07","arxiv_id":"2103.04350","n_code_links":1,"syntology":null},{"paper":null,"slug":"fine-tuning-pretrained-multilingual-bert","title":"Fine-tuning Pretrained Multilingual BERT Model for Indonesian Aspect-based Sentiment Analysis","date":"2021-03-05","arxiv_id":"2103.03732","n_code_links":0,"syntology":null},{"paper":null,"slug":"malbert-using-transformers-for-cybersecurity","title":"MalBERT: Using Transformers for Cybersecurity and Malicious Software Detection","date":"2021-03-05","arxiv_id":"2103.03806","n_code_links":0,"syntology":null},{"paper":null,"slug":"non-invasive-self-attention-for-side","title":"Non-invasive Self-attention for Side Information Fusion in Sequential Recommendation","date":"2021-03-05","arxiv_id":"2103.03578","n_code_links":0,"syntology":null},{"paper":null,"slug":"hardware-acceleration-of-fully-quantized-bert","title":"Hardware Acceleration of Fully Quantized BERT for Efficient Natural Language Processing","date":"2021-03-04","arxiv_id":"2103.02800","n_code_links":0,"syntology":null},{"paper":null,"slug":"few-shot-learning-for-slot-tagging-with","title":"Few-shot Learning for Slot Tagging with Attentive Relational Network","date":"2021-03-03","arxiv_id":"2103.02333","n_code_links":0,"syntology":null},{"paper":null,"slug":"natural-language-understanding-for","title":"Natural Language Understanding for Argumentative Dialogue Systems in the Opinion Building Domain","date":"2021-03-03","arxiv_id":"2103.02691","n_code_links":0,"syntology":null},{"paper":null,"slug":"hate-towards-the-political-opponent-a-twitter","title":"Hate Towards the Political Opponent: A Twitter Corpus Study of the 2020 US Elections on the Basis of Offensive Speech and Stance Detection","date":"2021-03-02","arxiv_id":"2103.01664","n_code_links":0,"syntology":null},{"paper":null,"slug":"bert-based-knowledge-extraction-method-of","title":"BERT-based knowledge extraction method of unstructured domain text","date":"2021-03-01","arxiv_id":"2103.00728","n_code_links":0,"syntology":null},{"paper":null,"slug":"bert-based-patent-novelty-search-by-training","title":"BERT based patent novelty search by training claims to their own description","date":"2021-03-01","arxiv_id":"2103.01126","n_code_links":0,"syntology":null},{"paper":null,"slug":"combat-covid-19-infodemic-using-explainable","title":"Combat COVID-19 Infodemic Using Explainable Natural Language Processing Models","date":"2021-03-01","arxiv_id":"2103.00747","n_code_links":0,"syntology":null},{"paper":"/paper/nlp-cuet-dravidianlangtech-eacl2021-offensive","slug":"nlp-cuet-dravidianlangtech-eacl2021-offensive","title":"NLP-CUET@DravidianLangTech-EACL2021: Offensive Language Detection from Multilingual Code-Mixed Text using Transformers","date":"2021-02-28","arxiv_id":"2103.00455","n_code_links":1,"syntology":null},{"paper":"/paper/nlp-cuet-lt-edi-eacl2021-multilingual-code","slug":"nlp-cuet-lt-edi-eacl2021-multilingual-code","title":"NLP-CUET@LT-EDI-EACL2021: Multilingual Code-Mixed Hope Speech Detection using Cross-lingual Representation Learner","date":"2021-02-28","arxiv_id":"2103.00464","n_code_links":1,"syntology":null},{"paper":"/paper/covid-19-tweets-analysis-through-transformer","slug":"covid-19-tweets-analysis-through-transformer","title":"COVID-19 Tweets Analysis through Transformer Language Models","date":"2021-02-27","arxiv_id":"2103.00199","n_code_links":1,"syntology":null},{"paper":null,"slug":"transformers-with-competitive-ensembles-of-1","title":"Transformers with Competitive Ensembles of Independent Mechanisms","date":"2021-02-27","arxiv_id":"2103.00336","n_code_links":0,"syntology":null},{"paper":null,"slug":"multi-task-transfer-learning-for-finding","title":"Multi-task transfer learning for finding actionable information from crisis-related messages on social media","date":"2021-02-26","arxiv_id":"2102.13395","n_code_links":0,"syntology":null},{"paper":null,"slug":"bert-based-acronym-disambiguation-with","title":"BERT-based Acronym Disambiguation with Multiple Training Strategies","date":"2021-02-25","arxiv_id":"2103.00488","n_code_links":0,"syntology":null},{"paper":null,"slug":"emotion-aware-emotion-agnostic-or-automatic","title":"Emotion-Aware, Emotion-Agnostic, or Automatic: Corpus Creation Strategies to Obtain Cognitive Event Appraisal Annotations","date":"2021-02-25","arxiv_id":"2102.12858","n_code_links":0,"syntology":null},{"paper":null,"slug":"pharmke-knowledge-extraction-platform-for","title":"PharmKE: Knowledge Extraction Platform for Pharmaceutical Texts using Transfer Learning","date":"2021-02-25","arxiv_id":"2102.13139","n_code_links":0,"syntology":null},{"paper":"/paper/sentiment-analysis-of-persian-english-code","slug":"sentiment-analysis-of-persian-english-code","title":"Sentiment Analysis of Persian-English Code-mixed Texts","date":"2021-02-25","arxiv_id":"2102.12700","n_code_links":1,"syntology":null},{"paper":null,"slug":"from-universal-language-model-to-downstream","title":"From Universal Language Model to Downstream Task: Improving RoBERTa-Based Vietnamese Hate Speech Detection","date":"2021-02-24","arxiv_id":"2102.12162","n_code_links":0,"syntology":null},{"paper":null,"slug":"hopeful-men-lt-edi-eacl2021-hope-speech","title":"Hopeful_Men@LT-EDI-EACL2021: Hope Speech Detection Using Indic Transliteration and Transformers","date":"2021-02-24","arxiv_id":"2102.12082","n_code_links":0,"syntology":null},{"paper":"/paper/lrg-at-semeval-2021-task-4-improving-reading","slug":"lrg-at-semeval-2021-task-4-improving-reading","title":"LRG at SemEval-2021 Task 4: Improving Reading Comprehension with Abstract Words using Augmentation, Linguistic Features and Voting","date":"2021-02-24","arxiv_id":"2102.12255","n_code_links":1,"syntology":null},{"paper":"/paper/nlrg-at-semeval-2021-task-5-toxic-spans","slug":"nlrg-at-semeval-2021-task-5-toxic-spans","title":"NLRG at SemEval-2021 Task 5: Toxic Spans Detection Leveraging BERT-based Token Classification and Span Prediction Techniques","date":"2021-02-24","arxiv_id":"2102.12254","n_code_links":1,"syntology":null},{"paper":null,"slug":"task-specific-pre-training-and-cross-lingual","title":"Task-Specific Pre-Training and Cross Lingual Transfer for Code-Switched Data","date":"2021-02-24","arxiv_id":"2102.12407","n_code_links":0,"syntology":null},{"paper":null,"slug":"minimally-supervised-structure-rich-text","title":"Minimally-Supervised Structure-Rich Text Categorization via Learning on Text-Rich Networks","date":"2021-02-23","arxiv_id":"2102.11479","n_code_links":0,"syntology":null},{"paper":null,"slug":"robust-and-transferable-anomaly-detection-in","title":"Robust and Transferable Anomaly Detection in Log Data using Pre-Trained Language Models","date":"2021-02-23","arxiv_id":"2102.11570","n_code_links":0,"syntology":null},{"paper":"/paper/visualchexbert-addressing-the-discrepancy","slug":"visualchexbert-addressing-the-discrepancy","title":"VisualCheXbert: Addressing the Discrepancy Between Radiology Report Labels and Image Labels","date":"2021-02-23","arxiv_id":"2102.11467","n_code_links":1,"syntology":{"ran":0,"of":2,"n_ran_checked":0,"n_instrument":0,"unverified":2,"pointer_only":2,"phrase":"0 ran · 2 unverified","official":{"repos":["stanfordmlgroup/VisualCheXbert"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":2,"ran_from_kinds":[]}}},{"paper":"/paper/evaluating-contextualized-language-models-for","slug":"evaluating-contextualized-language-models-for","title":"Evaluating Contextualized Language Models for Hungarian","date":"2021-02-22","arxiv_id":"2102.10848","n_code_links":1,"syntology":null},{"paper":null,"slug":"generating-human-readable-transcript-for","title":"Generating Human Readable Transcript for Automatic Speech Recognition with Pre-trained Language Model","date":"2021-02-22","arxiv_id":"2102.11114","n_code_links":0,"syntology":null},{"paper":null,"slug":"mixup-training-leads-to-reduced-overfitting","title":"MixUp Training Leads to Reduced Overfitting and Improved Calibration for the Transformer Architecture","date":"2021-02-22","arxiv_id":"2102.11402","n_code_links":0,"syntology":null},{"paper":"/paper/parallelizing-legendre-memory-unit-training","slug":"parallelizing-legendre-memory-unit-training","title":"Parallelizing Legendre Memory Unit Training","date":"2021-02-22","arxiv_id":"2102.11417","n_code_links":2,"syntology":null},{"paper":null,"slug":"rubert-a-bilingual-roman-urdu-bert-using","title":"RUBERT: A Bilingual Roman Urdu BERT Using Cross Lingual Transfer Learning","date":"2021-02-22","arxiv_id":"2102.11278","n_code_links":0,"syntology":null},{"paper":"/paper/using-prior-knowledge-to-guide-bert-s","slug":"using-prior-knowledge-to-guide-bert-s","title":"Using Prior Knowledge to Guide BERT's Attention in Semantic Textual Matching Tasks","date":"2021-02-22","arxiv_id":"2102.10934","n_code_links":1,"syntology":null},{"paper":null,"slug":"pre-training-bert-on-arabic-tweets-practical","title":"Pre-Training BERT on Arabic Tweets: Practical Considerations","date":"2021-02-21","arxiv_id":"2102.10684","n_code_links":0,"syntology":null},{"paper":null,"slug":"web-based-application-for-detecting","title":"Web-based Application for Detecting Indonesian Clickbait Headlines using IndoBERT","date":"2021-02-21","arxiv_id":"2102.10601","n_code_links":0,"syntology":null},{"paper":null,"slug":"learning-dynamic-bert-via-trainable-gate","title":"Learning Dynamic BERT via Trainable Gate Variables and a Bi-modal Regularizer","date":"2021-02-19","arxiv_id":"2102.09727","n_code_links":0,"syntology":null},{"paper":"/paper/towards-emotion-recognition-in-hindi-english","slug":"towards-emotion-recognition-in-hindi-english","title":"Towards Emotion Recognition in Hindi-English Code-Mixed Data: A Transformer Based Approach","date":"2021-02-19","arxiv_id":"2102.09943","n_code_links":1,"syntology":null},{"paper":"/paper/using-transformer-based-ensemble-learning-to","slug":"using-transformer-based-ensemble-learning-to","title":"Using Transformer based Ensemble Learning to classify Scientific Articles","date":"2021-02-19","arxiv_id":"2102.09991","n_code_links":2,"syntology":null},{"paper":"/paper/analysis-of-contextual-and-non-contextual","slug":"analysis-of-contextual-and-non-contextual","title":"Analysis Of Contextual and Non-Contextual Word Embedding Models For Hindi NER With Web Application For Data Collection","date":"2021-02-18","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/training-microsoft-news-recommenders-with","slug":"training-microsoft-news-recommenders-with","title":"Training Large-Scale News Recommenders with Pretrained Language Models in the Loop","date":"2021-02-18","arxiv_id":"2102.09268","n_code_links":1,"syntology":null},{"paper":null,"slug":"unibuckernel-geolocating-swiss-german-jodels","title":"UnibucKernel: Geolocating Swiss German Jodels Using Ensemble Learning","date":"2021-02-18","arxiv_id":"2102.09379","n_code_links":0,"syntology":null},{"paper":null,"slug":"leveraging-query-resolution-and-reading","title":"Leveraging Query Resolution and Reading Comprehension for Conversational Passage Retrieval","date":"2021-02-17","arxiv_id":"2102.08795","n_code_links":0,"syntology":null},{"paper":"/paper/scidr-at-sdu-2020-ideas-identifying-and","slug":"scidr-at-sdu-2020-ideas-identifying-and","title":"SciDr at SDU-2020: IDEAS -- Identifying and Disambiguating Everyday Acronyms for Scientific Domain","date":"2021-02-17","arxiv_id":"2102.08818","n_code_links":2,"syntology":null},{"paper":"/paper/tcn-table-convolutional-network-for-web-table","slug":"tcn-table-convolutional-network-for-web-table","title":"TCN: Table Convolutional Network for Web Table Interpretation","date":"2021-02-17","arxiv_id":"2102.09460","n_code_links":1,"syntology":null},{"paper":"/paper/coco-lm-correcting-and-contrasting-text","slug":"coco-lm-correcting-and-contrasting-text","title":"COCO-LM: Correcting and Contrasting Text Sequences for Language Model Pretraining","date":"2021-02-16","arxiv_id":"2102.08473","n_code_links":2,"syntology":{"ran":5,"of":6,"n_ran_checked":1,"n_instrument":4,"unverified":1,"pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 4 where Syntology's instrument failed) · 1 unverified","official":{"repos":["microsoft/coco-lm"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/exploring-transformers-in-natural-language","slug":"exploring-transformers-in-natural-language","title":"Exploring Transformers in Natural Language Generation: GPT, BERT, and XLNet","date":"2021-02-16","arxiv_id":"2102.08036","n_code_links":1,"syntology":null},{"paper":null,"slug":"have-attention-heads-in-bert-learned","title":"Have Attention Heads in BERT Learned Constituency Grammar?","date":"2021-02-16","arxiv_id":"2102.07926","n_code_links":0,"syntology":null},{"paper":"/paper/non-autoregressive-text-generation-with-pre","slug":"non-autoregressive-text-generation-with-pre","title":"Non-Autoregressive Text Generation with Pre-trained Language Models","date":"2021-02-16","arxiv_id":"2102.08220","n_code_links":1,"syntology":null},{"paper":"/paper/dobf-a-deobfuscation-pre-training-objective","slug":"dobf-a-deobfuscation-pre-training-objective","title":"DOBF: A Deobfuscation Pre-Training Objective for Programming Languages","date":"2021-02-15","arxiv_id":"2102.07492","n_code_links":2,"syntology":null},{"paper":null,"slug":"fast-end-to-end-speech-recognition-via-non","title":"Fast End-to-End Speech Recognition via Non-Autoregressive Models and Cross-Modal Knowledge Transferring from BERT","date":"2021-02-15","arxiv_id":"2102.07594","n_code_links":0,"syntology":null},{"paper":null,"slug":"improved-customer-transaction-classification","title":"Improved Customer Transaction Classification using Semi-Supervised Knowledge Distillation","date":"2021-02-15","arxiv_id":"2102.07635","n_code_links":0,"syntology":null},{"paper":null,"slug":"within-document-event-coreference-with-bert","title":"Within-Document Event Coreference with BERT-Based Contextualized Representations","date":"2021-02-15","arxiv_id":"2102.09600","n_code_links":0,"syntology":null}],"record_sha256":"07aed7b7308f16334c380eeca1a7b8da3c62ebbc21ffb2a53b7403678672fc9d","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}