{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/bert/papers/67","list_of":"/method/bert","method":"BERT","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":67,"pages_in_order":70,"rows_per_page":100,"rows":[6601,6700],"of":6938,"counts":{"archive_papers_tagged":6938,"with_a_code_link":2862,"where_syntology_ran_a_sample":640,"not_listed_spam_title":0,"listed":6938,"listed_where_code_ran":640,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":520,"every_run_a_failure_of_syntologys_instrument":120,"listed_with_a_run_with_no_instrument_failure":520,"listed_every_run_a_failure_of_syntologys_instrument":120,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/bert","prev":"/method/bert/papers/66","next":"/method/bert/papers/68","papers":[{"paper":"/paper/portuguese-named-entity-recognition-using-1","slug":"portuguese-named-entity-recognition-using-1","title":"Portuguese Named Entity Recognition using BERT-CRF","date":"2019-09-23","arxiv_id":"1909.10649","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":2,"n_instrument":1,"unverified":0,"pointer_only":3,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 2 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["neuralmind-ai/portuguese-bert"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"bert-meets-chinese-word-segmentation","title":"BERT Meets Chinese Word Segmentation","date":"2019-09-20","arxiv_id":"1909.09292","n_code_links":0,"syntology":null},{"paper":"/paper/allennlp-interpret-a-framework-for-explaining","slug":"allennlp-interpret-a-framework-for-explaining","title":"AllenNLP Interpret: A Framework for Explaining Predictions of NLP Models","date":"2019-09-19","arxiv_id":"1909.09251","n_code_links":1,"syntology":null},{"paper":null,"slug":"exploring-ways-to-incorporate-additional","title":"How Additional Knowledge can Improve Natural Language Commonsense Question Answering?","date":"2019-09-19","arxiv_id":"1909.08855","n_code_links":0,"syntology":null},{"paper":"/paper/summary-level-training-of-sentence-rewriting","slug":"summary-level-training-of-sentence-rewriting","title":"Summary Level Training of Sentence Rewriting for Abstractive Summarization","date":"2019-09-19","arxiv_id":"1909.08752","n_code_links":0,"syntology":null},{"paper":"/paper/enriching-bert-with-knowledge-graph","slug":"enriching-bert-with-knowledge-graph","title":"Enriching BERT with Knowledge Graph Embeddings for Document Classification","date":"2019-09-18","arxiv_id":"1909.08402","n_code_links":1,"syntology":null},{"paper":"/paper/improving-natural-language-inference-with-a","slug":"improving-natural-language-inference-with-a","title":"Improving Natural Language Inference with a Pretrained Parser","date":"2019-09-18","arxiv_id":"1909.08217","n_code_links":1,"syntology":null},{"paper":"/paper/language-models-and-automated-essay-scoring","slug":"language-models-and-automated-essay-scoring","title":"Language models and Automated Essay Scoring","date":"2019-09-18","arxiv_id":"1909.09482","n_code_links":1,"syntology":null},{"paper":null,"slug":"using-bert-for-word-sense-disambiguation","title":"Using BERT for Word Sense Disambiguation","date":"2019-09-18","arxiv_id":"1909.08358","n_code_links":0,"syntology":null},{"paper":"/paper/do-nlp-models-know-numbers-probing-numeracy","slug":"do-nlp-models-know-numbers-probing-numeracy","title":"Do NLP Models Know Numbers? Probing Numeracy in Embeddings","date":"2019-09-17","arxiv_id":"1909.07940","n_code_links":1,"syntology":null},{"paper":"/paper/extracting-evidence-of-supplement-drug","slug":"extracting-evidence-of-supplement-drug","title":"SUPP.AI: Finding Evidence for Supplement-Drug Interactions","date":"2019-09-17","arxiv_id":"1909.08135","n_code_links":1,"syntology":null},{"paper":"/paper/k-bert-enabling-language-representation-with","slug":"k-bert-enabling-language-representation-with","title":"K-BERT: Enabling Language Representation with Knowledge Graph","date":"2019-09-17","arxiv_id":"1909.07606","n_code_links":2,"syntology":null},{"paper":"/paper/megatron-lm-training-multi-billion-parameter","slug":"megatron-lm-training-multi-billion-parameter","title":"Megatron-LM: Training Multi-Billion Parameter Language Models Using Model Parallelism","date":"2019-09-17","arxiv_id":"1909.08053","n_code_links":10,"syntology":{"ran":12,"of":47,"n_ran_checked":7,"n_instrument":5,"unverified":35,"pointer_only":15,"phrase":"12 ran (of which 2 constructed an object rather than computing a result; 7 with no instrument failure: 4 honoured, 0 violated, 3 with no contract checked; 5 where Syntology's instrument failed) · 35 unverified","official":{"repos":["NVIDIA/Megatron-LM"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":null,"slug":"simple-yet-effective-bridge-reasoning-for","title":"Simple yet Effective Bridge Reasoning for Open-Domain Multi-Hop Question Answering","date":"2019-09-17","arxiv_id":"1909.07597","n_code_links":0,"syntology":null},{"paper":"/paper/span-based-joint-entity-and-relation","slug":"span-based-joint-entity-and-relation","title":"Span-based Joint Entity and Relation Extraction with Transformer Pre-training","date":"2019-09-17","arxiv_id":"1909.07755","n_code_links":3,"syntology":null},{"paper":"/paper/probing-natural-language-inference-models","slug":"probing-natural-language-inference-models","title":"Probing Natural Language Inference Models through Semantic Fragments","date":"2019-09-16","arxiv_id":"1909.07521","n_code_links":3,"syntology":null},{"paper":"/paper/cross-lingual-bert-transformation-for-zero","slug":"cross-lingual-bert-transformation-for-zero","title":"Cross-Lingual BERT Transformation for Zero-Shot Dependency Parsing","date":"2019-09-15","arxiv_id":"1909.06775","n_code_links":1,"syntology":null},{"paper":"/paper/tree-transformer-integrating-tree-structures","slug":"tree-transformer-integrating-tree-structures","title":"Tree Transformer: Integrating Tree Structures into Self-Attention","date":"2019-09-14","arxiv_id":"1909.06639","n_code_links":3,"syntology":{"ran":9,"of":10,"n_ran_checked":7,"n_instrument":2,"unverified":1,"pointer_only":2,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 2 where Syntology's instrument failed) · 1 unverified","official":{"repos":["yaushian/Tree-Transformer"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/addressing-semantic-drift-in-question","slug":"addressing-semantic-drift-in-question","title":"Addressing Semantic Drift in Question Generation for Semi-Supervised Question Answering","date":"2019-09-13","arxiv_id":"1909.06356","n_code_links":2,"syntology":{"ran":13,"of":19,"n_ran_checked":11,"n_instrument":2,"unverified":6,"pointer_only":1,"phrase":"13 ran (of which 0 constructed an object rather than computing a result; 11 with no instrument failure: 2 honoured, 1 violated, 8 with no contract checked; 2 where Syntology's instrument failed) · 6 unverified","official":{"repos":["ZhangShiyue/QGforQA"],"state":"official (archive's flag): 13 ran","n_ran":13,"n_constructed":0,"n_ran_no_instrument_failure":11,"n_unverified":6,"ran_from_kinds":["official"]}}},{"paper":"/paper/towards-open-domain-named-entity-recognition","slug":"towards-open-domain-named-entity-recognition","title":"Neural Correction Model for Open-Domain Named Entity Recognition","date":"2019-09-13","arxiv_id":"1909.06058","n_code_links":1,"syntology":null},{"paper":null,"slug":"measuring-domain-portability-and","title":"Measuring Domain Portability and ErrorPropagation in Biomedical QA","date":"2019-09-12","arxiv_id":"1909.09704","n_code_links":0,"syntology":null},{"paper":"/paper/q-bert-hessian-based-ultra-low-precision","slug":"q-bert-hessian-based-ultra-low-precision","title":"Q-BERT: Hessian Based Ultra Low Precision Quantization of BERT","date":"2019-09-12","arxiv_id":"1909.05840","n_code_links":0,"syntology":null},{"paper":"/paper/uer-an-open-source-toolkit-for-pre-training","slug":"uer-an-open-source-toolkit-for-pre-training","title":"UER: An Open-Source Toolkit for Pre-training Models","date":"2019-09-12","arxiv_id":"1909.05658","n_code_links":2,"syntology":{"ran":1,"of":3,"n_ran_checked":1,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["dbiir/UER-py"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/bertgrid-contextualized-embedding-for-2d","slug":"bertgrid-contextualized-embedding-for-2d","title":"BERTgrid: Contextualized Embedding for 2D Document Representation and Understanding","date":"2019-09-11","arxiv_id":"1909.04948","n_code_links":2,"syntology":null},{"paper":"/paper/from-english-to-code-switching-transfer","slug":"from-english-to-code-switching-transfer","title":"From English to Code-Switching: Transfer Learning with Strong Morphological Clues","date":"2019-09-11","arxiv_id":"1909.05158","n_code_links":1,"syntology":null},{"paper":"/paper/frustratingly-easy-natural-question-answering","slug":"frustratingly-easy-natural-question-answering","title":"Frustratingly Easy Natural Question Answering","date":"2019-09-11","arxiv_id":"1909.05286","n_code_links":0,"syntology":null},{"paper":"/paper/how-does-bert-answer-questions-a-layer-wise","slug":"how-does-bert-answer-questions-a-layer-wise","title":"How Does BERT Answer Questions? A Layer-Wise Analysis of Transformer Representations","date":"2019-09-11","arxiv_id":"1909.04925","n_code_links":2,"syntology":{"ran":1,"of":2,"n_ran_checked":1,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["bvanaken/explain-BERT-QA"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"bert-based-arabic-social-media","title":"BERT-Based Arabic Social Media Author Profiling","date":"2019-09-09","arxiv_id":"1909.04181","n_code_links":0,"syntology":null},{"paper":"/paper/knowledge-enhanced-contextual-word","slug":"knowledge-enhanced-contextual-word","title":"Knowledge Enhanced Contextual Word Representations","date":"2019-09-09","arxiv_id":"1909.04164","n_code_links":1,"syntology":null},{"paper":"/paper/pretrained-language-models-for-sequential","slug":"pretrained-language-models-for-sequential","title":"Pretrained Language Models for Sequential Sentence Classification","date":"2019-09-09","arxiv_id":"1909.04054","n_code_links":1,"syntology":null},{"paper":"/paper/reasoning-over-semantic-level-graph-for-fact","slug":"reasoning-over-semantic-level-graph-for-fact","title":"Reasoning Over Semantic-Level Graph for Fact Checking","date":"2019-09-09","arxiv_id":"1909.03745","n_code_links":0,"syntology":null},{"paper":"/paper/span-selection-pre-training-for-question","slug":"span-selection-pre-training-for-question","title":"Span Selection Pre-training for Question Answering","date":"2019-09-09","arxiv_id":"1909.04120","n_code_links":1,"syntology":{"ran":6,"of":6,"n_ran_checked":6,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["IBM/span-selection-pretraining"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"commonsense-knowledge-bert-for-level-2","title":"Commonsense Knowledge + BERT for Level 2 Reading Comprehension Ability Test","date":"2019-09-08","arxiv_id":"1909.03415","n_code_links":0,"syntology":null},{"paper":null,"slug":"czech-text-processing-with-contextual","title":"Czech Text Processing with Contextual Embeddings: POS Tagging, Lemmatization, Parsing and NER","date":"2019-09-08","arxiv_id":"1909.03544","n_code_links":0,"syntology":null},{"paper":"/paper/entity-relation-and-event-extraction-with","slug":"entity-relation-and-event-extraction-with","title":"Entity, Relation, and Event Extraction with Contextualized Span Representations","date":"2019-09-08","arxiv_id":"1909.03546","n_code_links":4,"syntology":{"ran":5,"of":7,"n_ran_checked":5,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["dwadden/dygiepp"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":null,"slug":"multi-task-bidirectional-transformer","title":"Multi-Task Bidirectional Transformer Representations for Irony Detection","date":"2019-09-08","arxiv_id":"1909.03526","n_code_links":0,"syntology":null},{"paper":"/paper/symmetric-regularization-based-bert-for-pair","slug":"symmetric-regularization-based-bert-for-pair","title":"Symmetric Regularization based BERT for Pair-wise Semantic Reasoning","date":"2019-09-08","arxiv_id":"1909.03405","n_code_links":1,"syntology":null},{"paper":"/paper/transfer-learning-robustness-in-multi-class","slug":"transfer-learning-robustness-in-multi-class","title":"Transfer Learning Robustness in Multi-Class Categorization by Fine-Tuning Pre-Trained Contextualized Language Models","date":"2019-09-08","arxiv_id":"1909.03564","n_code_links":1,"syntology":null},{"paper":"/paper/a-novel-hierarchical-binary-tagging-framework","slug":"a-novel-hierarchical-binary-tagging-framework","title":"A Novel Cascade Binary Tagging Framework for Relational Triple Extraction","date":"2019-09-07","arxiv_id":"1909.03227","n_code_links":5,"syntology":null},{"paper":"/paper/rnn-architecture-learning-with-sparse","slug":"rnn-architecture-learning-with-sparse","title":"RNN Architecture Learning with Sparse Regularization","date":"2019-09-06","arxiv_id":"1909.03011","n_code_links":1,"syntology":null},{"paper":"/paper/supervised-multimodal-bitransformers-for","slug":"supervised-multimodal-bitransformers-for","title":"Supervised Multimodal Bitransformers for Classifying Images and Text","date":"2019-09-06","arxiv_id":"1909.02950","n_code_links":6,"syntology":null},{"paper":"/paper/effective-use-of-transformer-networks-for","slug":"effective-use-of-transformer-networks-for","title":"Effective Use of Transformer Networks for Entity Tracking","date":"2019-09-05","arxiv_id":"1909.02635","n_code_links":1,"syntology":{"ran":3,"of":5,"n_ran_checked":2,"n_instrument":1,"unverified":2,"pointer_only":5,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 2 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","official":{"repos":["aditya2211/transformer-entity-tracking"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":2,"ran_from_kinds":["official","unlocated"]}}},{"paper":"/paper/in-plain-sight-media-bias-through-the-lens-of","slug":"in-plain-sight-media-bias-through-the-lens-of","title":"In Plain Sight: Media Bias Through the Lens of Factual Reporting","date":"2019-09-05","arxiv_id":"1909.02670","n_code_links":1,"syntology":null},{"paper":"/paper/informing-unsupervised-pretraining-with","slug":"informing-unsupervised-pretraining-with","title":"Specializing Unsupervised Pretraining Models for Word-Level Semantic Similarity","date":"2019-09-05","arxiv_id":"1909.02339","n_code_links":1,"syntology":null},{"paper":"/paper/investigating-berts-knowledge-of-language","slug":"investigating-berts-knowledge-of-language","title":"Investigating BERT's Knowledge of Language: Five Analysis Methods with NPIs","date":"2019-09-05","arxiv_id":"1909.02597","n_code_links":1,"syntology":{"ran":0,"of":3,"n_ran_checked":0,"n_instrument":0,"unverified":3,"pointer_only":3,"phrase":"0 ran · 3 unverified","official":{"repos":["alexwarstadt/data_generation"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":3,"ran_from_kinds":[]}}},{"paper":"/paper/semantics-aware-bert-for-language","slug":"semantics-aware-bert-for-language","title":"Semantics-aware BERT for Language Understanding","date":"2019-09-05","arxiv_id":"1909.02209","n_code_links":1,"syntology":{"ran":7,"of":12,"n_ran_checked":5,"n_instrument":2,"unverified":5,"pointer_only":3,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 2 honoured, 0 violated, 3 with no contract checked; 2 where Syntology's instrument failed) · 5 unverified","official":{"repos":["cooelf/SemBERT"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":5,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"syntax-aware-aspect-level-sentiment","title":"Syntax-Aware Aspect Level Sentiment Classification with Graph Attention Networks","date":"2019-09-05","arxiv_id":"1909.02606","n_code_links":0,"syntology":null},{"paper":"/paper/encode-tag-realize-high-precision-text","slug":"encode-tag-realize-high-precision-text","title":"Encode, Tag, Realize: High-Precision Text Editing","date":"2019-09-03","arxiv_id":"1909.01187","n_code_links":5,"syntology":{"ran":8,"of":11,"n_ran_checked":7,"n_instrument":1,"unverified":3,"pointer_only":1,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 1 honoured, 0 violated, 6 with no contract checked; 1 where Syntology's instrument failed) · 3 unverified","official":{"repos":["google-research/lasertagger"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":1,"ran_from_kinds":["listed","official"]}}},{"paper":"/paper/language-models-as-knowledge-bases","slug":"language-models-as-knowledge-bases","title":"Language Models as Knowledge Bases?","date":"2019-09-03","arxiv_id":"1909.01066","n_code_links":1,"syntology":null},{"paper":null,"slug":"multimodal-deep-learning-for-mental-disorders","title":"Multimodal Deep Learning for Mental Disorders Prediction from Audio Speech Samples","date":"2019-09-03","arxiv_id":"1909.01067","n_code_links":0,"syntology":null},{"paper":"/paper/transfer-fine-tuning-a-bert-case-study","slug":"transfer-fine-tuning-a-bert-case-study","title":"Transfer Fine-Tuning: A BERT Case Study","date":"2019-09-03","arxiv_id":"1909.00931","n_code_links":1,"syntology":null},{"paper":null,"slug":"unicoder-a-universal-language-encoder-by-pre","title":"Unicoder: A Universal Language Encoder by Pre-training with Multiple Cross-lingual Tasks","date":"2019-09-03","arxiv_id":"1909.00964","n_code_links":0,"syntology":null},{"paper":"/paper/how-contextual-are-contextualized-word","slug":"how-contextual-are-contextualized-word","title":"How Contextual are Contextualized Word Representations? Comparing the Geometry of BERT, ELMo, and GPT-2 Embeddings","date":"2019-09-02","arxiv_id":"1909.00512","n_code_links":1,"syntology":null},{"paper":"/paper/sumqe-a-bert-based-summary-quality-estimation","slug":"sumqe-a-bert-based-summary-quality-estimation","title":"SumQE: a BERT-based Summary Quality Estimation Model","date":"2019-09-02","arxiv_id":"1909.00578","n_code_links":1,"syntology":null},{"paper":null,"slug":"classification-approaches-to-identify","title":"Classification Approaches to Identify Informative Tweets","date":"2019-09-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/cross-lingual-machine-reading-comprehension","slug":"cross-lingual-machine-reading-comprehension","title":"Cross-Lingual Machine Reading Comprehension","date":"2019-09-01","arxiv_id":"1909.00361","n_code_links":1,"syntology":null},{"paper":null,"slug":"evaluating-the-cross-lingual-effectiveness-of","title":"Evaluating the Cross-Lingual Effectiveness of Massively Multilingual Neural Machine Translation","date":"2019-09-01","arxiv_id":"1909.00437","n_code_links":0,"syntology":null},{"paper":null,"slug":"evaluation-of-stacked-embeddings-for","title":"Evaluation of Stacked Embeddings for Bulgarian on the Downstream Tasks POS and NERC","date":"2019-09-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"evaluation-of-vector-embedding-models-in","title":"Evaluation of vector embedding models in clustering of text documents","date":"2019-09-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"friendsqa-open-domain-question-answering-on","title":"FriendsQA: Open-Domain Question Answering on TV Show Transcripts","date":"2019-09-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/incidental-supervision-from-question","slug":"incidental-supervision-from-question","title":"QuASE: Question-Answer Driven Sentence Encoding","date":"2019-09-01","arxiv_id":"1909.00333","n_code_links":1,"syntology":null},{"paper":null,"slug":"multilingual-language-models-for-named-entity","title":"Multilingual Language Models for Named Entity Recognition in German and English","date":"2019-09-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"multilingual-probing-of-deep-pre-trained","title":"Multilingual Probing of Deep Pre-Trained Contextual Encoders","date":"2019-09-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"predicting-sentiment-of-polish-language-short","title":"Predicting Sentiment of Polish Language Short Texts","date":"2019-09-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"semantic-role-labeling-with-pretrained","title":"Semantic Role Labeling with Pretrained Language Models for Known and Unknown Predicates","date":"2019-09-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"adversarial-learning-with-contextual","title":"Adversarial Learning with Contextual Embeddings for Zero-resource Cross-lingual Classification and NER","date":"2019-08-31","arxiv_id":"1909.00153","n_code_links":0,"syntology":null},{"paper":"/paper/evaluation-benchmarks-and-learning","slug":"evaluation-benchmarks-and-learning","title":"Evaluation Benchmarks and Learning Criteria for Discourse-Aware Sentence Representations","date":"2019-08-31","arxiv_id":"1909.00142","n_code_links":2,"syntology":{"ran":2,"of":2,"n_ran_checked":1,"n_instrument":1,"unverified":0,"pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["ZeweiChu/DiscoEval"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"knowledge-enhanced-attention-for-robust","title":"Knowledge Enhanced Attention for Robust Natural Language Inference","date":"2019-08-31","arxiv_id":"1909.00102","n_code_links":0,"syntology":null},{"paper":"/paper/nezha-neural-contextualized-representation","slug":"nezha-neural-contextualized-representation","title":"NEZHA: Neural Contextualized Representation for Chinese Language Understanding","date":"2019-08-31","arxiv_id":"1909.00204","n_code_links":10,"syntology":null},{"paper":null,"slug":"quantity-doesnt-buy-quality-syntax-with","title":"Quantity doesn't buy quality syntax with neural language models","date":"2019-08-31","arxiv_id":"1909.00111","n_code_links":0,"syntology":null},{"paper":null,"slug":"small-and-practical-bert-models-for-sequence","title":"Small and Practical BERT Models for Sequence Labeling","date":"2019-08-31","arxiv_id":"1909.00100","n_code_links":0,"syntology":null},{"paper":"/paper/adapt-or-get-left-behind-domain-adaptation","slug":"adapt-or-get-left-behind-domain-adaptation","title":"Adapt or Get Left Behind: Domain Adaptation through BERT Language Model Finetuning for Aspect-Target Sentiment Classification","date":"2019-08-30","arxiv_id":"1908.11860","n_code_links":3,"syntology":{"ran":6,"of":7,"n_ran_checked":6,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["deepopinion/domain-adapted-atsc"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":["listed"]}}},{"paper":"/paper/bilingual-is-at-least-monolingual-balm-a","slug":"bilingual-is-at-least-monolingual-balm-a","title":"Bilingual is At Least Monolingual (BALM): A Novel Translation Algorithm that Encodes Monolingual Priors","date":"2019-08-30","arxiv_id":"1909.01146","n_code_links":1,"syntology":null},{"paper":"/paper/paws-x-a-cross-lingual-adversarial-dataset","slug":"paws-x-a-cross-lingual-adversarial-dataset","title":"PAWS-X: A Cross-lingual Adversarial Dataset for Paraphrase Identification","date":"2019-08-30","arxiv_id":"1908.11828","n_code_links":3,"syntology":null},{"paper":null,"slug":"adversarial-representation-learning-for-text","title":"Adversarial Representation Learning for Text-to-Image Matching","date":"2019-08-28","arxiv_id":"1908.10534","n_code_links":0,"syntology":null},{"paper":"/paper/finbert-financial-sentiment-analysis-with-pre","slug":"finbert-financial-sentiment-analysis-with-pre","title":"FinBERT: Financial Sentiment Analysis with Pre-trained Language Models","date":"2019-08-27","arxiv_id":"1908.10063","n_code_links":3,"syntology":{"ran":3,"of":3,"n_ran_checked":3,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 1 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":null}},{"paper":"/paper/sentence-bert-sentence-embeddings-using","slug":"sentence-bert-sentence-embeddings-using","title":"Sentence-BERT: Sentence Embeddings using Siamese BERT-Networks","date":"2019-08-27","arxiv_id":"1908.10084","n_code_links":64,"syntology":{"ran":33,"of":58,"n_ran_checked":30,"n_instrument":3,"unverified":25,"pointer_only":11,"phrase":"33 ran (of which 9 constructed an object rather than computing a result; 30 with no instrument failure: 1 honoured, 0 violated, 29 with no contract checked; 3 where Syntology's instrument failed) · 25 unverified","official":{"repos":["UKPLab/sentence-transformers"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/attentive-history-selection-for","slug":"attentive-history-selection-for","title":"Attentive History Selection for Conversational Question Answering","date":"2019-08-26","arxiv_id":"1908.09456","n_code_links":2,"syntology":null},{"paper":"/paper/detecting-toxicity-in-news-articles","slug":"detecting-toxicity-in-news-articles","title":"Detecting Toxicity in News Articles: Application to Bulgarian","date":"2019-08-26","arxiv_id":"1908.09785","n_code_links":1,"syntology":null},{"paper":"/paper/does-bert-agree-evaluating-knowledge-of","slug":"does-bert-agree-evaluating-knowledge-of","title":"Does BERT agree? Evaluating knowledge of structure dependence through agreement relations","date":"2019-08-26","arxiv_id":"1908.09892","n_code_links":1,"syntology":null},{"paper":null,"slug":"measuring-patent-claim-generation-by-span","title":"Measuring Patent Claim Generation by Span Relevancy","date":"2019-08-26","arxiv_id":"1908.09591","n_code_links":0,"syntology":null},{"paper":"/paper/patient-knowledge-distillation-for-bert-model","slug":"patient-knowledge-distillation-for-bert-model","title":"Patient Knowledge Distillation for BERT Model Compression","date":"2019-08-25","arxiv_id":"1908.09355","n_code_links":5,"syntology":{"ran":19,"of":27,"n_ran_checked":13,"n_instrument":6,"unverified":8,"pointer_only":27,"phrase":"19 ran (of which 0 constructed an object rather than computing a result; 13 with no instrument failure: 2 honoured, 1 violated, 10 with no contract checked; 6 where Syntology's instrument failed) · 8 unverified","official":{"repos":["intersun/PKD-for-BERT-Model-Compression"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/bert-for-coreference-resolution-baselines-and","slug":"bert-for-coreference-resolution-baselines-and","title":"BERT for Coreference Resolution: Baselines and Analysis","date":"2019-08-24","arxiv_id":"1908.09091","n_code_links":2,"syntology":null},{"paper":"/paper/well-read-students-learn-better-the-impact-of","slug":"well-read-students-learn-better-the-impact-of","title":"Well-Read Students Learn Better: On the Importance of Pre-training Compact Models","date":"2019-08-23","arxiv_id":"1908.08962","n_code_links":40,"syntology":{"ran":21,"of":32,"n_ran_checked":16,"n_instrument":5,"unverified":11,"pointer_only":4,"phrase":"21 ran (of which 0 constructed an object rather than computing a result; 16 with no instrument failure: 1 honoured, 0 violated, 15 with no contract checked; 5 where Syntology's instrument failed) · 11 unverified","official":{"repos":["google-research/bert"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/multi-passage-bert-a-globally-normalized-bert","slug":"multi-passage-bert-a-globally-normalized-bert","title":"Multi-passage BERT: A Globally Normalized BERT Model for Open-domain Question Answering","date":"2019-08-22","arxiv_id":"1908.08167","n_code_links":0,"syntology":null},{"paper":"/paper/revisit-semantic-representation-and-tree","slug":"revisit-semantic-representation-and-tree","title":"Revisiting Semantic Representation and Tree Search for Similar Question Retrieval","date":"2019-08-22","arxiv_id":"1908.08326","n_code_links":1,"syntology":null},{"paper":"/paper/text-summarization-with-pretrained-encoders","slug":"text-summarization-with-pretrained-encoders","title":"Text Summarization with Pretrained Encoders","date":"2019-08-22","arxiv_id":"1908.08345","n_code_links":19,"syntology":{"ran":13,"of":21,"n_ran_checked":12,"n_instrument":1,"unverified":8,"pointer_only":5,"phrase":"13 ran (of which 0 constructed an object rather than computing a result; 12 with no instrument failure: 0 honoured, 1 violated, 11 with no contract checked; 1 where Syntology's instrument failed) · 8 unverified","official":{"repos":["nlpyang/PreSumm"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/vl-bert-pre-training-of-generic-visual","slug":"vl-bert-pre-training-of-generic-visual","title":"VL-BERT: Pre-training of Generic Visual-Linguistic Representations","date":"2019-08-22","arxiv_id":"1908.08530","n_code_links":3,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["jackroos/VL-BERT"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/190807721","slug":"190807721","title":"Fine-tuning BERT for Joint Entity and Relation Extraction in Chinese Medical Text","date":"2019-08-21","arxiv_id":"1908.07721","n_code_links":1,"syntology":null},{"paper":null,"slug":"revealing-the-dark-secrets-of-bert","title":"Revealing the Dark Secrets of BERT","date":"2019-08-21","arxiv_id":"1908.08593","n_code_links":0,"syntology":null},{"paper":"/paper/evaluating-contextualized-embeddings-on-54","slug":"evaluating-contextualized-embeddings-on-54","title":"Evaluating Contextualized Embeddings on 54 Languages in POS Tagging, Lemmatization and Dependency Parsing","date":"2019-08-20","arxiv_id":"1908.07448","n_code_links":0,"syntology":null},{"paper":"/paper/glossbert-bert-for-word-sense-disambiguation","slug":"glossbert-bert-for-word-sense-disambiguation","title":"GlossBERT: BERT for Word Sense Disambiguation with Gloss Knowledge","date":"2019-08-20","arxiv_id":"1908.07245","n_code_links":3,"syntology":null},{"paper":null,"slug":"a-study-of-bert-for-non-factoid-question","title":"A Study of BERT for Non-Factoid Question-Answering under Passage Length Constraints","date":"2019-08-19","arxiv_id":"1908.06780","n_code_links":0,"syntology":null},{"paper":"/paper/align-mask-and-select-a-simple-method-for","slug":"align-mask-and-select-a-simple-method-for","title":"Align, Mask and Select: A Simple Method for Incorporating Commonsense Knowledge into Language Representation Models","date":"2019-08-19","arxiv_id":"1908.06725","n_code_links":0,"syntology":null},{"paper":"/paper/neural-architectures-for-nested-ner-through-1","slug":"neural-architectures-for-nested-ner-through-1","title":"Neural Architectures for Nested NER through Linearization","date":"2019-08-19","arxiv_id":"1908.06926","n_code_links":1,"syntology":null},{"paper":"/paper/emotionx-idea-emotion-bert-an-affectional","slug":"emotionx-idea-emotion-bert-an-affectional","title":"EmotionX-IDEA: Emotion BERT -- an Affectional Model for Conversation","date":"2019-08-17","arxiv_id":"1908.06264","n_code_links":1,"syntology":null},{"paper":null,"slug":"language-features-matter-effective-language","title":"Language Features Matter: Effective Language Representations for Vision-Language Tasks","date":"2019-08-17","arxiv_id":"1908.06327","n_code_links":0,"syntology":null},{"paper":"/paper/bert-based-multi-head-selection-for-joint","slug":"bert-based-multi-head-selection-for-joint","title":"BERT-Based Multi-Head Selection for Joint Entity-Relation Extraction","date":"2019-08-16","arxiv_id":"1908.05908","n_code_links":1,"syntology":null},{"paper":null,"slug":"cfo-a-framework-for-building-production-nlp","title":"CFO: A Framework for Building Production NLP Systems","date":"2019-08-16","arxiv_id":"1908.06121","n_code_links":0,"syntology":null},{"paper":"/paper/clutrr-a-diagnostic-benchmark-for-inductive","slug":"clutrr-a-diagnostic-benchmark-for-inductive","title":"CLUTRR: A Diagnostic Benchmark for Inductive Reasoning from Text","date":"2019-08-16","arxiv_id":"1908.06177","n_code_links":5,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["facebookresearch/clutrr"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}}],"record_sha256":"73394a2ea122ed5ee9c37b4114c29c01cdb3544dca51d2bfd1c702e4c8a4be26","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}