{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/linear-warmup-with-linear-decay/papers/59","list_of":"/method/linear-warmup-with-linear-decay","method":"Linear Warmup With Linear Decay","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":59,"pages_in_order":71,"rows_per_page":100,"rows":[5801,5900],"of":7076,"counts":{"archive_papers_tagged":7076,"with_a_code_link":2913,"where_syntology_ran_a_sample":650,"not_listed_spam_title":0,"listed":7076,"listed_where_code_ran":650,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":531,"every_run_a_failure_of_syntologys_instrument":119,"listed_with_a_run_with_no_instrument_failure":531,"listed_every_run_a_failure_of_syntologys_instrument":119,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/linear-warmup-with-linear-decay","prev":"/method/linear-warmup-with-linear-decay/papers/58","next":"/method/linear-warmup-with-linear-decay/papers/60","papers":[{"paper":null,"slug":"a-token-wise-cnn-based-method-for-sentence","title":"A Token-wise CNN-based Method for Sentence Compression","date":"2020-09-23","arxiv_id":"2009.11260","n_code_links":0,"syntology":null},{"paper":null,"slug":"autorc-improving-bert-based-relation","title":"AutoRC: Improving BERT Based Relation Classification Models via Architecture Search","date":"2020-09-22","arxiv_id":"2009.10680","n_code_links":0,"syntology":null},{"paper":"/paper/constructing-interval-variables-via-faceted","slug":"constructing-interval-variables-via-faceted","title":"Constructing interval variables via faceted Rasch measurement and multitask deep learning: a hate speech application","date":"2020-09-22","arxiv_id":"2009.10277","n_code_links":2,"syntology":{"ran":3,"of":3,"n_ran_checked":3,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["ck37/coral-ordinal"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/grace-gradient-harmonized-and-cascaded","slug":"grace-gradient-harmonized-and-cascaded","title":"GRACE: Gradient Harmonized and Cascaded Labeling for Aspect-based Sentiment Analysis","date":"2020-09-22","arxiv_id":"2009.10557","n_code_links":1,"syntology":null},{"paper":null,"slug":"on-data-augmentation-for-extreme-multi-label","title":"On Data Augmentation for Extreme Multi-label Classification","date":"2020-09-22","arxiv_id":"2009.10778","n_code_links":0,"syntology":null},{"paper":"/paper/latin-bert-a-contextual-language-model-for","slug":"latin-bert-a-contextual-language-model-for","title":"Latin BERT: A Contextual Language Model for Classical Philology","date":"2020-09-21","arxiv_id":"2009.10053","n_code_links":1,"syntology":null},{"paper":"/paper/profile-consistency-identification-for-open","slug":"profile-consistency-identification-for-open","title":"Profile Consistency Identification for Open-domain Dialogue Agents","date":"2020-09-21","arxiv_id":"2009.09680","n_code_links":1,"syntology":{"ran":1,"of":3,"n_ran_checked":1,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"1 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified; the one sample that ran constructed an object rather than computing a result","official":{"repos":["songhaoyu/KvPI"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/ted-triple-supervision-decouples-end-to-end","slug":"ted-triple-supervision-decouples-end-to-end","title":"\"Listen, Understand and Translate\": Triple Supervision Decouples End-to-end Speech-to-text Translation","date":"2020-09-21","arxiv_id":"2009.09704","n_code_links":1,"syntology":null},{"paper":null,"slug":"when-they-say-weed-causes-depression-but-it-s","title":"\"When they say weed causes depression, but it's your fav antidepressant\": Knowledge-aware Attention Framework for Relationship Extraction","date":"2020-09-21","arxiv_id":"2009.10155","n_code_links":0,"syntology":null},{"paper":"/paper/dual-path-cnn-with-max-gated-block-for-text","slug":"dual-path-cnn-with-max-gated-block-for-text","title":"Dual-path CNN with Max Gated block for Text-Based Person Re-identification","date":"2020-09-20","arxiv_id":"2009.09343","n_code_links":1,"syntology":null},{"paper":"/paper/longformer-for-ms-marco-document-re-ranking","slug":"longformer-for-ms-marco-document-re-ranking","title":"Longformer for MS MARCO Document Re-ranking Task","date":"2020-09-20","arxiv_id":"2009.09392","n_code_links":1,"syntology":null},{"paper":"/paper/persian-ezafe-recognition-using-transformers","slug":"persian-ezafe-recognition-using-transformers","title":"Persian Ezafe Recognition Using Transformers and Its Role in Part-Of-Speech Tagging","date":"2020-09-20","arxiv_id":"2009.09474","n_code_links":1,"syntology":null},{"paper":null,"slug":"vicomtech-at-ehealth-kd-challenge-2020-deep","title":"Vicomtech at eHealth-KD Challenge 2020: Deep End-to-End Model for Entity and Relation Extraction in Medical Text","date":"2020-09-20","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"virtualflow-decoupling-deep-learning-model","title":"VirtualFlow: Decoupling Deep Learning Models from the Underlying Hardware","date":"2020-09-20","arxiv_id":"2009.09523","n_code_links":0,"syntology":null},{"paper":"/paper/conditionally-adaptive-multi-task-learning","slug":"conditionally-adaptive-multi-task-learning","title":"Conditionally Adaptive Multi-Task Learning: Improving Transfer Learning in NLP Using Fewer Parameters & Less Data","date":"2020-09-19","arxiv_id":"2009.09139","n_code_links":1,"syntology":null},{"paper":null,"slug":"nominal-compound-chain-extraction-a-new-task","title":"Nominal Compound Chain Extraction: A New Task for Semantic-enriched Lexical Chain","date":"2020-09-19","arxiv_id":"2009.09173","n_code_links":0,"syntology":null},{"paper":null,"slug":"prior-art-search-and-reranking-for-generated","title":"Prior Art Search and Reranking for Generated Patent Text","date":"2020-09-19","arxiv_id":"2009.09132","n_code_links":0,"syntology":null},{"paper":"/paper/fasthan-a-bert-based-joint-many-task-toolkit","slug":"fasthan-a-bert-based-joint-many-task-toolkit","title":"fastHan: A BERT-based Multi-Task Toolkit for Chinese NLP","date":"2020-09-18","arxiv_id":"2009.08633","n_code_links":1,"syntology":null},{"paper":null,"slug":"neu-at-wnut-2020-task-2-data-augmentation-to","title":"NEU at WNUT-2020 Task 2: Data Augmentation To Tell BERT That Death Is Not Necessarily Informative","date":"2020-09-18","arxiv_id":"2009.08590","n_code_links":0,"syntology":null},{"paper":"/paper/the-birth-of-romanian-bert","slug":"the-birth-of-romanian-bert","title":"The birth of Romanian BERT","date":"2020-09-18","arxiv_id":"2009.08712","n_code_links":1,"syntology":null},{"paper":"/paper/will-it-unblend","slug":"will-it-unblend","title":"Will it Unblend?","date":"2020-09-18","arxiv_id":"2009.09123","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-multimodal-memes-classification-a-survey","title":"A Multimodal Memes Classification: A Survey and Open Research Issues","date":"2020-09-17","arxiv_id":"2009.08395","n_code_links":0,"syntology":null},{"paper":null,"slug":"compositional-and-lexical-semantics-in","title":"Compositional and Lexical Semantics in RoBERTa, BERT and DistilBERT: A Case Study on CoQA","date":"2020-09-17","arxiv_id":"2009.08257","n_code_links":0,"syntology":null},{"paper":null,"slug":"cross-modal-alignment-with-mixture-experts","title":"Cross-Modal Alignment with Mixture Experts Neural Network for Intral-City Retail Recommendation","date":"2020-09-17","arxiv_id":"2009.09926","n_code_links":0,"syntology":null},{"paper":"/paper/dsc-iit-ism-at-semeval-2020-task-6-boosting","slug":"dsc-iit-ism-at-semeval-2020-task-6-boosting","title":"DSC IIT-ISM at SemEval-2020 Task 6: Boosting BERT with Dependencies for Definition Extraction","date":"2020-09-17","arxiv_id":"2009.08180","n_code_links":1,"syntology":null},{"paper":null,"slug":"efficient-transformer-based-large-scale","title":"Efficient Transformer-based Large Scale Language Representations using Hardware-friendly Block Structured Pruning","date":"2020-09-17","arxiv_id":"2009.08065","n_code_links":0,"syntology":null},{"paper":"/paper/multi-2oie-multilingual-open-information","slug":"multi-2oie-multilingual-open-information","title":"Multi$^2$OIE: Multilingual Open Information Extraction Based on Multi-Head Attention with BERT","date":"2020-09-17","arxiv_id":"2009.08128","n_code_links":1,"syntology":{"ran":6,"of":7,"n_ran_checked":2,"n_instrument":4,"unverified":1,"pointer_only":0,"phrase":"6 ran (of which 2 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 4 where Syntology's instrument failed) · 1 unverified","official":{"repos":["youngbin-ro/Multi2OIE"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":2,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"deep-learning-approaches-for-extracting","title":"Deep Learning Approaches for Extracting Adverse Events and Indications of Dietary Supplements from Clinical Text","date":"2020-09-16","arxiv_id":"2009.07780","n_code_links":0,"syntology":null},{"paper":"/paper/simplified-tinybert-knowledge-distillation","slug":"simplified-tinybert-knowledge-distillation","title":"Simplified TinyBERT: Knowledge Distillation for Document Retrieval","date":"2020-09-16","arxiv_id":"2009.07531","n_code_links":4,"syntology":null},{"paper":null,"slug":"solomon-at-semeval-2020-task-11-ensemble","title":"Solomon at SemEval-2020 Task 11: Ensemble Architecture for Fine-Tuned Propaganda Detection in News Articles","date":"2020-09-16","arxiv_id":"2009.07473","n_code_links":0,"syntology":null},{"paper":"/paper/union-an-unreferenced-metric-for-evaluating","slug":"union-an-unreferenced-metric-for-evaluating","title":"UNION: An Unreferenced Metric for Evaluating Open-ended Story Generation","date":"2020-09-16","arxiv_id":"2009.07602","n_code_links":1,"syntology":{"ran":0,"of":1,"n_ran_checked":0,"n_instrument":0,"unverified":1,"pointer_only":1,"phrase":"0 ran · 1 unverified","official":{"repos":["thu-coai/UNION"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":[]}}},{"paper":null,"slug":"achieving-real-time-execution-of-transformer","title":"Real-Time Execution of Large-scale Language Models on Mobile","date":"2020-09-15","arxiv_id":"2009.06823","n_code_links":0,"syntology":null},{"paper":null,"slug":"augmented-natural-language-for-generative","title":"Augmented Natural Language for Generative Sequence Labeling","date":"2020-09-15","arxiv_id":"2009.13272","n_code_links":0,"syntology":null},{"paper":"/paper/bert-qe-contextualized-query-expansion-for","slug":"bert-qe-contextualized-query-expansion-for","title":"BERT-QE: Contextualized Query Expansion for Document Re-ranking","date":"2020-09-15","arxiv_id":"2009.07258","n_code_links":1,"syntology":{"ran":2,"of":11,"n_ran_checked":0,"n_instrument":2,"unverified":9,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 9 unverified","official":{"repos":["zh-zheng/BERT-QE"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":9,"ran_from_kinds":["official"]}}},{"paper":"/paper/denert-kg-named-entity-and-relation","slug":"denert-kg-named-entity-and-relation","title":"DeNERT-KG: Named Entity and Relation Extraction Model Using DQN, Knowledge Graph, and BERT","date":"2020-09-15","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"event-presence-prediction-helps-trigger","title":"Event Presence Prediction Helps Trigger Detection Across Languages","date":"2020-09-15","arxiv_id":"2009.07188","n_code_links":0,"syntology":null},{"paper":null,"slug":"lessons-learned-from-applying-off-the-shelf","title":"Lessons Learned from Applying off-the-shelf BERT: There is no Silver Bullet","date":"2020-09-15","arxiv_id":"2009.07238","n_code_links":0,"syntology":null},{"paper":"/paper/mlmlm-link-prediction-with-mean-likelihood","slug":"mlmlm-link-prediction-with-mean-likelihood","title":"MLMLM: Link Prediction with Mean Likelihood Masked Language Model","date":"2020-09-15","arxiv_id":"2009.07058","n_code_links":0,"syntology":null},{"paper":null,"slug":"beyond-accuracy-roi-driven-data-analytics-of","title":"Beyond Accuracy: ROI-driven Data Analytics of Empirical Data","date":"2020-09-14","arxiv_id":"2009.06492","n_code_links":0,"syntology":null},{"paper":"/paper/can-fine-tuning-pre-trained-models-lead-to","slug":"can-fine-tuning-pre-trained-models-lead-to","title":"On Robustness and Bias Analysis of BERT-based Relation Extraction","date":"2020-09-14","arxiv_id":"2009.06206","n_code_links":1,"syntology":null},{"paper":null,"slug":"efficient-transformers-a-survey","title":"Efficient Transformers: A Survey","date":"2020-09-14","arxiv_id":"2009.06732","n_code_links":0,"syntology":null},{"paper":"/paper/filling-the-gap-of-utterance-aware-and","slug":"filling-the-gap-of-utterance-aware-and","title":"Filling the Gap of Utterance-aware and Speaker-aware Representation for Multi-turn Dialogue","date":"2020-09-14","arxiv_id":"2009.06504","n_code_links":1,"syntology":{"ran":2,"of":3,"n_ran_checked":2,"n_instrument":0,"unverified":1,"pointer_only":1,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":null}},{"paper":null,"slug":"boostingbert-integrating-multi-class-boosting","title":"BoostingBERT:Integrating Multi-Class Boosting into BERT for NLP Tasks","date":"2020-09-13","arxiv_id":"2009.05959","n_code_links":0,"syntology":null},{"paper":null,"slug":"cia-nitt-at-wnut-2020-task-2-classification","title":"CIA_NITT at WNUT-2020 Task 2: Classification of COVID-19 Tweets Using Pre-trained Language Models","date":"2020-09-12","arxiv_id":"2009.05782","n_code_links":0,"syntology":null},{"paper":"/paper/country-image-in-covid-19-pandemic-a-case","slug":"country-image-in-covid-19-pandemic-a-case","title":"Country Image in COVID-19 Pandemic: A Case Study of China","date":"2020-09-12","arxiv_id":"2009.05817","n_code_links":1,"syntology":null},{"paper":null,"slug":"fine-tuning-pre-trained-contextual-embeddings","title":"Fine-tuning Pre-trained Contextual Embeddings for Citation Content Analysis in Scholarly Publication","date":"2020-09-12","arxiv_id":"2009.05836","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-comparison-of-lstm-and-bert-for-small","title":"A Comparison of LSTM and BERT for Small Corpus","date":"2020-09-11","arxiv_id":"2009.05451","n_code_links":0,"syntology":null},{"paper":"/paper/compressed-deep-networks-goodbye-svd-hello","slug":"compressed-deep-networks-goodbye-svd-hello","title":"Compressed Deep Networks: Goodbye SVD, Hello Robust Low-Rank Approximation","date":"2020-09-11","arxiv_id":"2009.05647","n_code_links":1,"syntology":null},{"paper":null,"slug":"upb-at-semeval-2020-task-11-propaganda","title":"UPB at SemEval-2020 Task 11: Propaganda Detection with Domain-Specific Trained BERT","date":"2020-09-11","arxiv_id":"2009.05289","n_code_links":0,"syntology":null},{"paper":"/paper/upb-at-semeval-2020-task-6-pretrained","slug":"upb-at-semeval-2020-task-6-pretrained","title":"UPB at SemEval-2020 Task 6: Pretrained Language Models for Definition Extraction","date":"2020-09-11","arxiv_id":"2009.05603","n_code_links":3,"syntology":null},{"paper":"/paper/do-response-selection-models-really-know-what","slug":"do-response-selection-models-really-know-what","title":"Do Response Selection Models Really Know What's Next? Utterance Manipulation Strategies for Multi-turn Response Selection","date":"2020-09-10","arxiv_id":"2009.04703","n_code_links":1,"syntology":null},{"paper":null,"slug":"investigating-gender-bias-in-bert","title":"Investigating Gender Bias in BERT","date":"2020-09-10","arxiv_id":"2009.05021","n_code_links":0,"syntology":null},{"paper":"/paper/modern-methods-for-text-generation","slug":"modern-methods-for-text-generation","title":"Modern Methods for Text Generation","date":"2020-09-10","arxiv_id":"2009.04968","n_code_links":2,"syntology":null},{"paper":null,"slug":"comparative-study-of-language-models-on-cross","title":"Comparative Study of Language Models on Cross-Domain Data with Model Agnostic Explainability","date":"2020-09-09","arxiv_id":"2009.04095","n_code_links":0,"syntology":null},{"paper":"/paper/pay-attention-when-required","slug":"pay-attention-when-required","title":"Pay Attention when Required","date":"2020-09-09","arxiv_id":"2009.04534","n_code_links":2,"syntology":null},{"paper":null,"slug":"ernie-at-semeval-2020-task-10-learning-word","title":"ERNIE at SemEval-2020 Task 10: Learning Word Emphasis Selection by Pre-trained Language Model","date":"2020-09-08","arxiv_id":"2009.03706","n_code_links":0,"syntology":null},{"paper":null,"slug":"e-bert-a-phrase-and-product-knowledge","title":"E-BERT: A Phrase and Product Knowledge Enhanced Language Model for E-commerce","date":"2020-09-07","arxiv_id":"2009.02835","n_code_links":0,"syntology":null},{"paper":null,"slug":"edinburghnlp-at-wnut-2020-task-2-leveraging","title":"EdinburghNLP at WNUT-2020 Task 2: Leveraging Transformers with Generalized Augmentation for Identifying Informativeness in COVID-19 Tweets","date":"2020-09-06","arxiv_id":"2009.06375","n_code_links":0,"syntology":null},{"paper":null,"slug":"qiaoning-at-semeval-2020-task-4-commonsense","title":"QiaoNing at SemEval-2020 Task 4: Commonsense Validation and Explanation system based on ensemble of language model","date":"2020-09-06","arxiv_id":"2009.02645","n_code_links":0,"syntology":null},{"paper":null,"slug":"accenture-at-checkthat-2020-if-you-say-so","title":"Accenture at CheckThat! 2020: If you say so: Post-hoc fact-checking of claims using transformer-based models","date":"2020-09-05","arxiv_id":"2009.02431","n_code_links":0,"syntology":null},{"paper":"/paper/comparative-evaluation-of-pretrained-transfer","slug":"comparative-evaluation-of-pretrained-transfer","title":"Comparative Evaluation of Pretrained Transfer Learning Models on Automatic Short Answer Grading","date":"2020-09-02","arxiv_id":"2009.01303","n_code_links":1,"syntology":null},{"paper":"/paper/automatic-assignment-of-radiology-examination","slug":"automatic-assignment-of-radiology-examination","title":"Automatic Assignment of Radiology Examination Protocols Using Pre-trained Language Models with Knowledge Distillation","date":"2020-09-01","arxiv_id":"2009.00694","n_code_links":1,"syntology":null},{"paper":"/paper/sentimental-liar-extended-corpus-and-deep","slug":"sentimental-liar-extended-corpus-and-deep","title":"Sentimental LIAR: Extended Corpus and Deep Learning Models for Fake Claim Classification","date":"2020-09-01","arxiv_id":"2009.01047","n_code_links":2,"syntology":null},{"paper":"/paper/a-bidirectional-tree-tagging-scheme-for","slug":"a-bidirectional-tree-tagging-scheme-for","title":"A Bidirectional Tree Tagging Scheme for Joint Medical Relation Extraction","date":"2020-08-31","arxiv_id":"2008.13339","n_code_links":0,"syntology":null},{"paper":"/paper/soccogcom-at-semeval-2020-task-11","slug":"soccogcom-at-semeval-2020-task-11","title":"SocCogCom at SemEval-2020 Task 11: Characterizing and Detecting Propaganda using Sentence-Level Emotional Salience Features","date":"2020-08-29","arxiv_id":"2008.13012","n_code_links":1,"syntology":null},{"paper":null,"slug":"knowledge-efficient-deep-learning-for-natural","title":"Knowledge Efficient Deep Learning for Natural Language Processing","date":"2020-08-28","arxiv_id":"2008.12878","n_code_links":0,"syntology":null},{"paper":"/paper/rethinking-the-objectives-of-extractive","slug":"rethinking-the-objectives-of-extractive","title":"Rethinking the Objectives of Extractive Question Answering","date":"2020-08-28","arxiv_id":"2008.12804","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":3,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["KNOT-FIT-BUT/JointSpanExtraction"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/a-fast-and-robust-bert-based-dialogue-state","slug":"a-fast-and-robust-bert-based-dialogue-state","title":"A Fast and Robust BERT-based Dialogue State Tracker for Schema-Guided Dialogue Dataset","date":"2020-08-27","arxiv_id":"2008.12335","n_code_links":1,"syntology":null},{"paper":null,"slug":"ambert-a-pre-trained-language-model-with","title":"AMBERT: A Pre-trained Language Model with Multi-Grained Tokenization","date":"2020-08-27","arxiv_id":"2008.11869","n_code_links":0,"syntology":null},{"paper":"/paper/entity-and-evidence-guided-relation","slug":"entity-and-evidence-guided-relation","title":"Entity and Evidence Guided Relation Extraction for DocRED","date":"2020-08-27","arxiv_id":"2008.12283","n_code_links":0,"syntology":null},{"paper":"/paper/greek-bert-the-greeks-visiting-sesame-street","slug":"greek-bert-the-greeks-visiting-sesame-street","title":"GREEK-BERT: The Greeks visiting Sesame Street","date":"2020-08-27","arxiv_id":"2008.12014","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":2,"n_instrument":1,"unverified":0,"pointer_only":1,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["nlpaueb/greek-bert"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"multigbs-a-multi-layer-graph-approach-to","title":"MultiGBS: A multi-layer graph approach to biomedical summarization","date":"2020-08-27","arxiv_id":"2008.11908","n_code_links":0,"syntology":null},{"paper":"/paper/query-focused-multi-document-summarisation-of","slug":"query-focused-multi-document-summarisation-of","title":"Query Focused Multi-document Summarisation of Biomedical Texts","date":"2020-08-27","arxiv_id":"2008.11986","n_code_links":1,"syntology":null},{"paper":"/paper/query-focused-multi-document-summarisation-of-1","slug":"query-focused-multi-document-summarisation-of-1","title":"Query Focused Multi-document Summarisation of Biomedical Texts: Macquarie Universiy and the Australian National University at BioASQ8b","date":"2020-08-27","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"a-multitask-deep-learning-approach-for-user","title":"A Multitask Deep Learning Approach for User Depression Detection on Sina Weibo","date":"2020-08-26","arxiv_id":"2008.11708","n_code_links":0,"syntology":null},{"paper":null,"slug":"apmsqueeze-a-communication-efficient-adam","title":"APMSqueeze: A Communication Efficient Adam-Preconditioned Momentum SGD Algorithm","date":"2020-08-26","arxiv_id":"2008.11343","n_code_links":0,"syntology":null},{"paper":"/paper/language-models-and-word-sense-disambiguation","slug":"language-models-and-word-sense-disambiguation","title":"Analysis and Evaluation of Language Models for Word Sense Disambiguation","date":"2020-08-26","arxiv_id":"2008.11608","n_code_links":1,"syntology":null},{"paper":null,"slug":"conceptualized-representation-learning-for","title":"Conceptualized Representation Learning for Chinese Biomedical Text Mining","date":"2020-08-25","arxiv_id":"2008.10813","n_code_links":0,"syntology":null},{"paper":"/paper/etc-nlg-end-to-end-topic-conditioned-natural","slug":"etc-nlg-end-to-end-topic-conditioned-natural","title":"ETC-NLG: End-to-end Topic-Conditioned Natural Language Generation","date":"2020-08-25","arxiv_id":"2008.10875","n_code_links":1,"syntology":null},{"paper":"/paper/knowledge-empowered-representation-learning","slug":"knowledge-empowered-representation-learning","title":"Knowledge-Empowered Representation Learning for Chinese Medical Reading Comprehension: Task, Model and Resources","date":"2020-08-24","arxiv_id":"2008.10327","n_code_links":1,"syntology":null},{"paper":null,"slug":"prediction-of-icd-codes-with-clinical-bert","title":"Prediction of ICD Codes with Clinical BERT Embeddings and Text Augmentation with Label Balancing using MIMIC-III","date":"2020-08-24","arxiv_id":"2008.10492","n_code_links":0,"syntology":null},{"paper":null,"slug":"syrapropa-at-semeval-2020-task-11-bert-based","title":"syrapropa at SemEval-2020 Task 11: BERT-based Models Design For Propagandistic Technique and Span Detection","date":"2020-08-24","arxiv_id":"2008.10163","n_code_links":0,"syntology":null},{"paper":null,"slug":"two-stages-approach-for-tweet-engagement","title":"Two Stages Approach for Tweet Engagement Prediction","date":"2020-08-24","arxiv_id":"2008.10419","n_code_links":0,"syntology":null},{"paper":"/paper/ynu-hpcc-at-semeval-2020-task-11-lstm-network","slug":"ynu-hpcc-at-semeval-2020-task-11-lstm-network","title":"YNU-HPCC at SemEval-2020 Task 11: LSTM Network for Detection of Propaganda Techniques in News Articles","date":"2020-08-24","arxiv_id":"2008.10166","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":2,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["daojiaxu/semeval_11"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"applications-of-bert-based-sequence-tagging","title":"Applications of BERT Based Sequence Tagging Models on Chinese Medical Text Attributes Extraction","date":"2020-08-22","arxiv_id":"2008.09740","n_code_links":0,"syntology":null},{"paper":"/paper/cyberwalle-at-semeval-2020-task-11-an","slug":"cyberwalle-at-semeval-2020-task-11-an","title":"CyberWallE at SemEval-2020 Task 11: An Analysis of Feature Engineering for Ensemble Models for Propaganda Detection","date":"2020-08-22","arxiv_id":"2008.09859","n_code_links":1,"syntology":null},{"paper":"/paper/duth-at-semeval-2020-task-11-bert-with-entity","slug":"duth-at-semeval-2020-task-11-bert-with-entity","title":"DUTH at SemEval-2020 Task 11: BERT with Entity Mapping for Propaganda Classification","date":"2020-08-22","arxiv_id":"2008.09894","n_code_links":1,"syntology":null},{"paper":"/paper/fat-albert-finding-answers-in-large-texts","slug":"fat-albert-finding-answers-in-large-texts","title":"FAT ALBERT: Finding Answers in Large Texts using Semantic Similarity Attention Layer based on BERT","date":"2020-08-22","arxiv_id":"2009.01004","n_code_links":1,"syntology":null},{"paper":"/paper/hinglishnlp-fine-tuned-language-models-for","slug":"hinglishnlp-fine-tuned-language-models-for","title":"HinglishNLP: Fine-tuned Language Models for Hinglish Sentiment Detection","date":"2020-08-22","arxiv_id":"2008.09820","n_code_links":2,"syntology":null},{"paper":"/paper/abstractive-summarization-of-spoken","slug":"abstractive-summarization-of-spoken","title":"Abstractive Summarization of Spoken andWritten Instructions with BERT","date":"2020-08-21","arxiv_id":null,"n_code_links":2,"syntology":null},{"paper":null,"slug":"adapting-event-extractors-to-medical-data","title":"Adapting Event Extractors to Medical Data: Bridging the Covariate Shift","date":"2020-08-21","arxiv_id":"2008.09266","n_code_links":0,"syntology":null},{"paper":null,"slug":"an-experimental-study-of-deep-neural-network","title":"An Experimental Study of Deep Neural Network Models for Vietnamese Multiple-Choice Reading Comprehension","date":"2020-08-20","arxiv_id":"2008.08810","n_code_links":0,"syntology":null},{"paper":"/paper/parade-passage-representation-aggregation-for","slug":"parade-passage-representation-aggregation-for","title":"PARADE: Passage Representation Aggregation for Document Reranking","date":"2020-08-20","arxiv_id":"2008.09093","n_code_links":1,"syntology":{"ran":0,"of":3,"n_ran_checked":0,"n_instrument":0,"unverified":3,"pointer_only":0,"phrase":"0 ran · 3 unverified","official":{"repos":["canjiali/PARADE"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":3,"ran_from_kinds":[]}}},{"paper":"/paper/top2vec-distributed-representations-of-topics","slug":"top2vec-distributed-representations-of-topics","title":"Top2Vec: Distributed Representations of Topics","date":"2020-08-19","arxiv_id":"2008.09470","n_code_links":2,"syntology":{"ran":2,"of":2,"n_ran_checked":0,"n_instrument":2,"unverified":0,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":{"repos":["ddangelov/Top2Vec"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"uob-at-semeval-2020-task-12-boosting-bert","title":"UoB at SemEval-2020 Task 12: Boosting BERT with Corpus Level Information","date":"2020-08-19","arxiv_id":"2008.08547","n_code_links":0,"syntology":null},{"paper":null,"slug":"ranking-clarification-questions-via-natural","title":"Ranking Clarification Questions via Natural Language Inference","date":"2020-08-18","arxiv_id":"2008.07688","n_code_links":0,"syntology":null},{"paper":null,"slug":"stock-index-prediction-with-multi-task","title":"Stock Index Prediction with Multi-task Learning and Word Polarity Over Time","date":"2020-08-17","arxiv_id":"2008.07605","n_code_links":0,"syntology":null},{"paper":"/paper/devlbert-learning-deconfounded-visio","slug":"devlbert-learning-deconfounded-visio","title":"DeVLBert: Learning Deconfounded Visio-Linguistic Representations","date":"2020-08-16","arxiv_id":"2008.06884","n_code_links":1,"syntology":null},{"paper":null,"slug":"finding-fast-transformers-one-shot-neural","title":"Finding Fast Transformers: One-Shot Neural Architecture Search by Component Composition","date":"2020-08-15","arxiv_id":"2008.06808","n_code_links":0,"syntology":null},{"paper":"/paper/jointly-fine-tuning-bert-like-self-supervised","slug":"jointly-fine-tuning-bert-like-self-supervised","title":"Jointly Fine-Tuning \"BERT-like\" Self Supervised Models to Improve Multimodal Speech Emotion Recognition","date":"2020-08-15","arxiv_id":"2008.06682","n_code_links":1,"syntology":null}],"record_sha256":"7253aa2f733a7a4625ecfc89c0e42b4db536be4a5ee072cdd5078cfe3ecb433d","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}