{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/dropout/papers/221","list_of":"/method/dropout","method":"Dropout","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":221,"pages_in_order":275,"rows_per_page":100,"rows":[22001,22100],"of":27472,"counts":{"archive_papers_tagged":27472,"with_a_code_link":12129,"where_syntology_ran_a_sample":3620,"not_listed_spam_title":0,"listed":27472,"listed_where_code_ran":3620,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":3044,"every_run_a_failure_of_syntologys_instrument":576,"listed_with_a_run_with_no_instrument_failure":3044,"listed_every_run_a_failure_of_syntologys_instrument":576,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/dropout","prev":"/method/dropout/papers/220","next":"/method/dropout/papers/222","papers":[{"paper":null,"slug":"listen-read-and-identify-multimodal-singing","title":"Listen, Read, and Identify: Multimodal Singing Language Identification of Music","date":"2021-03-02","arxiv_id":"2103.01893","n_code_links":0,"syntology":null},{"paper":null,"slug":"probing-product-description-generation-via","title":"Probing Product Description Generation via Posterior Distillation","date":"2021-03-02","arxiv_id":"2103.01594","n_code_links":0,"syntology":null},{"paper":null,"slug":"bert-based-knowledge-extraction-method-of","title":"BERT-based knowledge extraction method of unstructured domain text","date":"2021-03-01","arxiv_id":"2103.00728","n_code_links":0,"syntology":null},{"paper":null,"slug":"bert-based-patent-novelty-search-by-training","title":"BERT based patent novelty search by training claims to their own description","date":"2021-03-01","arxiv_id":"2103.01126","n_code_links":0,"syntology":null},{"paper":null,"slug":"brain-programming-is-immune-to-adversarial","title":"Brain Programming is Immune to Adversarial Attacks: Towards Accurate and Robust Image Classification using Symbolic Learning","date":"2021-03-01","arxiv_id":"2103.01359","n_code_links":0,"syntology":null},{"paper":null,"slug":"combat-covid-19-infodemic-using-explainable","title":"Combat COVID-19 Infodemic Using Explainable Natural Language Processing Models","date":"2021-03-01","arxiv_id":"2103.00747","n_code_links":0,"syntology":null},{"paper":null,"slug":"crossmap-transformer-a-crossmodal-masked-path","title":"CrossMap Transformer: A Crossmodal Masked Path Transformer Using Double Back-Translation for Vision-and-Language Navigation","date":"2021-03-01","arxiv_id":"2103.00852","n_code_links":0,"syntology":null},{"paper":null,"slug":"localdrop-a-hybrid-regularization-for-deep","title":"LocalDrop: A Hybrid Regularization for Deep Neural Networks","date":"2021-03-01","arxiv_id":"2103.00719","n_code_links":0,"syntology":null},{"paper":null,"slug":"long-document-summarization-in-a-low-resource","title":"Long Document Summarization in a Low Resource Setting using Pretrained Language Models","date":"2021-03-01","arxiv_id":"2103.00751","n_code_links":0,"syntology":null},{"paper":null,"slug":"nlp-cuet-dravidianlangtech-eacl2021","title":"NLP-CUET@DravidianLangTech-EACL2021: Investigating Visual and Textual Features to Identify Trolls from Multimodal Social Media Memes","date":"2021-02-28","arxiv_id":"2103.00466","n_code_links":0,"syntology":null},{"paper":"/paper/nlp-cuet-dravidianlangtech-eacl2021-offensive","slug":"nlp-cuet-dravidianlangtech-eacl2021-offensive","title":"NLP-CUET@DravidianLangTech-EACL2021: Offensive Language Detection from Multilingual Code-Mixed Text using Transformers","date":"2021-02-28","arxiv_id":"2103.00455","n_code_links":1,"syntology":null},{"paper":"/paper/nlp-cuet-lt-edi-eacl2021-multilingual-code","slug":"nlp-cuet-lt-edi-eacl2021-multilingual-code","title":"NLP-CUET@LT-EDI-EACL2021: Multilingual Code-Mixed Hope Speech Detection using Cross-lingual Representation Learner","date":"2021-02-28","arxiv_id":"2103.00464","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-novel-adaptive-deep-network-for-building","title":"A Novel Adaptive Deep Network for Building Footprint Segmentation","date":"2021-02-27","arxiv_id":"2103.00286","n_code_links":0,"syntology":null},{"paper":"/paper/covid-19-tweets-analysis-through-transformer","slug":"covid-19-tweets-analysis-through-transformer","title":"COVID-19 Tweets Analysis through Transformer Language Models","date":"2021-02-27","arxiv_id":"2103.00199","n_code_links":1,"syntology":null},{"paper":"/paper/generative-chemical-transformer-attention","slug":"generative-chemical-transformer-attention","title":"Generative Chemical Transformer: Neural Machine Learning of Molecular Geometric Structures from Chemical Language via Attention","date":"2021-02-27","arxiv_id":"2103.00213","n_code_links":2,"syntology":null},{"paper":"/paper/transformer-in-transformer","slug":"transformer-in-transformer","title":"Transformer in Transformer","date":"2021-02-27","arxiv_id":"2103.00112","n_code_links":12,"syntology":{"ran":16,"of":24,"n_ran_checked":15,"n_instrument":1,"unverified":8,"pointer_only":5,"phrase":"16 ran (of which 12 constructed an object rather than computing a result; 15 with no instrument failure: 1 honoured, 1 violated, 13 with no contract checked; 1 where Syntology's instrument failed) · 8 unverified","official":{"repos":["huawei-noah/CV-backbones"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":4,"n_ran_no_instrument_failure":4,"n_unverified":1,"ran_from_kinds":["listed","official"]}}},{"paper":null,"slug":"transformers-with-competitive-ensembles-of-1","title":"Transformers with Competitive Ensembles of Independent Mechanisms","date":"2021-02-27","arxiv_id":"2103.00336","n_code_links":0,"syntology":null},{"paper":"/paper/fjord-fair-and-accurate-federated-learning","slug":"fjord-fair-and-accurate-federated-learning","title":"FjORD: Fair and Accurate Federated Learning under heterogeneous targets with Ordered Dropout","date":"2021-02-26","arxiv_id":"2102.13451","n_code_links":2,"syntology":{"ran":2,"of":3,"n_ran_checked":2,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"2 ran (of which 2 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified; every one of the 2 samples that ran constructed an object rather than computing a result","official":{"repos":["adap/flower"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":null,"slug":"multi-task-transfer-learning-for-finding","title":"Multi-task transfer learning for finding actionable information from crisis-related messages on social media","date":"2021-02-26","arxiv_id":"2102.13395","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-framework-for-pruning-deep-neural-networks","title":"A Framework For Pruning Deep Neural Networks Using Energy-Based Models","date":"2021-02-25","arxiv_id":"2102.13188","n_code_links":0,"syntology":null},{"paper":null,"slug":"bert-based-acronym-disambiguation-with","title":"BERT-based Acronym Disambiguation with Multiple Training Strategies","date":"2021-02-25","arxiv_id":"2103.00488","n_code_links":0,"syntology":null},{"paper":null,"slug":"emotion-aware-emotion-agnostic-or-automatic","title":"Emotion-Aware, Emotion-Agnostic, or Automatic: Corpus Creation Strategies to Obtain Cognitive Event Appraisal Annotations","date":"2021-02-25","arxiv_id":"2102.12858","n_code_links":0,"syntology":null},{"paper":null,"slug":"lazyformer-self-attention-with-lazy-update","title":"LazyFormer: Self Attention with Lazy Update","date":"2021-02-25","arxiv_id":"2102.12702","n_code_links":0,"syntology":null},{"paper":"/paper/let-linguistic-knowledge-enhanced-graph","slug":"let-linguistic-knowledge-enhanced-graph","title":"LET: Linguistic Knowledge Enhanced Graph Transformer for Chinese Short Text Matching","date":"2021-02-25","arxiv_id":"2102.12671","n_code_links":1,"syntology":null},{"paper":null,"slug":"mixspeech-data-augmentation-for-low-resource","title":"MixSpeech: Data Augmentation for Low-resource Automatic Speech Recognition","date":"2021-02-25","arxiv_id":"2102.12664","n_code_links":0,"syntology":null},{"paper":null,"slug":"pharmke-knowledge-extraction-platform-for","title":"PharmKE: Knowledge Extraction Platform for Pharmaceutical Texts using Transfer Learning","date":"2021-02-25","arxiv_id":"2102.13139","n_code_links":0,"syntology":null},{"paper":null,"slug":"robust-pollen-imagery-classification-with","title":"Robust Pollen Imagery Classification with Generative Modeling and Mixup Training","date":"2021-02-25","arxiv_id":"2102.13143","n_code_links":0,"syntology":null},{"paper":"/paper/sentiment-analysis-of-persian-english-code","slug":"sentiment-analysis-of-persian-english-code","title":"Sentiment Analysis of Persian-English Code-mixed Texts","date":"2021-02-25","arxiv_id":"2102.12700","n_code_links":1,"syntology":null},{"paper":null,"slug":"from-universal-language-model-to-downstream","title":"From Universal Language Model to Downstream Task: Improving RoBERTa-Based Vietnamese Hate Speech Detection","date":"2021-02-24","arxiv_id":"2102.12162","n_code_links":0,"syntology":null},{"paper":null,"slug":"hopeful-men-lt-edi-eacl2021-hope-speech","title":"Hopeful_Men@LT-EDI-EACL2021: Hope Speech Detection Using Indic Transliteration and Transformers","date":"2021-02-24","arxiv_id":"2102.12082","n_code_links":0,"syntology":null},{"paper":"/paper/lrg-at-semeval-2021-task-4-improving-reading","slug":"lrg-at-semeval-2021-task-4-improving-reading","title":"LRG at SemEval-2021 Task 4: Improving Reading Comprehension with Abstract Words using Augmentation, Linguistic Features and Voting","date":"2021-02-24","arxiv_id":"2102.12255","n_code_links":1,"syntology":null},{"paper":"/paper/nlrg-at-semeval-2021-task-5-toxic-spans","slug":"nlrg-at-semeval-2021-task-5-toxic-spans","title":"NLRG at SemEval-2021 Task 5: Toxic Spans Detection Leveraging BERT-based Token Classification and Span Prediction Techniques","date":"2021-02-24","arxiv_id":"2102.12254","n_code_links":1,"syntology":null},{"paper":"/paper/pada-a-prompt-based-autoregressive-approach","slug":"pada-a-prompt-based-autoregressive-approach","title":"PADA: Example-based Prompt Learning for on-the-fly Adaptation to Unseen Domains","date":"2021-02-24","arxiv_id":"2102.12206","n_code_links":1,"syntology":null},{"paper":"/paper/pyramid-vision-transformer-a-versatile","slug":"pyramid-vision-transformer-a-versatile","title":"Pyramid Vision Transformer: A Versatile Backbone for Dense Prediction without Convolutions","date":"2021-02-24","arxiv_id":"2102.12122","n_code_links":11,"syntology":{"ran":22,"of":30,"n_ran_checked":18,"n_instrument":4,"unverified":8,"pointer_only":1,"phrase":"22 ran (of which 16 constructed an object rather than computing a result; 18 with no instrument failure: 2 honoured, 0 violated, 16 with no contract checked; 4 where Syntology's instrument failed) · 8 unverified","official":{"repos":["whai362/PVT"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":4,"n_ran_no_instrument_failure":4,"n_unverified":0,"ran_from_kinds":["listed","official"]}}},{"paper":null,"slug":"task-specific-pre-training-and-cross-lingual","title":"Task-Specific Pre-Training and Cross Lingual Transfer for Code-Switched Data","date":"2021-02-24","arxiv_id":"2102.12407","n_code_links":0,"syntology":null},{"paper":"/paper/when-attention-meets-fast-recurrence-training","slug":"when-attention-meets-fast-recurrence-training","title":"When Attention Meets Fast Recurrence: Training Language Models with Reduced Compute","date":"2021-02-24","arxiv_id":"2102.12459","n_code_links":1,"syntology":null},{"paper":"/paper/accurate-learning-of-graph-representations-1","slug":"accurate-learning-of-graph-representations-1","title":"Accurate Learning of Graph Representations with Graph Multiset Pooling","date":"2021-02-23","arxiv_id":"2102.11533","n_code_links":1,"syntology":{"ran":4,"of":5,"n_ran_checked":4,"n_instrument":0,"unverified":1,"pointer_only":5,"phrase":"4 ran (of which 4 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified; every one of the 4 samples that ran constructed an object rather than computing a result","official":{"repos":["JinheonBaek/GMT"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":4,"n_ran_no_instrument_failure":4,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"automatic-ship-classification-utilizing-bag","title":"Automatic Ship Classification Utilizing Bag of Deep Features","date":"2021-02-23","arxiv_id":"2102.11520","n_code_links":0,"syntology":null},{"paper":null,"slug":"deep-deformation-detail-synthesis-for-thin","title":"Deep Deformation Detail Synthesis for Thin Shell Models","date":"2021-02-23","arxiv_id":"2102.11541","n_code_links":0,"syntology":null},{"paper":"/paper/do-transformer-modifications-transfer-across","slug":"do-transformer-modifications-transfer-across","title":"Do Transformer Modifications Transfer Across Implementations and Applications?","date":"2021-02-23","arxiv_id":"2102.11972","n_code_links":1,"syntology":null},{"paper":null,"slug":"minimally-supervised-structure-rich-text","title":"Minimally-Supervised Structure-Rich Text Categorization via Learning on Text-Rich Networks","date":"2021-02-23","arxiv_id":"2102.11479","n_code_links":0,"syntology":null},{"paper":null,"slug":"robust-and-transferable-anomaly-detection-in","title":"Robust and Transferable Anomaly Detection in Log Data using Pre-Trained Language Models","date":"2021-02-23","arxiv_id":"2102.11570","n_code_links":0,"syntology":null},{"paper":"/paper/visualchexbert-addressing-the-discrepancy","slug":"visualchexbert-addressing-the-discrepancy","title":"VisualCheXbert: Addressing the Discrepancy Between Radiology Report Labels and Image Labels","date":"2021-02-23","arxiv_id":"2102.11467","n_code_links":1,"syntology":{"ran":0,"of":2,"n_ran_checked":0,"n_instrument":0,"unverified":2,"pointer_only":2,"phrase":"0 ran · 2 unverified","official":{"repos":["stanfordmlgroup/VisualCheXbert"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":2,"ran_from_kinds":[]}}},{"paper":null,"slug":"wavelet-transform-analytics-for-rf-based-uav","title":"Wavelet Transform Analytics for RF-Based UAV Detection and Identification System Using Machine Learning","date":"2021-02-23","arxiv_id":"2102.11894","n_code_links":0,"syntology":null},{"paper":"/paper/deepfake-video-detection-using-convolutional","slug":"deepfake-video-detection-using-convolutional","title":"Deepfake Video Detection Using Convolutional Vision Transformer","date":"2021-02-22","arxiv_id":"2102.11126","n_code_links":1,"syntology":null},{"paper":null,"slug":"determination-of-fault-location-in","title":"Determination of Fault Location in Transmission Lines with Image Processing and Artificial Neural Networks","date":"2021-02-22","arxiv_id":"2102.11073","n_code_links":0,"syntology":null},{"paper":"/paper/do-we-really-need-explicit-position-encodings","slug":"do-we-really-need-explicit-position-encodings","title":"Conditional Positional Encodings for Vision Transformers","date":"2021-02-22","arxiv_id":"2102.10882","n_code_links":2,"syntology":null},{"paper":"/paper/evaluating-contextualized-language-models-for","slug":"evaluating-contextualized-language-models-for","title":"Evaluating Contextualized Language Models for Hungarian","date":"2021-02-22","arxiv_id":"2102.10848","n_code_links":1,"syntology":null},{"paper":null,"slug":"few-shot-learning-for-information","title":"Few Shot Learning for Information Verification","date":"2021-02-22","arxiv_id":"2102.10956","n_code_links":0,"syntology":null},{"paper":null,"slug":"generating-human-readable-transcript-for","title":"Generating Human Readable Transcript for Automatic Speech Recognition with Pre-trained Language Model","date":"2021-02-22","arxiv_id":"2102.11114","n_code_links":0,"syntology":null},{"paper":"/paper/lightweight-combinational-machine-learning","slug":"lightweight-combinational-machine-learning","title":"Lightweight Combinational Machine Learning Algorithm for Sorting Canine Torso Radiographs","date":"2021-02-22","arxiv_id":"2102.11385","n_code_links":1,"syntology":null},{"paper":"/paper/lvcnet-efficient-condition-dependent-modeling","slug":"lvcnet-efficient-condition-dependent-modeling","title":"LVCNet: Efficient Condition-Dependent Modeling Network for Waveform Generation","date":"2021-02-22","arxiv_id":"2102.10815","n_code_links":5,"syntology":null},{"paper":null,"slug":"mixup-training-leads-to-reduced-overfitting","title":"MixUp Training Leads to Reduced Overfitting and Improved Calibration for the Transformer Architecture","date":"2021-02-22","arxiv_id":"2102.11402","n_code_links":0,"syntology":null},{"paper":"/paper/parallelizing-legendre-memory-unit-training","slug":"parallelizing-legendre-memory-unit-training","title":"Parallelizing Legendre Memory Unit Training","date":"2021-02-22","arxiv_id":"2102.11417","n_code_links":2,"syntology":null},{"paper":null,"slug":"position-information-in-transformers-an","title":"Position Information in Transformers: An Overview","date":"2021-02-22","arxiv_id":"2102.11090","n_code_links":0,"syntology":null},{"paper":null,"slug":"rconet-deformable-mutual-information","title":"RCoNet: Deformable Mutual Information Maximization and High-order Uncertainty-aware Learning for Robust COVID-19 Detection","date":"2021-02-22","arxiv_id":"2102.11099","n_code_links":0,"syntology":null},{"paper":null,"slug":"rubert-a-bilingual-roman-urdu-bert-using","title":"RUBERT: A Bilingual Roman Urdu BERT Using Cross Lingual Transfer Learning","date":"2021-02-22","arxiv_id":"2102.11278","n_code_links":0,"syntology":null},{"paper":"/paper/transformer-is-all-you-need-multimodal","slug":"transformer-is-all-you-need-multimodal","title":"UniT: Multimodal Multitask Learning with a Unified Transformer","date":"2021-02-22","arxiv_id":"2102.10772","n_code_links":1,"syntology":null},{"paper":"/paper/using-prior-knowledge-to-guide-bert-s","slug":"using-prior-knowledge-to-guide-bert-s","title":"Using Prior Knowledge to Guide BERT's Attention in Semantic Textual Matching Tasks","date":"2021-02-22","arxiv_id":"2102.10934","n_code_links":1,"syntology":null},{"paper":"/paper/medical-transformer-gated-axial-attention-for","slug":"medical-transformer-gated-axial-attention-for","title":"Medical Transformer: Gated Axial-Attention for Medical Image Segmentation","date":"2021-02-21","arxiv_id":"2102.10662","n_code_links":2,"syntology":null},{"paper":null,"slug":"pre-training-bert-on-arabic-tweets-practical","title":"Pre-Training BERT on Arabic Tweets: Practical Considerations","date":"2021-02-21","arxiv_id":"2102.10684","n_code_links":0,"syntology":null},{"paper":null,"slug":"web-based-application-for-detecting","title":"Web-based Application for Detecting Indonesian Clickbait Headlines using IndoBERT","date":"2021-02-21","arxiv_id":"2102.10601","n_code_links":0,"syntology":null},{"paper":null,"slug":"multilingual-answer-sentence-reranking-via","title":"Multilingual Answer Sentence Reranking via Automatically Translated Data","date":"2021-02-20","arxiv_id":"2102.10250","n_code_links":0,"syntology":null},{"paper":"/paper/towards-accurate-and-compact-architectures","slug":"towards-accurate-and-compact-architectures","title":"Towards Accurate and Compact Architectures via Neural Architecture Transformer","date":"2021-02-20","arxiv_id":"2102.10301","n_code_links":2,"syntology":null},{"paper":null,"slug":"unsupervised-medical-image-alignment-with","title":"Unsupervised Medical Image Alignment with Curriculum Learning","date":"2021-02-20","arxiv_id":"2102.10438","n_code_links":0,"syntology":null},{"paper":"/paper/calibrate-before-use-improving-few-shot","slug":"calibrate-before-use-improving-few-shot","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models","date":"2021-02-19","arxiv_id":"2102.09690","n_code_links":5,"syntology":{"ran":0,"of":4,"n_ran_checked":0,"n_instrument":0,"unverified":4,"pointer_only":0,"phrase":"0 ran · 4 unverified","official":{"repos":["tonyzhaozh/few-shot-learning"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":4,"ran_from_kinds":[]}}},{"paper":null,"slug":"dialect-identification-in-nuanced-arabic","title":"Dialect Identification in Nuanced Arabic Tweets Using Farasa Segmentation and AraBERT","date":"2021-02-19","arxiv_id":"2102.09749","n_code_links":0,"syntology":null},{"paper":"/paper/latent-variable-nested-set-transformers","slug":"latent-variable-nested-set-transformers","title":"Latent Variable Sequential Set Transformers For Joint Multi-Agent Motion Prediction","date":"2021-02-19","arxiv_id":"2104.00563","n_code_links":2,"syntology":{"ran":11,"of":17,"n_ran_checked":10,"n_instrument":1,"unverified":6,"pointer_only":2,"phrase":"11 ran (of which 0 constructed an object rather than computing a result; 10 with no instrument failure: 0 honoured, 0 violated, 10 with no contract checked; 1 where Syntology's instrument failed) · 6 unverified","official":{"repos":["roggirg/AutoBots"],"state":"official (archive's flag): 11 ran","n_ran":11,"n_constructed":0,"n_ran_no_instrument_failure":10,"n_unverified":6,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"learning-dynamic-bert-via-trainable-gate","title":"Learning Dynamic BERT via Trainable Gate Variables and a Bi-modal Regularizer","date":"2021-02-19","arxiv_id":"2102.09727","n_code_links":0,"syntology":null},{"paper":"/paper/towards-emotion-recognition-in-hindi-english","slug":"towards-emotion-recognition-in-hindi-english","title":"Towards Emotion Recognition in Hindi-English Code-Mixed Data: A Transformer Based Approach","date":"2021-02-19","arxiv_id":"2102.09943","n_code_links":1,"syntology":null},{"paper":"/paper/using-transformer-based-ensemble-learning-to","slug":"using-transformer-based-ensemble-learning-to","title":"Using Transformer based Ensemble Learning to classify Scientific Articles","date":"2021-02-19","arxiv_id":"2102.09991","n_code_links":2,"syntology":null},{"paper":"/paper/analysis-of-contextual-and-non-contextual","slug":"analysis-of-contextual-and-non-contextual","title":"Analysis Of Contextual and Non-Contextual Word Embedding Models For Hindi NER With Web Application For Data Collection","date":"2021-02-18","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/densely-nested-top-down-flows-for-salient","slug":"densely-nested-top-down-flows-for-salient","title":"Densely Nested Top-Down Flows for Salient Object Detection","date":"2021-02-18","arxiv_id":"2102.09133","n_code_links":1,"syntology":null},{"paper":"/paper/going-full-tilt-boogie-on-document","slug":"going-full-tilt-boogie-on-document","title":"Going Full-TILT Boogie on Document Understanding with Text-Image-Layout Transformer","date":"2021-02-18","arxiv_id":"2102.09550","n_code_links":1,"syntology":null},{"paper":"/paper/on-connectivity-of-solutions-in-deep-learning","slug":"on-connectivity-of-solutions-in-deep-learning","title":"When Are Solutions Connected in Deep Networks?","date":"2021-02-18","arxiv_id":"2102.09671","n_code_links":1,"syntology":null},{"paper":"/paper/quiz-style-question-generation-for-news","slug":"quiz-style-question-generation-for-news","title":"Quiz-Style Question Generation for News Stories","date":"2021-02-18","arxiv_id":"2102.09094","n_code_links":2,"syntology":null},{"paper":"/paper/training-microsoft-news-recommenders-with","slug":"training-microsoft-news-recommenders-with","title":"Training Large-Scale News Recommenders with Pretrained Language Models in the Loop","date":"2021-02-18","arxiv_id":"2102.09268","n_code_links":1,"syntology":null},{"paper":null,"slug":"unibuckernel-geolocating-swiss-german-jodels","title":"UnibucKernel: Geolocating Swiss German Jodels Using Ensemble Learning","date":"2021-02-18","arxiv_id":"2102.09379","n_code_links":0,"syntology":null},{"paper":"/paper/verifying-probabilistic-specifications-with","slug":"verifying-probabilistic-specifications-with","title":"Make Sure You're Unsure: A Framework for Verifying Probabilistic Specifications","date":"2021-02-18","arxiv_id":"2102.09479","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-dataset-and-benchmark-for-malaria-life","title":"A Dataset and Benchmark for Malaria Life-Cycle Classification in Thin Blood Smear Images","date":"2021-02-17","arxiv_id":"2102.08708","n_code_links":0,"syntology":null},{"paper":"/paper/beyond-fully-connected-layers-with","slug":"beyond-fully-connected-layers-with","title":"Beyond Fully-Connected Layers with Quaternions: Parameterization of Hypercomplex Multiplications with $1/n$ Parameters","date":"2021-02-17","arxiv_id":"2102.08597","n_code_links":3,"syntology":{"ran":0,"of":1,"n_ran_checked":0,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"0 ran · 1 unverified","official":{"repos":["astonzhang/Parameterization-of-Hypercomplex-Multiplications"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":[]}}},{"paper":null,"slug":"leveraging-query-resolution-and-reading","title":"Leveraging Query Resolution and Reading Comprehension for Conversational Passage Retrieval","date":"2021-02-17","arxiv_id":"2102.08795","n_code_links":0,"syntology":null},{"paper":null,"slug":"rethinking-co-design-of-neural-architectures","title":"Rethinking Co-design of Neural Architectures and Hardware Accelerators","date":"2021-02-17","arxiv_id":"2102.08619","n_code_links":0,"syntology":null},{"paper":"/paper/scidr-at-sdu-2020-ideas-identifying-and","slug":"scidr-at-sdu-2020-ideas-identifying-and","title":"SciDr at SDU-2020: IDEAS -- Identifying and Disambiguating Everyday Acronyms for Scientific Domain","date":"2021-02-17","arxiv_id":"2102.08818","n_code_links":2,"syntology":null},{"paper":"/paper/tcn-table-convolutional-network-for-web-table","slug":"tcn-table-convolutional-network-for-web-table","title":"TCN: Table Convolutional Network for Web Table Interpretation","date":"2021-02-17","arxiv_id":"2102.09460","n_code_links":1,"syntology":null},{"paper":null,"slug":"theaitre-1-0-interactive-generation-of","title":"THEaiTRE 1.0: Interactive generation of theatre play scripts","date":"2021-02-17","arxiv_id":"2102.08892","n_code_links":0,"syntology":null},{"paper":"/paper/an-automl-based-approach-to-multimodal-image","slug":"an-automl-based-approach-to-multimodal-image","title":"An AutoML-based Approach to Multimodal Image Sentiment Analysis","date":"2021-02-16","arxiv_id":"2102.08092","n_code_links":0,"syntology":null},{"paper":"/paper/coco-lm-correcting-and-contrasting-text","slug":"coco-lm-correcting-and-contrasting-text","title":"COCO-LM: Correcting and Contrasting Text Sequences for Language Model Pretraining","date":"2021-02-16","arxiv_id":"2102.08473","n_code_links":2,"syntology":{"ran":5,"of":6,"n_ran_checked":1,"n_instrument":4,"unverified":1,"pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 4 where Syntology's instrument failed) · 1 unverified","official":{"repos":["microsoft/coco-lm"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/exploring-transformers-in-natural-language","slug":"exploring-transformers-in-natural-language","title":"Exploring Transformers in Natural Language Generation: GPT, BERT, and XLNet","date":"2021-02-16","arxiv_id":"2102.08036","n_code_links":1,"syntology":null},{"paper":"/paper/gradinit-learning-to-initialize-neural","slug":"gradinit-learning-to-initialize-neural","title":"GradInit: Learning to Initialize Neural Networks for Stable and Efficient Training","date":"2021-02-16","arxiv_id":"2102.08098","n_code_links":2,"syntology":{"ran":7,"of":13,"n_ran_checked":4,"n_instrument":3,"unverified":6,"pointer_only":12,"phrase":"7 ran (of which 4 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 3 where Syntology's instrument failed) · 6 unverified","official":{"repos":["zhuchen03/gradinit"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":3,"n_ran_no_instrument_failure":3,"n_unverified":6,"ran_from_kinds":["listed","official"]}}},{"paper":null,"slug":"have-attention-heads-in-bert-learned","title":"Have Attention Heads in BERT Learned Constituency Grammar?","date":"2021-02-16","arxiv_id":"2102.07926","n_code_links":0,"syntology":null},{"paper":null,"slug":"improving-bayesian-inference-in-deep-neural","title":"Structured Dropout Variational Inference for Bayesian Neural Networks","date":"2021-02-16","arxiv_id":"2102.07927","n_code_links":0,"syntology":null},{"paper":"/paper/non-autoregressive-text-generation-with-pre","slug":"non-autoregressive-text-generation-with-pre","title":"Non-Autoregressive Text Generation with Pre-trained Language Models","date":"2021-02-16","arxiv_id":"2102.08220","n_code_links":1,"syntology":null},{"paper":"/paper/revisiting-language-encoding-in-learning","slug":"revisiting-language-encoding-in-learning","title":"Revisiting Language Encoding in Learning Multilingual Representations","date":"2021-02-16","arxiv_id":"2102.08357","n_code_links":1,"syntology":null},{"paper":"/paper/terapipe-token-level-pipeline-parallelism-for","slug":"terapipe-token-level-pipeline-parallelism-for","title":"TeraPipe: Token-Level Pipeline Parallelism for Training Large-Scale Language Models","date":"2021-02-16","arxiv_id":"2102.07988","n_code_links":1,"syntology":{"ran":7,"of":8,"n_ran_checked":3,"n_instrument":4,"unverified":1,"pointer_only":8,"phrase":"7 ran (of which 3 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 4 where Syntology's instrument failed) · 1 unverified","official":{"repos":["zhuohan123/terapipe"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":3,"n_ran_no_instrument_failure":3,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"training-larger-networks-for-deep","title":"Training Larger Networks for Deep Reinforcement Learning","date":"2021-02-16","arxiv_id":"2102.07920","n_code_links":0,"syntology":null},{"paper":null,"slug":"colored-kimia-path24-dataset-configurations","title":"Colored Kimia Path24 Dataset: Configurations and Benchmarks with Deep Embeddings","date":"2021-02-15","arxiv_id":"2102.07611","n_code_links":0,"syntology":null},{"paper":null,"slug":"detection-and-severity-classification-of","title":"Detection and severity classification of COVID-19 in CT images using deep learning","date":"2021-02-15","arxiv_id":"2102.07726","n_code_links":0,"syntology":null},{"paper":"/paper/dobf-a-deobfuscation-pre-training-objective","slug":"dobf-a-deobfuscation-pre-training-objective","title":"DOBF: A Deobfuscation Pre-Training Objective for Programming Languages","date":"2021-02-15","arxiv_id":"2102.07492","n_code_links":2,"syntology":null},{"paper":null,"slug":"fast-end-to-end-speech-recognition-via-non","title":"Fast End-to-End Speech Recognition via Non-Autoregressive Models and Cross-Modal Knowledge Transferring from BERT","date":"2021-02-15","arxiv_id":"2102.07594","n_code_links":0,"syntology":null}],"record_sha256":"88b70482f85339b3b80735158dbe661e0202f22f5f3fdf07a8a1234b1a017c1e","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}