{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/residual-connection/papers/270","list_of":"/method/residual-connection","method":"Residual Connection","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":270,"pages_in_order":285,"rows_per_page":100,"rows":[26901,27000],"of":28401,"counts":{"archive_papers_tagged":28401,"with_a_code_link":12847,"where_syntology_ran_a_sample":3897,"not_listed_spam_title":0,"listed":28401,"listed_where_code_ran":3897,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":3291,"every_run_a_failure_of_syntologys_instrument":606,"listed_with_a_run_with_no_instrument_failure":3291,"listed_every_run_a_failure_of_syntologys_instrument":606,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/residual-connection","prev":"/method/residual-connection/papers/269","next":"/method/residual-connection/papers/271","papers":[{"paper":"/paper/improving-neural-machine-translation-with-3","slug":"improving-neural-machine-translation-with-3","title":"Enhancing Machine Translation with Dependency-Aware Self-Attention","date":"2019-09-06","arxiv_id":"1909.03149","n_code_links":1,"syntology":null},{"paper":"/paper/rnn-architecture-learning-with-sparse","slug":"rnn-architecture-learning-with-sparse","title":"RNN Architecture Learning with Sparse Regularization","date":"2019-09-06","arxiv_id":"1909.03011","n_code_links":1,"syntology":null},{"paper":"/paper/supervised-multimodal-bitransformers-for","slug":"supervised-multimodal-bitransformers-for","title":"Supervised Multimodal Bitransformers for Classifying Images and Text","date":"2019-09-06","arxiv_id":"1909.02950","n_code_links":6,"syntology":null},{"paper":"/paper/a-stack-propagation-framework-with-token","slug":"a-stack-propagation-framework-with-token","title":"A Stack-Propagation Framework with Token-Level Intent Detection for Spoken Language Understanding","date":"2019-09-05","arxiv_id":"1909.02188","n_code_links":2,"syntology":null},{"paper":null,"slug":"accelerating-transformer-decoding-via-a","title":"Accelerating Transformer Decoding via a Hybrid of Self-attention and Recurrent Neural Network","date":"2019-09-05","arxiv_id":"1909.02279","n_code_links":0,"syntology":null},{"paper":"/paper/effective-use-of-transformer-networks-for","slug":"effective-use-of-transformer-networks-for","title":"Effective Use of Transformer Networks for Entity Tracking","date":"2019-09-05","arxiv_id":"1909.02635","n_code_links":1,"syntology":{"ran":3,"of":5,"n_ran_checked":2,"n_instrument":1,"unverified":2,"pointer_only":5,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 2 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","official":{"repos":["aditya2211/transformer-entity-tracking"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":2,"ran_from_kinds":["official","unlocated"]}}},{"paper":null,"slug":"efficient-neural-architecture-transformation","title":"Efficient Neural Architecture Transformation Searchin Channel-Level for Object Detection","date":"2019-09-05","arxiv_id":"1909.02293","n_code_links":0,"syntology":null},{"paper":"/paper/freeanchor-learning-to-match-anchors-for","slug":"freeanchor-learning-to-match-anchors-for","title":"FreeAnchor: Learning to Match Anchors for Visual Object Detection","date":"2019-09-05","arxiv_id":"1909.02466","n_code_links":4,"syntology":{"ran":3,"of":3,"n_ran_checked":0,"n_instrument":3,"unverified":0,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 3 where Syntology's instrument failed) · 0 unverified","official":{"repos":["zhangxiaosong18/FreeAnchor"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/in-plain-sight-media-bias-through-the-lens-of","slug":"in-plain-sight-media-bias-through-the-lens-of","title":"In Plain Sight: Media Bias Through the Lens of Factual Reporting","date":"2019-09-05","arxiv_id":"1909.02670","n_code_links":1,"syntology":null},{"paper":"/paper/informing-unsupervised-pretraining-with","slug":"informing-unsupervised-pretraining-with","title":"Specializing Unsupervised Pretraining Models for Word-Level Semantic Similarity","date":"2019-09-05","arxiv_id":"1909.02339","n_code_links":1,"syntology":null},{"paper":"/paper/investigating-berts-knowledge-of-language","slug":"investigating-berts-knowledge-of-language","title":"Investigating BERT's Knowledge of Language: Five Analysis Methods with NPIs","date":"2019-09-05","arxiv_id":"1909.02597","n_code_links":1,"syntology":{"ran":0,"of":3,"n_ran_checked":0,"n_instrument":0,"unverified":3,"pointer_only":3,"phrase":"0 ran · 3 unverified","official":{"repos":["alexwarstadt/data_generation"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":3,"ran_from_kinds":[]}}},{"paper":"/paper/semantic-correlation-promoted-shape-variant-1","slug":"semantic-correlation-promoted-shape-variant-1","title":"Semantic Correlation Promoted Shape-Variant Context for Segmentation","date":"2019-09-05","arxiv_id":"1909.02651","n_code_links":1,"syntology":null},{"paper":"/paper/semantics-aware-bert-for-language","slug":"semantics-aware-bert-for-language","title":"Semantics-aware BERT for Language Understanding","date":"2019-09-05","arxiv_id":"1909.02209","n_code_links":1,"syntology":{"ran":7,"of":12,"n_ran_checked":5,"n_instrument":2,"unverified":5,"pointer_only":3,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 2 honoured, 0 violated, 3 with no contract checked; 2 where Syntology's instrument failed) · 5 unverified","official":{"repos":["cooelf/SemBERT"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":5,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"source-dependency-aware-transformer-with","title":"Source Dependency-Aware Transformer with Supervised Self-Attention","date":"2019-09-05","arxiv_id":"1909.02273","n_code_links":0,"syntology":null},{"paper":null,"slug":"syntax-aware-aspect-level-sentiment","title":"Syntax-Aware Aspect Level Sentiment Classification with Graph Attention Networks","date":"2019-09-05","arxiv_id":"1909.02606","n_code_links":0,"syntology":null},{"paper":null,"slug":"deep-learning-aided-tabu-search-detection-for","title":"Deep Learning-Aided Tabu Search Detection for Large MIMO Systems","date":"2019-09-04","arxiv_id":"1909.01683","n_code_links":0,"syntology":null},{"paper":"/paper/dense-extreme-inception-network-towards-a","slug":"dense-extreme-inception-network-towards-a","title":"Dense Extreme Inception Network: Towards a Robust CNN Model for Edge Detection","date":"2019-09-04","arxiv_id":"1909.01955","n_code_links":4,"syntology":{"ran":17,"of":19,"n_ran_checked":16,"n_instrument":1,"unverified":2,"pointer_only":0,"phrase":"17 ran (of which 0 constructed an object rather than computing a result; 16 with no instrument failure: 0 honoured, 0 violated, 16 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","official":{"repos":["xavysp/DexiNed"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed","official"]}}},{"paper":"/paper/jointly-learning-to-align-and-translate-with","slug":"jointly-learning-to-align-and-translate-with","title":"Jointly Learning to Align and Translate with Transformer Models","date":"2019-09-04","arxiv_id":"1909.02074","n_code_links":1,"syntology":null},{"paper":"/paper/mogrifier-lstm","slug":"mogrifier-lstm","title":"Mogrifier LSTM","date":"2019-09-04","arxiv_id":"1909.01792","n_code_links":3,"syntology":null},{"paper":"/paper/encode-tag-realize-high-precision-text","slug":"encode-tag-realize-high-precision-text","title":"Encode, Tag, Realize: High-Precision Text Editing","date":"2019-09-03","arxiv_id":"1909.01187","n_code_links":5,"syntology":{"ran":8,"of":11,"n_ran_checked":7,"n_instrument":1,"unverified":3,"pointer_only":1,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 1 honoured, 0 violated, 6 with no contract checked; 1 where Syntology's instrument failed) · 3 unverified","official":{"repos":["google-research/lasertagger"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":1,"ran_from_kinds":["listed","official"]}}},{"paper":"/paper/hardnet-a-low-memory-traffic-network","slug":"hardnet-a-low-memory-traffic-network","title":"HarDNet: A Low Memory Traffic Network","date":"2019-09-03","arxiv_id":"1909.00948","n_code_links":24,"syntology":null},{"paper":"/paper/language-models-as-knowledge-bases","slug":"language-models-as-knowledge-bases","title":"Language Models as Knowledge Bases?","date":"2019-09-03","arxiv_id":"1909.01066","n_code_links":1,"syntology":null},{"paper":null,"slug":"multimodal-deep-learning-for-mental-disorders","title":"Multimodal Deep Learning for Mental Disorders Prediction from Audio Speech Samples","date":"2019-09-03","arxiv_id":"1909.01067","n_code_links":0,"syntology":null},{"paper":null,"slug":"psdnet-and-dpdnet-efficient-channel-expansion","title":"PSDNet and DPDNet: Efficient channel expansion, Depthwise-Pointwise-Depthwise Inverted Bottleneck Block","date":"2019-09-03","arxiv_id":"1909.01026","n_code_links":0,"syntology":null},{"paper":null,"slug":"the-bottom-up-evolution-of-representations-in","title":"The Bottom-up Evolution of Representations in the Transformer: A Study with Machine Translation and Language Modeling Objectives","date":"2019-09-03","arxiv_id":"1909.01380","n_code_links":0,"syntology":null},{"paper":"/paper/transfer-fine-tuning-a-bert-case-study","slug":"transfer-fine-tuning-a-bert-case-study","title":"Transfer Fine-Tuning: A BERT Case Study","date":"2019-09-03","arxiv_id":"1909.00931","n_code_links":1,"syntology":null},{"paper":null,"slug":"unicoder-a-universal-language-encoder-by-pre","title":"Unicoder: A Universal Language Encoder by Pre-training with Multiple Cross-lingual Tasks","date":"2019-09-03","arxiv_id":"1909.00964","n_code_links":0,"syntology":null},{"paper":null,"slug":"hishabnet-detection-localization-and","title":"HishabNet: Detection, Localization and Calculation of Handwritten Bengali Mathematical Expressions","date":"2019-09-02","arxiv_id":"1909.00823","n_code_links":0,"syntology":null},{"paper":"/paper/how-contextual-are-contextualized-word","slug":"how-contextual-are-contextualized-word","title":"How Contextual are Contextualized Word Representations? Comparing the Geometry of BERT, ELMo, and GPT-2 Embeddings","date":"2019-09-02","arxiv_id":"1909.00512","n_code_links":1,"syntology":null},{"paper":null,"slug":"logic-and-the-2-simplicial-transformer","title":"Logic and the $2$-Simplicial Transformer","date":"2019-09-02","arxiv_id":"1909.00668","n_code_links":0,"syntology":null},{"paper":"/paper/sumqe-a-bert-based-summary-quality-estimation","slug":"sumqe-a-bert-based-summary-quality-estimation","title":"SumQE: a BERT-based Summary Quality Estimation Model","date":"2019-09-02","arxiv_id":"1909.00578","n_code_links":1,"syntology":null},{"paper":"/paper/training-time-friendly-network-for-real-time","slug":"training-time-friendly-network-for-real-time","title":"Training-Time-Friendly Network for Real-Time Object Detection","date":"2019-09-02","arxiv_id":"1909.00700","n_code_links":6,"syntology":{"ran":5,"of":8,"n_ran_checked":5,"n_instrument":0,"unverified":3,"pointer_only":1,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":{"repos":["ZJULearning/ttfnet"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":null,"slug":"classification-approaches-to-identify","title":"Classification Approaches to Identify Informative Tweets","date":"2019-09-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/cross-lingual-machine-reading-comprehension","slug":"cross-lingual-machine-reading-comprehension","title":"Cross-Lingual Machine Reading Comprehension","date":"2019-09-01","arxiv_id":"1909.00361","n_code_links":1,"syntology":null},{"paper":null,"slug":"dependency-based-relative-positional-encoding","title":"Dependency-Based Relative Positional Encoding for Transformer NMT","date":"2019-09-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"dependency-based-self-attention-for","title":"Dependency-Based Self-Attention for Transformer NMT","date":"2019-09-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"evaluating-the-cross-lingual-effectiveness-of","title":"Evaluating the Cross-Lingual Effectiveness of Massively Multilingual Neural Machine Translation","date":"2019-09-01","arxiv_id":"1909.00437","n_code_links":0,"syntology":null},{"paper":null,"slug":"evaluation-of-stacked-embeddings-for","title":"Evaluation of Stacked Embeddings for Bulgarian on the Downstream Tasks POS and NERC","date":"2019-09-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"evaluation-of-vector-embedding-models-in","title":"Evaluation of vector embedding models in clustering of text documents","date":"2019-09-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"friendsqa-open-domain-question-answering-on","title":"FriendsQA: Open-Domain Question Answering on TV Show Transcripts","date":"2019-09-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/incidental-supervision-from-question","slug":"incidental-supervision-from-question","title":"QuASE: Question-Answer Driven Sentence Encoding","date":"2019-09-01","arxiv_id":"1909.00333","n_code_links":1,"syntology":null},{"paper":null,"slug":"multilingual-language-models-for-named-entity","title":"Multilingual Language Models for Named Entity Recognition in German and English","date":"2019-09-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"multilingual-probing-of-deep-pre-trained","title":"Multilingual Probing of Deep Pre-Trained Contextual Encoders","date":"2019-09-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/pre-training-of-deep-contextualized","slug":"pre-training-of-deep-contextualized","title":"Global Entity Disambiguation with BERT","date":"2019-09-01","arxiv_id":"1909.00426","n_code_links":1,"syntology":null},{"paper":null,"slug":"predicting-sentiment-of-polish-language-short","title":"Predicting Sentiment of Polish Language Short Texts","date":"2019-09-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"semantic-role-labeling-with-pretrained","title":"Semantic Role Labeling with Pretrained Language Models for Known and Unknown Predicates","date":"2019-09-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"turkish-tweet-classification-with-transformer","title":"Turkish Tweet Classification with Transformer Encoder","date":"2019-09-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"adversarial-learning-with-contextual","title":"Adversarial Learning with Contextual Embeddings for Zero-resource Cross-lingual Classification and NER","date":"2019-08-31","arxiv_id":"1909.00153","n_code_links":0,"syntology":null},{"paper":"/paper/evaluation-benchmarks-and-learning","slug":"evaluation-benchmarks-and-learning","title":"Evaluation Benchmarks and Learning Criteria for Discourse-Aware Sentence Representations","date":"2019-08-31","arxiv_id":"1909.00142","n_code_links":2,"syntology":{"ran":2,"of":2,"n_ran_checked":1,"n_instrument":1,"unverified":0,"pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["ZeweiChu/DiscoEval"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/humor-detection-a-transformer-gets-the-last","slug":"humor-detection-a-transformer-gets-the-last","title":"Humor Detection: A Transformer Gets the Last Laugh","date":"2019-08-31","arxiv_id":"1909.00252","n_code_links":2,"syntology":{"ran":3,"of":4,"n_ran_checked":1,"n_instrument":2,"unverified":1,"pointer_only":2,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 1 unverified","official":{"repos":["orionw/RedditHumorDetection"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":["official","unlocated"]}}},{"paper":null,"slug":"improving-multi-head-attention-with-capsule","title":"Improving Multi-Head Attention with Capsule Networks","date":"2019-08-31","arxiv_id":"1909.00188","n_code_links":0,"syntology":null},{"paper":null,"slug":"knowledge-enhanced-attention-for-robust","title":"Knowledge Enhanced Attention for Robust Natural Language Inference","date":"2019-08-31","arxiv_id":"1909.00102","n_code_links":0,"syntology":null},{"paper":"/paper/modeling-graph-structure-in-transformer-for","slug":"modeling-graph-structure-in-transformer-for","title":"Modeling Graph Structure in Transformer for Better AMR-to-Text Generation","date":"2019-08-31","arxiv_id":"1909.00136","n_code_links":1,"syntology":null},{"paper":"/paper/nezha-neural-contextualized-representation","slug":"nezha-neural-contextualized-representation","title":"NEZHA: Neural Contextualized Representation for Chinese Language Understanding","date":"2019-08-31","arxiv_id":"1909.00204","n_code_links":10,"syntology":null},{"paper":null,"slug":"quantity-doesnt-buy-quality-syntax-with","title":"Quantity doesn't buy quality syntax with neural language models","date":"2019-08-31","arxiv_id":"1909.00111","n_code_links":0,"syntology":null},{"paper":null,"slug":"small-and-practical-bert-models-for-sequence","title":"Small and Practical BERT Models for Sequence Labeling","date":"2019-08-31","arxiv_id":"1909.00100","n_code_links":0,"syntology":null},{"paper":"/paper/adapt-or-get-left-behind-domain-adaptation","slug":"adapt-or-get-left-behind-domain-adaptation","title":"Adapt or Get Left Behind: Domain Adaptation through BERT Language Model Finetuning for Aspect-Target Sentiment Classification","date":"2019-08-30","arxiv_id":"1908.11860","n_code_links":3,"syntology":{"ran":6,"of":7,"n_ran_checked":6,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["deepopinion/domain-adapted-atsc"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":["listed"]}}},{"paper":"/paper/adaptively-sparse-transformers","slug":"adaptively-sparse-transformers","title":"Adaptively Sparse Transformers","date":"2019-08-30","arxiv_id":"1909.00015","n_code_links":3,"syntology":{"ran":3,"of":3,"n_ran_checked":0,"n_instrument":3,"unverified":0,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 3 where Syntology's instrument failed) · 0 unverified","official":{"repos":["deep-spin/entmax"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed","official"]}}},{"paper":null,"slug":"answering-conversational-questions-on","title":"Answering Conversational Questions on Structured Data without Logical Forms","date":"2019-08-30","arxiv_id":"1908.11787","n_code_links":0,"syntology":null},{"paper":"/paper/bilingual-is-at-least-monolingual-balm-a","slug":"bilingual-is-at-least-monolingual-balm-a","title":"Bilingual is At Least Monolingual (BALM): A Novel Translation Algorithm that Encodes Monolingual Priors","date":"2019-08-30","arxiv_id":"1909.01146","n_code_links":1,"syntology":null},{"paper":null,"slug":"learning-rich-representations-for-structured","title":"Learning Rich Representations For Structured Visual Prediction Tasks","date":"2019-08-30","arxiv_id":"1908.11820","n_code_links":0,"syntology":null},{"paper":null,"slug":"neural-language-model-for-automated","title":"Pre-training A Neural Language Model Improves The Sample Efficiency of an Emergency Room Classification Model","date":"2019-08-30","arxiv_id":"1909.01136","n_code_links":0,"syntology":null},{"paper":"/paper/paws-x-a-cross-lingual-adversarial-dataset","slug":"paws-x-a-cross-lingual-adversarial-dataset","title":"PAWS-X: A Cross-lingual Adversarial Dataset for Paraphrase Identification","date":"2019-08-30","arxiv_id":"1908.11828","n_code_links":3,"syntology":null},{"paper":"/paper/revisiting-cyclegan-for-semi-supervised","slug":"revisiting-cyclegan-for-semi-supervised","title":"Revisiting CycleGAN for semi-supervised segmentation","date":"2019-08-30","arxiv_id":"1908.11569","n_code_links":1,"syntology":null},{"paper":"/paper/transformer-dissection-an-unified","slug":"transformer-dissection-an-unified","title":"Transformer Dissection: A Unified Understanding of Transformer's Attention via the Lens of Kernel","date":"2019-08-30","arxiv_id":"1908.11775","n_code_links":1,"syntology":null},{"paper":"/paper/improving-deep-transformer-with-depth-scaled","slug":"improving-deep-transformer-with-depth-scaled","title":"Improving Deep Transformer with Depth-Scaled Initialization and Merged Attention","date":"2019-08-29","arxiv_id":"1908.11365","n_code_links":1,"syntology":null},{"paper":null,"slug":"probing-representations-learned-by-multimodal","title":"Probing Representations Learned by Multimodal Recurrent and Transformer Models","date":"2019-08-29","arxiv_id":"1908.11125","n_code_links":0,"syntology":null},{"paper":null,"slug":"regularized-context-gates-on-transformer-for","title":"Regularized Context Gates on Transformer for Machine Translation","date":"2019-08-29","arxiv_id":"1908.11020","n_code_links":0,"syntology":null},{"paper":"/paper/temporal-consistency-objectives-regularize","slug":"temporal-consistency-objectives-regularize","title":"Temporal Consistency Objectives Regularize the Learning of Disentangled Representations","date":"2019-08-29","arxiv_id":"1908.11330","n_code_links":1,"syntology":null},{"paper":null,"slug":"adversarial-representation-learning-for-text","title":"Adversarial Representation Learning for Text-to-Image Matching","date":"2019-08-28","arxiv_id":"1908.10534","n_code_links":0,"syntology":null},{"paper":null,"slug":"approxnet-content-and-contention-aware-video","title":"ApproxNet: Content and Contention-Aware Video Analytics System for Embedded Clients","date":"2019-08-28","arxiv_id":"1909.02068","n_code_links":0,"syntology":null},{"paper":null,"slug":"solving-math-word-problems-with-double","title":"Solving Math Word Problems with Double-Decoder Transformer","date":"2019-08-28","arxiv_id":"1908.10924","n_code_links":0,"syntology":null},{"paper":"/paper/bottom-up-higher-resolution-networks-for","slug":"bottom-up-higher-resolution-networks-for","title":"HigherHRNet: Scale-Aware Representation Learning for Bottom-Up Human Pose Estimation","date":"2019-08-27","arxiv_id":"1908.10357","n_code_links":19,"syntology":{"ran":17,"of":27,"n_ran_checked":13,"n_instrument":4,"unverified":10,"pointer_only":2,"phrase":"17 ran (of which 0 constructed an object rather than computing a result; 13 with no instrument failure: 0 honoured, 0 violated, 13 with no contract checked; 4 where Syntology's instrument failed) · 10 unverified","official":{"repos":["HRNet/Higher-HRNet-Human-Pose-Estimation"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":3,"ran_from_kinds":["listed","official"]}}},{"paper":null,"slug":"bridging-the-gap-for-tokenizer-free-language","title":"Bridging the Gap for Tokenizer-Free Language Models","date":"2019-08-27","arxiv_id":"1908.10322","n_code_links":0,"syntology":null},{"paper":null,"slug":"domain-agnostic-learning-with-anatomy","title":"Domain-Agnostic Learning with Anatomy-Consistent Embedding for Cross-Modality Liver Segmentation","date":"2019-08-27","arxiv_id":"1908.10489","n_code_links":0,"syntology":null},{"paper":"/paper/finbert-financial-sentiment-analysis-with-pre","slug":"finbert-financial-sentiment-analysis-with-pre","title":"FinBERT: Financial Sentiment Analysis with Pre-trained Language Models","date":"2019-08-27","arxiv_id":"1908.10063","n_code_links":3,"syntology":{"ran":3,"of":3,"n_ran_checked":3,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 1 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":null}},{"paper":null,"slug":"multiresolution-transformer-networks","title":"Multiresolution Transformer Networks: Recurrence is Not Essential for Modeling Hierarchical Structure","date":"2019-08-27","arxiv_id":"1908.10408","n_code_links":0,"syntology":null},{"paper":null,"slug":"on-nmt-search-errors-and-model-errors-cat-got","title":"On NMT Search Errors and Model Errors: Cat Got Your Tongue?","date":"2019-08-27","arxiv_id":"1908.10090","n_code_links":0,"syntology":null},{"paper":"/paper/sentence-bert-sentence-embeddings-using","slug":"sentence-bert-sentence-embeddings-using","title":"Sentence-BERT: Sentence Embeddings using Siamese BERT-Networks","date":"2019-08-27","arxiv_id":"1908.10084","n_code_links":64,"syntology":{"ran":33,"of":58,"n_ran_checked":30,"n_instrument":3,"unverified":25,"pointer_only":11,"phrase":"33 ran (of which 9 constructed an object rather than computing a result; 30 with no instrument failure: 1 honoured, 0 violated, 29 with no contract checked; 3 where Syntology's instrument failed) · 25 unverified","official":{"repos":["UKPLab/sentence-transformers"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/temporal-reasoning-graph-for-activity","slug":"temporal-reasoning-graph-for-activity","title":"Temporal Reasoning Graph for Activity Recognition","date":"2019-08-27","arxiv_id":"1908.09995","n_code_links":0,"syntology":null},{"paper":"/paper/attentive-history-selection-for","slug":"attentive-history-selection-for","title":"Attentive History Selection for Conversational Question Answering","date":"2019-08-26","arxiv_id":"1908.09456","n_code_links":2,"syntology":null},{"paper":"/paper/confidence-regularized-self-training","slug":"confidence-regularized-self-training","title":"Confidence Regularized Self-Training","date":"2019-08-26","arxiv_id":"1908.09822","n_code_links":2,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["yzou2/CRST"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/constructing-self-motivated-pyramid","slug":"constructing-self-motivated-pyramid","title":"Constructing Self-motivated Pyramid Curriculums for Cross-Domain Semantic Segmentation: A Non-Adversarial Approach","date":"2019-08-26","arxiv_id":"1908.09547","n_code_links":1,"syntology":null},{"paper":null,"slug":"cross-modality-knowledge-transfer-for","title":"Cross-modality Knowledge Transfer for Prostate Segmentation from CT Scans","date":"2019-08-26","arxiv_id":"1908.10208","n_code_links":0,"syntology":null},{"paper":null,"slug":"cyclegan-with-a-blur-kernel-for-deconvolution","title":"CycleGAN with a Blur Kernel for Deconvolution Microscopy: Optimal Transport Geometry","date":"2019-08-26","arxiv_id":"1908.09414","n_code_links":0,"syntology":null},{"paper":"/paper/detecting-toxicity-in-news-articles","slug":"detecting-toxicity-in-news-articles","title":"Detecting Toxicity in News Articles: Application to Bulgarian","date":"2019-08-26","arxiv_id":"1908.09785","n_code_links":1,"syntology":null},{"paper":"/paper/does-bert-agree-evaluating-knowledge-of","slug":"does-bert-agree-evaluating-knowledge-of","title":"Does BERT agree? Evaluating knowledge of structure dependence through agreement relations","date":"2019-08-26","arxiv_id":"1908.09892","n_code_links":1,"syntology":null},{"paper":null,"slug":"learning-disentangled-representations-via","title":"Learning Disentangled Representations via Independent Subspaces","date":"2019-08-26","arxiv_id":"1908.08989","n_code_links":0,"syntology":null},{"paper":null,"slug":"measuring-patent-claim-generation-by-span","title":"Measuring Patent Claim Generation by Span Relevancy","date":"2019-08-26","arxiv_id":"1908.09591","n_code_links":0,"syntology":null},{"paper":null,"slug":"see-more-than-once-kernel-sharing-atrous","title":"See More Than Once -- Kernel-Sharing Atrous Convolution for Semantic Segmentation","date":"2019-08-26","arxiv_id":"1908.09443","n_code_links":0,"syntology":null},{"paper":null,"slug":"slidergan-synthesizing-expressive-face-images","title":"SliderGAN: Synthesizing Expressive Face Images by Sliding 3D Blendshape Parameters","date":"2019-08-26","arxiv_id":"1908.09638","n_code_links":0,"syntology":null},{"paper":"/paper/patient-knowledge-distillation-for-bert-model","slug":"patient-knowledge-distillation-for-bert-model","title":"Patient Knowledge Distillation for BERT Model Compression","date":"2019-08-25","arxiv_id":"1908.09355","n_code_links":5,"syntology":{"ran":19,"of":27,"n_ran_checked":13,"n_instrument":6,"unverified":8,"pointer_only":27,"phrase":"19 ran (of which 0 constructed an object rather than computing a result; 13 with no instrument failure: 2 honoured, 1 violated, 10 with no contract checked; 6 where Syntology's instrument failed) · 8 unverified","official":{"repos":["intersun/PKD-for-BERT-Model-Compression"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/transforming-delete-retrieve-generate","slug":"transforming-delete-retrieve-generate","title":"Transforming Delete, Retrieve, Generate Approach for Controlled Text Style Transfer","date":"2019-08-25","arxiv_id":"1908.09368","n_code_links":1,"syntology":null},{"paper":"/paper/bert-for-coreference-resolution-baselines-and","slug":"bert-for-coreference-resolution-baselines-and","title":"BERT for Coreference Resolution: Baselines and Analysis","date":"2019-08-24","arxiv_id":"1908.09091","n_code_links":2,"syntology":null},{"paper":null,"slug":"release-strategies-and-the-social-impacts-of","title":"Release Strategies and the Social Impacts of Language Models","date":"2019-08-24","arxiv_id":"1908.09203","n_code_links":0,"syntology":null},{"paper":"/paper/deep-learning-based-chatbot-models","slug":"deep-learning-based-chatbot-models","title":"Deep Learning Based Chatbot Models","date":"2019-08-23","arxiv_id":"1908.08835","n_code_links":1,"syntology":null},{"paper":"/paper/mish-a-self-regularized-non-monotonic-neural","slug":"mish-a-self-regularized-non-monotonic-neural","title":"Mish: A Self Regularized Non-Monotonic Activation Function","date":"2019-08-23","arxiv_id":"1908.08681","n_code_links":9,"syntology":{"ran":9,"of":12,"n_ran_checked":8,"n_instrument":1,"unverified":3,"pointer_only":1,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 1 honoured, 0 violated, 7 with no contract checked; 1 where Syntology's instrument failed) · 3 unverified","official":{"repos":["digantamisra98/Mish"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":3,"ran_from_kinds":["listed","official"]}}},{"paper":"/paper/neural-data-to-text-generation-a-comparison","slug":"neural-data-to-text-generation-a-comparison","title":"Neural data-to-text generation: A comparison between pipeline and end-to-end architectures","date":"2019-08-23","arxiv_id":"1908.09022","n_code_links":1,"syntology":null},{"paper":"/paper/shadow-removal-via-shadow-image-decomposition","slug":"shadow-removal-via-shadow-image-decomposition","title":"Shadow Removal via Shadow Image Decomposition","date":"2019-08-23","arxiv_id":"1908.08628","n_code_links":3,"syntology":{"ran":8,"of":11,"n_ran_checked":8,"n_instrument":0,"unverified":3,"pointer_only":2,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 0 violated, 8 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":{"repos":["lmhieu612/SID"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":3,"ran_from_kinds":["listed"]}}},{"paper":null,"slug":"training-optimus-prime-md-generating-medical","title":"Training Optimus Prime, M.D.: Generating Medical Certification Items by Fine-Tuning OpenAI's gpt2 Transformer Model","date":"2019-08-23","arxiv_id":"1908.08594","n_code_links":0,"syntology":null}],"record_sha256":"8396201090c6dde2ac4ebeb6fe039480a64e43dba16ed032de9b391d6f30ba4b","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}