{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/dropout/papers/251","list_of":"/method/dropout","method":"Dropout","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":251,"pages_in_order":275,"rows_per_page":100,"rows":[25001,25100],"of":27472,"counts":{"archive_papers_tagged":27472,"with_a_code_link":12129,"where_syntology_ran_a_sample":3620,"not_listed_spam_title":0,"listed":27472,"listed_where_code_ran":3620,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":3044,"every_run_a_failure_of_syntologys_instrument":576,"listed_with_a_run_with_no_instrument_failure":3044,"listed_every_run_a_failure_of_syntologys_instrument":576,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/dropout","prev":"/method/dropout/papers/250","next":"/method/dropout/papers/252","papers":[{"paper":null,"slug":"learning-long-and-short-term-user-literal","title":"Learning Long- and Short-Term User Literal-Preference with Multimodal Hierarchical Transformer Network for Personalized Image Caption","date":"2020-02-04","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"multistage-model-for-robust-face-alignment","title":"Multistage Model for Robust Face Alignment Using Deep Neural Networks","date":"2020-02-04","arxiv_id":"2002.01075","n_code_links":0,"syntology":null},{"paper":"/paper/bertrand-dr-improving-text-to-sql-using-a","slug":"bertrand-dr-improving-text-to-sql-using-a","title":"Bertrand-DR: Improving Text-to-SQL using a Discriminative Re-ranker","date":"2020-02-03","arxiv_id":"2002.00557","n_code_links":1,"syntology":null},{"paper":null,"slug":"exponential-discretization-of-weights-of","title":"Exponential discretization of weights of neural network connections in pre-trained neural networks","date":"2020-02-03","arxiv_id":"2002.00623","n_code_links":0,"syntology":null},{"paper":"/paper/iart-intent-aware-response-ranking-with","slug":"iart-intent-aware-response-ranking-with","title":"IART: Intent-aware Response Ranking with Transformers in Information-seeking Conversation Systems","date":"2020-02-03","arxiv_id":"2002.00571","n_code_links":1,"syntology":null},{"paper":"/paper/pix2pix-based-stain-to-stain-translation-a","slug":"pix2pix-based-stain-to-stain-translation-a","title":"Pix2Pix-based Stain-to-Stain Translation: A Solution for Robust Stain Normalization in Histopathology Images Analysis","date":"2020-02-03","arxiv_id":"2002.00647","n_code_links":1,"syntology":null},{"paper":"/paper/robust-saliency-maps-with-decoy-enhanced","slug":"robust-saliency-maps-with-decoy-enhanced","title":"DANCE: Enhancing saliency maps using decoys","date":"2020-02-03","arxiv_id":"2002.00526","n_code_links":1,"syntology":null},{"paper":"/paper/beat-the-ai-investigating-adversarial-human","slug":"beat-the-ai-investigating-adversarial-human","title":"Beat the AI: Investigating Adversarial Human Annotation for Reading Comprehension","date":"2020-02-02","arxiv_id":"2002.00293","n_code_links":1,"syntology":null},{"paper":"/paper/non-linear-neurons-with-human-like-apical","slug":"non-linear-neurons-with-human-like-apical","title":"Non-linear Neurons with Human-like Apical Dendrite Activations","date":"2020-02-02","arxiv_id":"2003.03229","n_code_links":1,"syntology":null},{"paper":"/paper/bridging-text-and-video-a-universal","slug":"bridging-text-and-video-a-universal","title":"Bridging Text and Video: A Universal Multimodal Transformer for Video-Audio Scene-Aware Dialog","date":"2020-02-01","arxiv_id":"2002.00163","n_code_links":1,"syntology":{"ran":5,"of":5,"n_ran_checked":3,"n_instrument":2,"unverified":0,"pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":null}},{"paper":null,"slug":"fine-tuning-bert-for-schema-guided-zero-shot","title":"Fine-Tuning BERT for Schema-Guided Zero-Shot Dialogue State Tracking","date":"2020-02-01","arxiv_id":"2002.00181","n_code_links":0,"syntology":null},{"paper":"/paper/pop-music-transformer-generating-music-with","slug":"pop-music-transformer-generating-music-with","title":"Pop Music Transformer: Beat-based Modeling and Generation of Expressive Pop Piano Compositions","date":"2020-02-01","arxiv_id":"2002.00212","n_code_links":7,"syntology":null},{"paper":null,"slug":"fast-monte-carlo-dropout-and-error-correction","title":"Fast Monte Carlo Dropout and Error Correction for Radio Transmitter Classification","date":"2020-01-31","arxiv_id":"2001.11963","n_code_links":0,"syntology":null},{"paper":"/paper/pretrained-transformers-for-simple-question","slug":"pretrained-transformers-for-simple-question","title":"Pretrained Transformers for Simple Question Answering over Knowledge Graphs","date":"2020-01-31","arxiv_id":"2001.11985","n_code_links":1,"syntology":null},{"paper":"/paper/adversarial-training-for-aspect-based","slug":"adversarial-training-for-aspect-based","title":"Adversarial Training for Aspect-Based Sentiment Analysis with BERT","date":"2020-01-30","arxiv_id":"2001.11316","n_code_links":4,"syntology":null},{"paper":null,"slug":"do-we-need-word-order-information-for-cross","title":"On the Importance of Word Order Information in Cross-lingual Sequence Labeling","date":"2020-01-30","arxiv_id":"2001.11164","n_code_links":0,"syntology":null},{"paper":"/paper/interpretable-rumor-detection-in-microblogs","slug":"interpretable-rumor-detection-in-microblogs","title":"Interpretable Rumor Detection in Microblogs by Attending to User Interactions","date":"2020-01-29","arxiv_id":"2001.10667","n_code_links":1,"syntology":null},{"paper":"/paper/pre-defined-sparsity-for-low-complexity","slug":"pre-defined-sparsity-for-low-complexity","title":"Pre-defined Sparsity for Low-Complexity Convolutional Neural Networks","date":"2020-01-29","arxiv_id":"2001.10710","n_code_links":1,"syntology":null},{"paper":null,"slug":"joint-contextual-modeling-for-asr-correction","title":"Joint Contextual Modeling for ASR Correction and Language Understanding","date":"2020-01-28","arxiv_id":"2002.00750","n_code_links":0,"syntology":null},{"paper":null,"slug":"pel-bert-a-joint-model-for-protocol-entity","title":"PEL-BERT: A Joint Model for Protocol Entity Linking","date":"2020-01-28","arxiv_id":"2002.00744","n_code_links":0,"syntology":null},{"paper":"/paper/artificial-neural-networks-for-cloud-masking","slug":"artificial-neural-networks-for-cloud-masking","title":"Artificial neural networks for cloud masking of Sentinel-2 ocean images with noise and sunglint","date":"2020-01-26","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"further-boosting-bert-based-models-by","title":"BERT's output layer recognizes all hidden layers? Some Intriguing Phenomena and a simple way to boost BERT","date":"2020-01-25","arxiv_id":"2001.09309","n_code_links":0,"syntology":null},{"paper":null,"slug":"generation-distillation-for-efficient-natural-1","title":"Generation-Distillation for Efficient Natural Language Understanding in Low-Data Settings","date":"2020-01-25","arxiv_id":"2002.00733","n_code_links":0,"syntology":null},{"paper":"/paper/gesticulator-a-framework-for-semantically","slug":"gesticulator-a-framework-for-semantically","title":"Gesticulator: A framework for semantically-aware speech-driven gesture generation","date":"2020-01-25","arxiv_id":"2001.09326","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["Svito-zar/gesticulator"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"compressing-language-models-using-doped","title":"Compressing Language Models using Doped Kronecker Products","date":"2020-01-24","arxiv_id":"2001.08896","n_code_links":0,"syntology":null},{"paper":"/paper/power-bert-accelerating-bert-inference-for","slug":"power-bert-accelerating-bert-inference-for","title":"PoWER-BERT: Accelerating BERT Inference via Progressive Word-vector Elimination","date":"2020-01-24","arxiv_id":"2001.08950","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":3,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"3 ran (of which 2 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["IBM/PoWER-BERT"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":2,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/stochastic-optimization-of-plain","slug":"stochastic-optimization-of-plain","title":"Stochastic Optimization of Plain Convolutional Neural Networks with Simple methods","date":"2020-01-24","arxiv_id":"2001.08856","n_code_links":1,"syntology":null},{"paper":null,"slug":"applying-recent-innovations-from-nlp-to-mooc","title":"Applying Recent Innovations from NLP to MOOC Student Course Trajectory Modeling","date":"2020-01-23","arxiv_id":"2001.08333","n_code_links":0,"syntology":null},{"paper":null,"slug":"fine-tuning-a-transformer-based-language","title":"Reducing Non-Normative Text Generation from Language Models","date":"2020-01-23","arxiv_id":"2001.08764","n_code_links":0,"syntology":null},{"paper":null,"slug":"navigation-based-candidate-expansion-and","title":"Navigation-Based Candidate Expansion and Pretrained Language Models for Citation Recommendation","date":"2020-01-23","arxiv_id":"2001.08687","n_code_links":0,"syntology":null},{"paper":"/paper/attention-a-lightweight-2d-hand-pose","slug":"attention-a-lightweight-2d-hand-pose","title":"Attention! A Lightweight 2D Hand Pose Estimation Approach","date":"2020-01-22","arxiv_id":"2001.08047","n_code_links":1,"syntology":null},{"paper":null,"slug":"autofcl-automatically-tuning-fully-connected","title":"AutoFCL: Automatically Tuning Fully Connected Layers for Handling Small Dataset","date":"2020-01-22","arxiv_id":"2001.11951","n_code_links":0,"syntology":null},{"paper":"/paper/multilingual-denoising-pre-training-for","slug":"multilingual-denoising-pre-training-for","title":"Multilingual Denoising Pre-training for Neural Machine Translation","date":"2020-01-22","arxiv_id":"2001.08210","n_code_links":8,"syntology":null},{"paper":"/paper/on-last-layer-algorithms-for-classification","slug":"on-last-layer-algorithms-for-classification","title":"On Last-Layer Algorithms for Classification: Decoupling Representation from Uncertainty Estimation","date":"2020-01-22","arxiv_id":"2001.08049","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["nbrosse/uncertainties"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/variational-dropout-sparsification-for","slug":"variational-dropout-sparsification-for","title":"Variational Dropout Sparsification for Particle Identification speed-up","date":"2020-01-21","arxiv_id":"2001.07493","n_code_links":1,"syntology":null},{"paper":null,"slug":"multi-level-head-wise-match-and-aggregation","title":"Multi-level Head-wise Match and Aggregation in Transformer for Textual Sequence Matching","date":"2020-01-20","arxiv_id":"2001.07234","n_code_links":0,"syntology":null},{"paper":"/paper/recommending-themes-for-ad-creative-design","slug":"recommending-themes-for-ad-creative-design","title":"Recommending Themes for Ad Creative Design via Visual-Linguistic Representations","date":"2020-01-20","arxiv_id":"2001.07194","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-multimodal-deep-learning-approach-for-named","title":"A multimodal deep learning approach for named entity recognition from social media","date":"2020-01-19","arxiv_id":"2001.06888","n_code_links":0,"syntology":null},{"paper":null,"slug":"deep-learning-for-hindi-text-classification-a","title":"Deep Learning for Hindi Text Classification: A Comparison","date":"2020-01-19","arxiv_id":"2001.10340","n_code_links":0,"syntology":null},{"paper":null,"slug":"capturing-evolution-in-word-usage-just-add","title":"Capturing Evolution in Word Usage: Just Add More Clusters?","date":"2020-01-18","arxiv_id":"2001.06629","n_code_links":0,"syntology":null},{"paper":"/paper/harmonic-convolutional-networks-based-on","slug":"harmonic-convolutional-networks-based-on","title":"Harmonic Convolutional Networks based on Discrete Cosine Transform","date":"2020-01-18","arxiv_id":"2001.06570","n_code_links":1,"syntology":{"ran":10,"of":10,"n_ran_checked":10,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"10 ran (of which 0 constructed an object rather than computing a result; 10 with no instrument failure: 0 honoured, 0 violated, 10 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["matej-ulicny/harmonic-networks"],"state":"official (archive's flag): 10 ran","n_ran":10,"n_constructed":0,"n_ran_no_instrument_failure":10,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"modality-balanced-models-for-visual-dialogue","title":"Modality-Balanced Models for Visual Dialogue","date":"2020-01-17","arxiv_id":"2001.06354","n_code_links":0,"syntology":null},{"paper":"/paper/robbert-a-dutch-roberta-based-language-model","slug":"robbert-a-dutch-roberta-based-language-model","title":"RobBERT: a Dutch RoBERTa-based Language Model","date":"2020-01-17","arxiv_id":"2001.06286","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["iPieter/RobBERT"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/delving-deeper-into-the-decoder-for-video","slug":"delving-deeper-into-the-decoder-for-video","title":"Delving Deeper into the Decoder for Video Captioning","date":"2020-01-16","arxiv_id":"2001.05614","n_code_links":1,"syntology":{"ran":0,"of":2,"n_ran_checked":0,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"0 ran · 2 unverified","official":{"repos":["WingsBrokenAngel/delving-deeper-into-the-decoder-for-video-captioning"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":2,"ran_from_kinds":[]}}},{"paper":"/paper/schema2qa-answering-complex-queries-on-the","slug":"schema2qa-answering-complex-queries-on-the","title":"Schema2QA: High-Quality and Low-Cost Q&A Agents for the Structured Web","date":"2020-01-16","arxiv_id":"2001.05609","n_code_links":3,"syntology":null},{"paper":null,"slug":"shifted-and-squeezed-8-bit-floating-point-1","title":"Shifted and Squeezed 8-bit Floating Point format for Low-Precision Training of Deep Neural Networks","date":"2020-01-16","arxiv_id":"2001.05674","n_code_links":0,"syntology":null},{"paper":"/paper/deep-residual-flow-for-novelty-detection","slug":"deep-residual-flow-for-novelty-detection","title":"Deep Residual Flow for Out of Distribution Detection","date":"2020-01-15","arxiv_id":"2001.05419","n_code_links":1,"syntology":{"ran":1,"of":2,"n_ran_checked":0,"n_instrument":1,"unverified":1,"pointer_only":2,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":{"repos":["EvZissel/Residual-Flow"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/fgn-fusion-glyph-network-for-chinese-named","slug":"fgn-fusion-glyph-network-for-chinese-named","title":"FGN: Fusion Glyph Network for Chinese Named Entity Recognition","date":"2020-01-15","arxiv_id":"2001.05272","n_code_links":1,"syntology":null},{"paper":null,"slug":"insertion-deletion-transformer","title":"Insertion-Deletion Transformer","date":"2020-01-15","arxiv_id":"2001.05540","n_code_links":0,"syntology":null},{"paper":"/paper/parallel-machine-translation-with","slug":"parallel-machine-translation-with","title":"Non-Autoregressive Machine Translation with Disentangled Context Transformer","date":"2020-01-15","arxiv_id":"2001.05136","n_code_links":1,"syntology":null},{"paper":null,"slug":"transformer-based-online-ctcattention-end-to","title":"Transformer-based Online CTC/attention End-to-End Speech Recognition Architecture","date":"2020-01-15","arxiv_id":"2001.08290","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-bert-based-sentiment-analysis-and-key","title":"A BERT based Sentiment Analysis and Key Entity Detection Approach for Online Financial Texts","date":"2020-01-14","arxiv_id":"2001.05326","n_code_links":0,"syntology":null},{"paper":null,"slug":"auto-completion-of-user-interface-layout-1","title":"Auto Completion of User Interface Layout Design Using Transformer-Based Tree Decoders","date":"2020-01-14","arxiv_id":"2001.05308","n_code_links":0,"syntology":null},{"paper":"/paper/quantisation-and-pruning-for-neural-network","slug":"quantisation-and-pruning-for-neural-network","title":"Quantisation and Pruning for Neural Network Compression and Regularisation","date":"2020-01-14","arxiv_id":"2001.04850","n_code_links":1,"syntology":null},{"paper":null,"slug":"the-problems-with-using-stns-to-align-cnn","title":"The problems with using STNs to align CNN feature maps","date":"2020-01-14","arxiv_id":"2001.05858","n_code_links":0,"syntology":null},{"paper":"/paper/adabert-task-adaptive-bert-compression-with","slug":"adabert-task-adaptive-bert-compression-with","title":"AdaBERT: Task-Adaptive BERT Compression with Differentiable Neural Architecture Search","date":"2020-01-13","arxiv_id":"2001.04246","n_code_links":1,"syntology":{"ran":0,"of":3,"n_ran_checked":0,"n_instrument":0,"unverified":3,"pointer_only":0,"phrase":"0 ran · 3 unverified","official":null}},{"paper":"/paper/reformer-the-efficient-transformer-1","slug":"reformer-the-efficient-transformer-1","title":"Reformer: The Efficient Transformer","date":"2020-01-13","arxiv_id":"2001.04451","n_code_links":10,"syntology":{"ran":6,"of":8,"n_ran_checked":1,"n_instrument":5,"unverified":2,"pointer_only":0,"phrase":"6 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 5 where Syntology's instrument failed) · 2 unverified","official":{"repos":["google/trax"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/representations-lexicales-pour-la-detection","slug":"representations-lexicales-pour-la-detection","title":"Représentations lexicales pour la détection non supervisée d'événements dans un flux de tweets : étude sur des corpus français et anglais","date":"2020-01-13","arxiv_id":"2001.04139","n_code_links":1,"syntology":null},{"paper":null,"slug":"urdu-english-machine-transliteration-using","title":"Urdu-English Machine Transliteration using Neural Networks","date":"2020-01-12","arxiv_id":"2001.05296","n_code_links":0,"syntology":null},{"paper":"/paper/exploring-and-improving-robustness-of-multi","slug":"exploring-and-improving-robustness-of-multi","title":"Exploring and Improving Robustness of Multi Task Deep Neural Networks via Domain Agnostic Defenses","date":"2020-01-11","arxiv_id":"2001.05286","n_code_links":1,"syntology":null},{"paper":null,"slug":"patenttransformer-2-controlling-patent-text","title":"PatentTransformer-2: Controlling Patent Text Generation by Structural Metadata","date":"2020-01-11","arxiv_id":"2001.03708","n_code_links":0,"syntology":null},{"paper":"/paper/resolving-the-scope-of-speculation-and","slug":"resolving-the-scope-of-speculation-and","title":"Resolving the Scope of Speculation and Negation using Transformer-Based Architectures","date":"2020-01-09","arxiv_id":"2001.02885","n_code_links":1,"syntology":null},{"paper":"/paper/spatial-temporal-transformer-networks-for","slug":"spatial-temporal-transformer-networks-for","title":"Spatial-Temporal Transformer Networks for Traffic Flow Forecasting","date":"2020-01-09","arxiv_id":"2001.02908","n_code_links":1,"syntology":null},{"paper":null,"slug":"streaming-automatic-speech-recognition-with","title":"Streaming automatic speech recognition with the transformer model","date":"2020-01-08","arxiv_id":"2001.02674","n_code_links":0,"syntology":null},{"paper":null,"slug":"to-transfer-or-not-to-transfer","title":"To Transfer or Not to Transfer: Misclassification Attacks Against Transfer Learned Text Classifiers","date":"2020-01-08","arxiv_id":"2001.02438","n_code_links":0,"syntology":null},{"paper":null,"slug":"recast-interactive-auditing-of-automatic","title":"RECAST: Interactive Auditing of Automatic Toxicity Detection Models","date":"2020-01-07","arxiv_id":"2001.01819","n_code_links":0,"syntology":null},{"paper":"/paper/sparse-weight-activation-training-1","slug":"sparse-weight-activation-training-1","title":"Sparse Weight Activation Training","date":"2020-01-07","arxiv_id":"2001.01969","n_code_links":1,"syntology":{"ran":3,"of":4,"n_ran_checked":3,"n_instrument":0,"unverified":1,"pointer_only":4,"phrase":"3 ran (of which 3 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified; every one of the 3 samples that ran constructed an object rather than computing a result","official":{"repos":["AamirRaihan/SWAT"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":3,"n_ran_no_instrument_failure":3,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"facial-emotions-recognition-using","title":"Facial Emotions Recognition using Convolutional Neural Net","date":"2020-01-06","arxiv_id":"2001.01456","n_code_links":0,"syntology":null},{"paper":"/paper/improving-entity-linking-by-modeling-latent-2","slug":"improving-entity-linking-by-modeling-latent-2","title":"Improving Entity Linking by Modeling Latent Entity Type Information","date":"2020-01-06","arxiv_id":"2001.01447","n_code_links":0,"syntology":null},{"paper":"/paper/fdftnet-facing-off-fake-images-using-fake","slug":"fdftnet-facing-off-fake-images-using-fake","title":"FDFtNet: Facing Off Fake Images using Fake Detection Fine-tuning Network","date":"2020-01-05","arxiv_id":"2001.01265","n_code_links":2,"syntology":{"ran":8,"of":13,"n_ran_checked":8,"n_instrument":0,"unverified":5,"pointer_only":0,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 0 violated, 8 with no contract checked; 0 where Syntology's instrument failed) · 5 unverified","official":{"repos":["cutz-j/FDFtNet"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":4,"ran_from_kinds":["listed","official"]}}},{"paper":null,"slug":"rpr-random-partition-relaxation-for-training","title":"RPR: Random Partition Relaxation for Training; Binary and Ternary Weight Neural Networks","date":"2020-01-04","arxiv_id":"2001.01091","n_code_links":0,"syntology":null},{"paper":null,"slug":"learning-accurate-integer-transformer-machine","title":"Learning Accurate Integer Transformer Machine-Translation Models","date":"2020-01-03","arxiv_id":"2001.00926","n_code_links":0,"syntology":null},{"paper":null,"slug":"multi-layer-content-interaction-through","title":"Multi-Layer Content Interaction Through Quaternion Product For Visual Question Answering","date":"2020-01-03","arxiv_id":"2001.05840","n_code_links":0,"syntology":null},{"paper":"/paper/two-level-transformer-and-auxiliary-coherence","slug":"two-level-transformer-and-auxiliary-coherence","title":"Two-Level Transformer and Auxiliary Coherence Modeling for Improved Text Segmentation","date":"2020-01-03","arxiv_id":"2001.00891","n_code_links":1,"syntology":null},{"paper":null,"slug":"representing-unordered-data-using-multiset-1","title":"Representing Unordered Data Using Complex-Weighted Multiset Automata","date":"2020-01-02","arxiv_id":"2001.00610","n_code_links":0,"syntology":null},{"paper":null,"slug":"robust-marine-buoy-placement-for-ship","title":"Robust Marine Buoy Placement for Ship Detection Using Dropout K-Means","date":"2020-01-02","arxiv_id":"2001.00564","n_code_links":0,"syntology":null},{"paper":null,"slug":"an-efficient-homotopy-training-algorithm-for","title":"AN EFFICIENT HOMOTOPY TRAINING ALGORITHM FOR NEURAL NETWORKS","date":"2020-01-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"analytical-moment-regularizer-for-training","title":"Analytical Moment Regularizer for Training Robust Networks","date":"2020-01-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"assessment-modeling-fundamental-pre-training","title":"Assessment Modeling: Fundamental Pre-training Tasks for Interactive Educational Systems","date":"2020-01-01","arxiv_id":"2002.05505","n_code_links":0,"syntology":null},{"paper":null,"slug":"bert-al-bert-for-arbitrarily-long-document","title":"BERT-AL: BERT for Arbitrarily Long Document Understanding","date":"2020-01-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"differentiable-architecture-compression","title":"Differentiable Architecture Compression","date":"2020-01-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"insights-on-visual-representations-for","title":"Insights on Visual Representations for Embodied Navigation Tasks","date":"2020-01-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"neural-execution-engines","title":"NEURAL EXECUTION ENGINES","date":"2020-01-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"out-of-distribution-detection-using-layerwise","title":"Out-of-Distribution Detection Using Layerwise Uncertainty in Deep Neural Networks","date":"2020-01-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"parallel-neural-text-to-speech-1","title":"Parallel Neural Text-to-Speech","date":"2020-01-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/poly-encoders-architectures-and-pre-training","slug":"poly-encoders-architectures-and-pre-training","title":"Poly-encoders: Architectures and Pre-training Strategies for Fast and Accurate Multi-sentence Scoring","date":"2020-01-01","arxiv_id":null,"n_code_links":2,"syntology":null},{"paper":"/paper/stacked-debert-all-attention-in-incomplete","slug":"stacked-debert-all-attention-in-incomplete","title":"Stacked DeBERT: All Attention in Incomplete Data for Text Classification","date":"2020-01-01","arxiv_id":"2001.00137","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":2,"n_instrument":1,"unverified":0,"pointer_only":1,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 2 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["gcunhase/StackedDeBERT"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official","unlocated"]}}},{"paper":null,"slug":"uw-net-an-inception-attention-network-for","title":"UW-NET: AN INCEPTION-ATTENTION NETWORK FOR UNDERWATER IMAGE CLASSIFICATION","date":"2020-01-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/zeroq-a-novel-zero-shot-quantization","slug":"zeroq-a-novel-zero-shot-quantization","title":"ZeroQ: A Novel Zero Shot Quantization Framework","date":"2020-01-01","arxiv_id":"2001.00281","n_code_links":3,"syntology":{"ran":7,"of":19,"n_ran_checked":4,"n_instrument":3,"unverified":12,"pointer_only":2,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 1 honoured, 0 violated, 3 with no contract checked; 3 where Syntology's instrument failed) · 12 unverified","official":{"repos":["amirgholami/ZeroQ"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":null,"slug":"deep-attentive-ranking-networks-for-learning","title":"Deep Attentive Ranking Networks for Learning to Order Sentences","date":"2019-12-31","arxiv_id":"2001.00056","n_code_links":0,"syntology":null},{"paper":null,"slug":"eeg-based-continuous-speech-recognition-using","title":"EEG based Continuous Speech Recognition using Transformers","date":"2019-12-31","arxiv_id":"2001.00501","n_code_links":0,"syntology":null},{"paper":"/paper/olmpics-on-what-language-model-pre-training","slug":"olmpics-on-what-language-model-pre-training","title":"oLMpics -- On what Language Model Pre-training Captures","date":"2019-12-31","arxiv_id":"1912.13283","n_code_links":2,"syntology":null},{"paper":"/paper/oteann-estimating-the-transparency-of","slug":"oteann-estimating-the-transparency-of","title":"OTEANN: Estimating the Transparency of Orthographies with an Artificial Neural Network","date":"2019-12-31","arxiv_id":"1912.13321","n_code_links":2,"syntology":null},{"paper":"/paper/aranet-a-deep-learning-toolkit-for-arabic","slug":"aranet-a-deep-learning-toolkit-for-arabic","title":"AraNet: A Deep Learning Toolkit for Arabic Social Media","date":"2019-12-30","arxiv_id":"1912.13072","n_code_links":1,"syntology":null},{"paper":"/paper/autodiscern-rating-the-quality-of-online","slug":"autodiscern-rating-the-quality-of-online","title":"AutoDiscern: Rating the Quality of Online Health Information with Hierarchical Encoder Attention-based Neural Networks","date":"2019-12-30","arxiv_id":"1912.12999","n_code_links":1,"syntology":null},{"paper":null,"slug":"searching-for-stage-wise-neural-graphs-in-the-1","title":"Searching for Stage-wise Neural Graphs In the Limit","date":"2019-12-30","arxiv_id":"1912.12860","n_code_links":0,"syntology":null},{"paper":null,"slug":"pipelined-training-with-stale-weights-of-deep-1","title":"Pipelined Training with Stale Weights of Deep Convolutional Neural Networks","date":"2019-12-29","arxiv_id":"1912.12675","n_code_links":0,"syntology":null},{"paper":null,"slug":"all-in-one-image-grounded-conversational","title":"All-in-One Image-Grounded Conversational Agents","date":"2019-12-28","arxiv_id":"1912.12394","n_code_links":0,"syntology":null},{"paper":null,"slug":"natural-language-processing-of-mimic-iii","title":"Natural language processing of MIMIC-III clinical notes for identifying diagnosis and procedures with neural networks","date":"2019-12-28","arxiv_id":"1912.12397","n_code_links":0,"syntology":null},{"paper":"/paper/clinical-xlnet-modeling-sequential-clinical","slug":"clinical-xlnet-modeling-sequential-clinical","title":"Clinical XLNet: Modeling Sequential Clinical Notes and Predicting Prolonged Mechanical Ventilation","date":"2019-12-27","arxiv_id":"1912.11975","n_code_links":3,"syntology":null}],"record_sha256":"be14bff3dc730b79ba86baa66d6c36034ca3e1a4c9e6270c45a88b797035be33","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}