{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/position-wise-feed-forward-layer/papers/128","list_of":"/method/position-wise-feed-forward-layer","method":"Position-Wise Feed-Forward Layer","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":128,"pages_in_order":139,"rows_per_page":100,"rows":[12701,12800],"of":13895,"counts":{"archive_papers_tagged":13895,"with_a_code_link":6514,"where_syntology_ran_a_sample":2229,"not_listed_spam_title":0,"listed":13895,"listed_where_code_ran":2229,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":1902,"every_run_a_failure_of_syntologys_instrument":327,"listed_with_a_run_with_no_instrument_failure":1902,"listed_every_run_a_failure_of_syntologys_instrument":327,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/position-wise-feed-forward-layer","prev":"/method/position-wise-feed-forward-layer/papers/127","next":"/method/position-wise-feed-forward-layer/papers/129","papers":[{"paper":null,"slug":"towards-fully-8-bit-integer-inference-for-the","title":"Towards Fully 8-bit Integer Inference for the Transformer Model","date":"2020-09-17","arxiv_id":"2009.08034","n_code_links":0,"syntology":null},{"paper":"/paper/automated-source-code-generation-and-auto","slug":"automated-source-code-generation-and-auto","title":"Automated Source Code Generation and Auto-completion Using Deep Learning: Comparing and Discussing Current Language-Model-Related Approaches","date":"2020-09-16","arxiv_id":"2009.07740","n_code_links":1,"syntology":null},{"paper":"/paper/cogtree-cognition-tree-loss-for-unbiased","slug":"cogtree-cognition-tree-loss-for-unbiased","title":"CogTree: Cognition Tree Loss for Unbiased Scene Graph Generation","date":"2020-09-16","arxiv_id":"2009.07526","n_code_links":1,"syntology":null},{"paper":null,"slug":"document-level-neural-machine-translation-1","title":"Document-level Neural Machine Translation with Document Embeddings","date":"2020-09-16","arxiv_id":"2009.08775","n_code_links":0,"syntology":null},{"paper":null,"slug":"extremely-low-bit-transformer-quantization","title":"Extremely Low Bit Transformer Quantization for On-Device Neural Machine Translation","date":"2020-09-16","arxiv_id":"2009.07453","n_code_links":0,"syntology":null},{"paper":null,"slug":"graph-to-sequence-neural-machine-translation","title":"Graph-to-Sequence Neural Machine Translation","date":"2020-09-16","arxiv_id":"2009.07489","n_code_links":0,"syntology":null},{"paper":null,"slug":"nabu-multilingual-graph-based-neural-rdf","title":"NABU $\\mathrm{-}$ Multilingual Graph-based Neural RDF Verbalizer","date":"2020-09-16","arxiv_id":"2009.07728","n_code_links":0,"syntology":null},{"paper":null,"slug":"retrofitting-structure-aware-transformer","title":"Retrofitting Structure-aware Transformer Language Model for End Tasks","date":"2020-09-16","arxiv_id":"2009.07408","n_code_links":0,"syntology":null},{"paper":"/paper/attention-aware-inference-for-neural","slug":"attention-aware-inference-for-neural","title":"Global-aware Beam Search for Neural Abstractive Summarization","date":"2020-09-15","arxiv_id":"2009.06891","n_code_links":2,"syntology":null},{"paper":null,"slug":"event-presence-prediction-helps-trigger","title":"Event Presence Prediction Helps Trigger Detection Across Languages","date":"2020-09-15","arxiv_id":"2009.07188","n_code_links":0,"syntology":null},{"paper":null,"slug":"efficient-transformers-a-survey","title":"Efficient Transformers: A Survey","date":"2020-09-14","arxiv_id":"2009.06732","n_code_links":0,"syntology":null},{"paper":null,"slug":"boostingbert-integrating-multi-class-boosting","title":"BoostingBERT:Integrating Multi-Class Boosting into BERT for NLP Tasks","date":"2020-09-13","arxiv_id":"2009.05959","n_code_links":0,"syntology":null},{"paper":"/paper/gtea-representation-learning-for-temporal","slug":"gtea-representation-learning-for-temporal","title":"GTEA: Inductive Representation Learning on Temporal Interaction Graphs via Temporal Edge Aggregation","date":"2020-09-11","arxiv_id":"2009.05266","n_code_links":2,"syntology":null},{"paper":null,"slug":"learning-universal-representations-from-word","title":"Learning Universal Representations from Word to Sentence","date":"2020-09-10","arxiv_id":"2009.04656","n_code_links":0,"syntology":null},{"paper":"/paper/rank-over-class-the-untapped-potential-of","slug":"rank-over-class-the-untapped-potential-of","title":"Rank over Class: The Untapped Potential of Ranking in Natural Language Processing","date":"2020-09-10","arxiv_id":"2009.05160","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":2,"n_instrument":1,"unverified":0,"pointer_only":1,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["atapour/rank-over-class"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/sparsifying-transformer-models-with","slug":"sparsifying-transformer-models-with","title":"Sparsifying Transformer Models with Trainable Representation Pooling","date":"2020-09-10","arxiv_id":"2009.05169","n_code_links":1,"syntology":null},{"paper":"/paper/pay-attention-when-required","slug":"pay-attention-when-required","title":"Pay Attention when Required","date":"2020-09-09","arxiv_id":"2009.04534","n_code_links":2,"syntology":null},{"paper":"/paper/masked-label-prediction-unified-massage","slug":"masked-label-prediction-unified-massage","title":"Masked Label Prediction: Unified Message Passing Model for Semi-Supervised Classification","date":"2020-09-08","arxiv_id":"2009.03509","n_code_links":3,"syntology":{"ran":7,"of":7,"n_ran_checked":6,"n_instrument":1,"unverified":0,"pointer_only":2,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 2 violated, 4 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["PaddlePaddle/PGL"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["listed","official"]}}},{"paper":"/paper/adversarial-watermarking-transformer-towards","slug":"adversarial-watermarking-transformer-towards","title":"Adversarial Watermarking Transformer: Towards Tracing Text Provenance with Data Hiding","date":"2020-09-07","arxiv_id":"2009.03015","n_code_links":1,"syntology":{"ran":2,"of":4,"n_ran_checked":1,"n_instrument":1,"unverified":2,"pointer_only":4,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","official":null}},{"paper":null,"slug":"robust-conversational-ai-with-grounded-text","title":"Robust Conversational AI with Grounded Text Generation","date":"2020-09-07","arxiv_id":"2009.03457","n_code_links":0,"syntology":null},{"paper":"/paper/transmodality-an-end2end-fusion-method-with","slug":"transmodality-an-end2end-fusion-method-with","title":"TransModality: An End2End Fusion Method with Transformer for Multimodal Sentiment Analysis","date":"2020-09-07","arxiv_id":"2009.02902","n_code_links":0,"syntology":null},{"paper":null,"slug":"voice-conversion-by-cascading-automatic","title":"Voice Conversion by Cascading Automatic Speech Recognition and Text-to-Speech Synthesis with Prosody Transfer","date":"2020-09-03","arxiv_id":"2009.01475","n_code_links":0,"syntology":null},{"paper":"/paper/liftformer-3d-human-pose-estimation-using","slug":"liftformer-3d-human-pose-estimation-using","title":"LiftFormer: 3D Human Pose Estimation using attention models","date":"2020-09-01","arxiv_id":"2009.00348","n_code_links":0,"syntology":null},{"paper":null,"slug":"parallel-rescoring-with-transformer-for","title":"Parallel Rescoring with Transformer for Streaming On-Device Speech Recognition","date":"2020-08-30","arxiv_id":"2008.13093","n_code_links":0,"syntology":null},{"paper":"/paper/hitter-hierarchical-transformers-for","slug":"hitter-hierarchical-transformers-for","title":"HittER: Hierarchical Transformers for Knowledge Graph Embeddings","date":"2020-08-28","arxiv_id":"2008.12813","n_code_links":3,"syntology":{"ran":7,"of":7,"n_ran_checked":7,"n_instrument":0,"unverified":0,"pointer_only":1,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":null}},{"paper":null,"slug":"tatl-at-w-nut-2020-task-2-a-transformer-based","title":"TATL at W-NUT 2020 Task 2: A Transformer-based Baseline System for Identification of Informative COVID-19 English Tweets","date":"2020-08-28","arxiv_id":"2008.12854","n_code_links":0,"syntology":null},{"paper":null,"slug":"text-conditioned-transformer-for-automatic","title":"Text-Conditioned Transformer for Automatic Pronunciation Error Detection","date":"2020-08-28","arxiv_id":"2008.12424","n_code_links":0,"syntology":null},{"paper":"/paper/improvement-of-a-dedicated-model-for-open","slug":"improvement-of-a-dedicated-model-for-open","title":"Improvement of a dedicated model for open domain persona-aware dialogue generation","date":"2020-08-27","arxiv_id":"2008.11970","n_code_links":1,"syntology":null},{"paper":null,"slug":"end-to-end-dialogue-transformer","title":"End to End Dialogue Transformer","date":"2020-08-24","arxiv_id":"2008.10392","n_code_links":0,"syntology":null},{"paper":"/paper/identity-aware-multi-sentence-video","slug":"identity-aware-multi-sentence-video","title":"Identity-Aware Multi-Sentence Video Description","date":"2020-08-22","arxiv_id":"2008.09791","n_code_links":1,"syntology":null},{"paper":"/paper/are-neural-open-domain-dialog-systems-robust","slug":"are-neural-open-domain-dialog-systems-robust","title":"Are Neural Open-Domain Dialog Systems Robust to Speech Recognition Errors in the Dialog History? An Empirical Study","date":"2020-08-18","arxiv_id":"2008.07683","n_code_links":1,"syntology":null},{"paper":"/paper/glancing-transformer-for-non-autoregressive","slug":"glancing-transformer-for-non-autoregressive","title":"Glancing Transformer for Non-Autoregressive Neural Machine Translation","date":"2020-08-18","arxiv_id":"2008.07905","n_code_links":2,"syntology":null},{"paper":"/paper/very-deep-transformers-for-neural-machine","slug":"very-deep-transformers-for-neural-machine","title":"Very Deep Transformers for Neural Machine Translation","date":"2020-08-18","arxiv_id":"2008.07772","n_code_links":4,"syntology":{"ran":8,"of":9,"n_ran_checked":3,"n_instrument":5,"unverified":1,"pointer_only":3,"phrase":"8 ran (of which 2 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 5 where Syntology's instrument failed) · 1 unverified","official":{"repos":["namisan/exdeep-nmt"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/spatial-temporal-transformer-network-for","slug":"spatial-temporal-transformer-network-for","title":"Skeleton-based Action Recognition via Spatial and Temporal Transformer Networks","date":"2020-08-17","arxiv_id":"2008.07404","n_code_links":1,"syntology":null},{"paper":null,"slug":"dcr-net-a-deep-co-interactive-relation","title":"DCR-Net: A Deep Co-Interactive Relation Network for Joint Dialog Act Recognition and Sentiment Classification","date":"2020-08-16","arxiv_id":"2008.06914","n_code_links":0,"syntology":null},{"paper":null,"slug":"topicbert-a-transformer-transfer-learning","title":"TopicBERT: A Transformer transfer learning based memory-graph approach for multimodal streaming social media topic detection","date":"2020-08-16","arxiv_id":"2008.06877","n_code_links":0,"syntology":null},{"paper":null,"slug":"finding-fast-transformers-one-shot-neural","title":"Finding Fast Transformers: One-Shot Neural Architecture Search by Component Composition","date":"2020-08-15","arxiv_id":"2008.06808","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-hybrid-bert-and-lightgbm-based-model-for","title":"A Hybrid BERT and LightGBM based Model for Predicting Emotion GIF Categories on Twitter","date":"2020-08-14","arxiv_id":"2008.06176","n_code_links":0,"syntology":null},{"paper":null,"slug":"adaptable-multi-domain-language-model-for","title":"Adaptable Multi-Domain Language Model for Transformer ASR","date":"2020-08-14","arxiv_id":"2008.06208","n_code_links":0,"syntology":null},{"paper":"/paper/a-community-powered-search-of-machine","slug":"a-community-powered-search-of-machine","title":"A community-powered search of machine learning strategy space to find NMR property prediction models","date":"2020-08-13","arxiv_id":"2008.05994","n_code_links":1,"syntology":null},{"paper":null,"slug":"conv-transformer-transducer-low-latency-low","title":"Conv-Transformer Transducer: Low Latency, Low Frame Rate, Streamable End-to-End Speech Recognition","date":"2020-08-13","arxiv_id":"2008.05750","n_code_links":0,"syntology":null},{"paper":null,"slug":"end-to-end-contextual-perception-and","title":"End-to-end Contextual Perception and Prediction with Interaction Transformer","date":"2020-08-13","arxiv_id":"2008.05927","n_code_links":0,"syntology":null},{"paper":null,"slug":"large-scale-transfer-learning-for-low","title":"Large-scale Transfer Learning for Low-resource Spoken Language Understanding","date":"2020-08-13","arxiv_id":"2008.05671","n_code_links":0,"syntology":null},{"paper":"/paper/mmm-exploring-conditional-multi-track-music","slug":"mmm-exploring-conditional-multi-track-music","title":"MMM : Exploring Conditional Multi-Track Music Generation with the Transformer","date":"2020-08-13","arxiv_id":"2008.06048","n_code_links":3,"syntology":{"ran":9,"of":9,"n_ran_checked":9,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 9 with no instrument failure: 0 honoured, 0 violated, 9 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":null}},{"paper":null,"slug":"compression-of-deep-learning-models-for-text","title":"Compression of Deep Learning Models for Text: A Survey","date":"2020-08-12","arxiv_id":"2008.05221","n_code_links":0,"syntology":null},{"paper":"/paper/evaluating-the-impact-of-knowledge-graph","slug":"evaluating-the-impact-of-knowledge-graph","title":"Evaluating the Impact of Knowledge Graph Context on Entity Disambiguation Models","date":"2020-08-12","arxiv_id":"2008.05190","n_code_links":1,"syntology":null},{"paper":"/paper/fine-grained-visual-textual-alignment-for","slug":"fine-grained-visual-textual-alignment-for","title":"Fine-grained Visual Textual Alignment for Cross-Modal Retrieval using Transformer Encoders","date":"2020-08-12","arxiv_id":"2008.05231","n_code_links":1,"syntology":{"ran":14,"of":16,"n_ran_checked":13,"n_instrument":1,"unverified":2,"pointer_only":2,"phrase":"14 ran (of which 0 constructed an object rather than computing a result; 13 with no instrument failure: 0 honoured, 1 violated, 12 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","official":{"repos":["mesnico/TERAN"],"state":"official (archive's flag): 14 ran","n_ran":14,"n_constructed":0,"n_ran_no_instrument_failure":13,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/pretraining-techniques-for-sequence-to","slug":"pretraining-techniques-for-sequence-to","title":"Pretraining Techniques for Sequence-to-Sequence Voice Conversion","date":"2020-08-07","arxiv_id":"2008.03088","n_code_links":2,"syntology":null},{"paper":null,"slug":"6veclm-language-modeling-in-vector-space-for","title":"6VecLM: Language Modeling in Vector Space for IPv6 Target Generation","date":"2020-08-05","arxiv_id":"2008.02213","n_code_links":0,"syntology":null},{"paper":"/paper/designing-the-business-conversation-corpus-1","slug":"designing-the-business-conversation-corpus-1","title":"Designing the Business Conversation Corpus","date":"2020-08-05","arxiv_id":"2008.01940","n_code_links":1,"syntology":null},{"paper":null,"slug":"select-extract-and-generate-neural-keyphrase","title":"Select, Extract and Generate: Neural Keyphrase Generation with Layer-wise Coverage Attention","date":"2020-08-04","arxiv_id":"2008.01739","n_code_links":0,"syntology":null},{"paper":"/paper/the-jazz-transformer-on-the-front-line","slug":"the-jazz-transformer-on-the-front-line","title":"The Jazz Transformer on the Front Line: Exploring the Shortcomings of AI-composed Music through Quantitative Measures","date":"2020-08-04","arxiv_id":"2008.01307","n_code_links":2,"syntology":{"ran":13,"of":14,"n_ran_checked":13,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"13 ran (of which 0 constructed an object rather than computing a result; 13 with no instrument failure: 0 honoured, 0 violated, 13 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["slSeanWU/MusDr","slSeanWU/jazz_transformer"],"state":"official (archive's flag): 13 ran","n_ran":13,"n_constructed":0,"n_ran_no_instrument_failure":13,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"lt-helsinki-at-semeval-2020-task-12","title":"LT@Helsinki at SemEval-2020 Task 12: Multilingual or language-specific BERT?","date":"2020-08-03","arxiv_id":"2008.00805","n_code_links":0,"syntology":null},{"paper":null,"slug":"self-attention-encoding-and-pooling-for","title":"Self-attention encoding and pooling for speaker recognition","date":"2020-08-03","arxiv_id":"2008.01077","n_code_links":0,"syntology":null},{"paper":"/paper/seqdialn-sequential-visual-dialog-networks-in","slug":"seqdialn-sequential-visual-dialog-networks-in","title":"SeqDialN: Sequential Visual Dialog Networks in Joint Visual-Linguistic Representation Space","date":"2020-08-02","arxiv_id":"2008.00397","n_code_links":1,"syntology":null},{"paper":"/paper/the-chess-transformer-mastering-play-using","slug":"the-chess-transformer-mastering-play-using","title":"The Chess Transformer: Mastering Play using Generative Language Models","date":"2020-08-02","arxiv_id":"2008.04057","n_code_links":2,"syntology":{"ran":1,"of":2,"n_ran_checked":0,"n_instrument":1,"unverified":1,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":null}},{"paper":"/paper/stochastic-fine-grained-labeling-of-multi","slug":"stochastic-fine-grained-labeling-of-multi","title":"Stochastic Fine-grained Labeling of Multi-state Sign Glosses for Continuous Sign Language Recognition","date":"2020-08-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"deep-multi-view-spatiotemporal-virtual-graph","title":"Deep Multi-View Spatiotemporal Virtual Graph Neural Network for Significant Citywide Ride-hailing Demand Prediction","date":"2020-07-30","arxiv_id":"2007.15189","n_code_links":0,"syntology":null},{"paper":"/paper/interpretable-contextual-team-aware-item","slug":"interpretable-contextual-team-aware-item","title":"Interpretable Contextual Team-aware Item Recommendation: Application in Multiplayer Online Battle Arena Games","date":"2020-07-30","arxiv_id":"2007.15236","n_code_links":1,"syntology":null},{"paper":null,"slug":"tensorcoder-dimension-wise-attention-via","title":"TensorCoder: Dimension-Wise Attention via Tensor Representation for Natural Language Modeling","date":"2020-07-28","arxiv_id":"2008.01547","n_code_links":0,"syntology":null},{"paper":null,"slug":"to-bert-or-not-to-bert-comparing-speech-and","title":"To BERT or Not To BERT: Comparing Speech and Language-based Approaches for Alzheimer's Disease Detection","date":"2020-07-26","arxiv_id":"2008.01551","n_code_links":0,"syntology":null},{"paper":"/paper/fissa-at-semeval-2020-task-9-fine-tuned-for","slug":"fissa-at-semeval-2020-task-9-fine-tuned-for","title":"FiSSA at SemEval-2020 Task 9: Fine-tuned For Feelings","date":"2020-07-24","arxiv_id":"2007.12544","n_code_links":1,"syntology":null},{"paper":"/paper/exploring-swedish-english-fasttext-embeddings","slug":"exploring-swedish-english-fasttext-embeddings","title":"Exploring Swedish & English fastText Embeddings for NER with the Transformer","date":"2020-07-23","arxiv_id":"2007.16007","n_code_links":1,"syntology":null},{"paper":null,"slug":"analogical-reasoning-for-visually-grounded","title":"Analogical Reasoning for Visually Grounded Language Acquisition","date":"2020-07-22","arxiv_id":"2007.11668","n_code_links":0,"syntology":null},{"paper":"/paper/crosstransformers-spatially-aware-few-shot","slug":"crosstransformers-spatially-aware-few-shot","title":"CrossTransformers: spatially-aware few-shot transfer","date":"2020-07-22","arxiv_id":"2007.11498","n_code_links":6,"syntology":{"ran":5,"of":6,"n_ran_checked":5,"n_instrument":0,"unverified":1,"pointer_only":2,"phrase":"5 ran (of which 2 constructed an object rather than computing a result; 5 with no instrument failure: 2 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["google-research/meta-dataset"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["community","listed"]}}},{"paper":"/paper/deepsvg-a-hierarchical-generative-network-for","slug":"deepsvg-a-hierarchical-generative-network-for","title":"DeepSVG: A Hierarchical Generative Network for Vector Graphics Animation","date":"2020-07-22","arxiv_id":"2007.11301","n_code_links":2,"syntology":{"ran":5,"of":25,"n_ran_checked":4,"n_instrument":1,"unverified":20,"pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 1 violated, 3 with no contract checked; 1 where Syntology's instrument failed) · 20 unverified","official":{"repos":["alexandre01/deepsvg"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":20,"ran_from_kinds":["official"]}}},{"paper":"/paper/neural-machine-translation-with-error","slug":"neural-machine-translation-with-error","title":"Neural Machine Translation with Error Correction","date":"2020-07-21","arxiv_id":"2007.10681","n_code_links":1,"syntology":null},{"paper":null,"slug":"sliceout-training-transformers-and-cnns","title":"Improving compute efficacy frontiers with SliceOut","date":"2020-07-21","arxiv_id":"2007.10909","n_code_links":0,"syntology":null},{"paper":"/paper/conformer-kernel-with-query-term-independence","slug":"conformer-kernel-with-query-term-independence","title":"Conformer-Kernel with Query Term Independence for Document Retrieval","date":"2020-07-20","arxiv_id":"2007.10434","n_code_links":1,"syntology":null},{"paper":"/paper/learning-joint-spatial-temporal","slug":"learning-joint-spatial-temporal","title":"Learning Joint Spatial-Temporal Transformations for Video Inpainting","date":"2020-07-20","arxiv_id":"2007.10247","n_code_links":2,"syntology":{"ran":3,"of":7,"n_ran_checked":2,"n_instrument":1,"unverified":4,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 1 where Syntology's instrument failed) · 4 unverified","official":{"repos":["researchmm/STTN"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"feature-pyramid-transformer","title":"Feature Pyramid Transformer","date":"2020-07-18","arxiv_id":"2007.09451","n_code_links":0,"syntology":null},{"paper":"/paper/temporal-pointwise-convolutional-networks-for","slug":"temporal-pointwise-convolutional-networks-for","title":"Temporal Pointwise Convolutional Networks for Length of Stay Prediction in the Intensive Care Unit","date":"2020-07-18","arxiv_id":"2007.09483","n_code_links":1,"syntology":null},{"paper":null,"slug":"deep-learning-based-traffic-surveillance","title":"Deep Learning Based Traffic Surveillance System For Missing and Suspicious Car Detection","date":"2020-07-17","arxiv_id":"2007.08783","n_code_links":0,"syntology":null},{"paper":"/paper/hopfield-networks-is-all-you-need","slug":"hopfield-networks-is-all-you-need","title":"Hopfield Networks is All You Need","date":"2020-07-16","arxiv_id":"2008.02217","n_code_links":3,"syntology":{"ran":8,"of":9,"n_ran_checked":8,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 0 violated, 8 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["ml-jku/hopfield-layers"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/infoxlm-an-information-theoretic-framework","slug":"infoxlm-an-information-theoretic-framework","title":"InfoXLM: An Information-Theoretic Framework for Cross-Lingual Language Model Pre-Training","date":"2020-07-15","arxiv_id":"2007.07834","n_code_links":4,"syntology":null},{"paper":null,"slug":"the-monte-carlo-transformer-a-stochastic-self","title":"The Monte Carlo Transformer: a stochastic self-attention model for sequence prediction","date":"2020-07-15","arxiv_id":"2007.08620","n_code_links":0,"syntology":null},{"paper":"/paper/contextualized-code-representation-learning","slug":"contextualized-code-representation-learning","title":"CoreGen: Contextualized Code Representation Learning for Commit Message Generation","date":"2020-07-14","arxiv_id":"2007.06934","n_code_links":1,"syntology":null},{"paper":null,"slug":"deep-transformer-based-data-augmentation-with","title":"Deep Transformer based Data Augmentation with Subword Units for Morphologically Rich Online ASR","date":"2020-07-14","arxiv_id":"2007.06949","n_code_links":0,"syntology":null},{"paper":"/paper/emoji-prediction-extensions-and-benchmarking","slug":"emoji-prediction-extensions-and-benchmarking","title":"Emoji Prediction: Extensions and Benchmarking","date":"2020-07-14","arxiv_id":"2007.07389","n_code_links":1,"syntology":null},{"paper":"/paper/paranoid-transformer-reading-narrative-of","slug":"paranoid-transformer-reading-narrative-of","title":"Paranoid Transformer: Reading Narrative of Madness as Computational Approach to Creativity","date":"2020-07-13","arxiv_id":"2007.06290","n_code_links":1,"syntology":null},{"paper":null,"slug":"transformer-with-depth-wise-lstm","title":"Rewiring the Transformer with Depth-Wise LSTMs","date":"2020-07-13","arxiv_id":"2007.06257","n_code_links":0,"syntology":null},{"paper":null,"slug":"sparse-graph-to-sequence-learning-for-vision","title":"Sparse Graph to Sequence Learning for Vision Conditioned Long Textual Sequence Generation","date":"2020-07-12","arxiv_id":"2007.06077","n_code_links":0,"syntology":null},{"paper":"/paper/tera-self-supervised-learning-of-transformer","slug":"tera-self-supervised-learning-of-transformer","title":"TERA: Self-Supervised Learning of Transformer Encoder Representation for Speech","date":"2020-07-12","arxiv_id":"2007.06028","n_code_links":7,"syntology":{"ran":11,"of":14,"n_ran_checked":7,"n_instrument":4,"unverified":3,"pointer_only":3,"phrase":"11 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 1 honoured, 0 violated, 6 with no contract checked; 4 where Syntology's instrument failed) · 3 unverified","official":{"repos":["andi611/Self-Supervised-Speech-Pretraining-and-Representation-Learning","s3prl/s3prl"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/sequence-generation-with-mixed","slug":"sequence-generation-with-mixed","title":"Sequence Generation with Mixed Representations","date":"2020-07-11","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/bison-bm25-weighted-self-attention-framework","slug":"bison-bm25-weighted-self-attention-framework","title":"GLOW : Global Weighted Self-Attention Network for Web Search","date":"2020-07-10","arxiv_id":"2007.05186","n_code_links":1,"syntology":null},{"paper":"/paper/advances-of-transformer-based-models-for-news","slug":"advances-of-transformer-based-models-for-news","title":"Advances of Transformer-Based Models for News Headline Generation","date":"2020-07-09","arxiv_id":"2007.05044","n_code_links":2,"syntology":null},{"paper":null,"slug":"deepsinger-singing-voice-synthesis-with-data","title":"DeepSinger: Singing Voice Synthesis with Data Mined From the Web","date":"2020-07-09","arxiv_id":"2007.04590","n_code_links":0,"syntology":null},{"paper":null,"slug":"spatio-temporal-scene-graphs-for-video-dialog","title":"Dynamic Graph Representation Learning for Video Dialog via Multi-Modal Shuffled Transformers","date":"2020-07-08","arxiv_id":"2007.03848","n_code_links":0,"syntology":null},{"paper":"/paper/do-transformers-need-deep-long-range-memory-1","slug":"do-transformers-need-deep-long-range-memory-1","title":"Do Transformers Need Deep Long-Range Memory","date":"2020-07-07","arxiv_id":"2007.03356","n_code_links":1,"syntology":null},{"paper":null,"slug":"the-go-transformer-natural-language-modeling","title":"The Go Transformer: Natural Language Modeling for Game Play","date":"2020-07-07","arxiv_id":"2007.03500","n_code_links":0,"syntology":null},{"paper":null,"slug":"learning-to-segment-anatomical-structures","title":"Learning to Segment Anatomical Structures Accurately from One Exemplar","date":"2020-07-06","arxiv_id":"2007.03052","n_code_links":0,"syntology":null},{"paper":null,"slug":"relevance-transformer-generating-concise-code","title":"Relevance Transformer: Generating Concise Code Snippets with Relevance Feedback","date":"2020-07-06","arxiv_id":"2007.02609","n_code_links":0,"syntology":null},{"paper":null,"slug":"abstractive-and-mixed-summarization-for-long","title":"Abstractive and mixed summarization for long-single documents","date":"2020-07-03","arxiv_id":"2007.01918","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-transformer-approach-to-contextual-sarcasm","title":"A Transformer Approach to Contextual Sarcasm Detection in Twitter","date":"2020-07-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"adaptation-of-multilingual-transformer","title":"Adaptation of Multilingual Transformer Encoder for Robust Enhanced Universal Dependency Parsing","date":"2020-07-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"addressing-posterior-collapse-with-mutual","title":"Addressing Posterior Collapse with Mutual Information for Improved Variational Neural Machine Translation","date":"2020-07-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"an-empirical-investigation-of-neural-methods","title":"An empirical investigation of neural methods for content scoring of science explanations","date":"2020-07-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"casia-s-system-for-iwslt-2020-open-domain","title":"CASIA's System for IWSLT 2020 Open Domain Translation","date":"2020-07-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"character-aware-models-with-similarity","title":"Character aware models with similarity learning for metaphor detection","date":"2020-07-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"combining-subword-representations-into-word","title":"Combining Subword Representations into Word-level Representations in the Transformer Architecture","date":"2020-07-01","arxiv_id":null,"n_code_links":0,"syntology":null}],"record_sha256":"70841b735246d8ac9cb62e689cb69d11ca7254fea771fa28087d8c1fea22aea3","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}