{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/dense-connections/papers/213","list_of":"/method/dense-connections","method":"Dense Connections","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":213,"pages_in_order":293,"rows_per_page":100,"rows":[21201,21300],"of":29230,"counts":{"archive_papers_tagged":29230,"with_a_code_link":12972,"where_syntology_ran_a_sample":3929,"not_listed_spam_title":0,"listed":29230,"listed_where_code_ran":3929,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":3303,"every_run_a_failure_of_syntologys_instrument":626,"listed_with_a_run_with_no_instrument_failure":3303,"listed_every_run_a_failure_of_syntologys_instrument":626,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/dense-connections","prev":"/method/dense-connections/papers/212","next":"/method/dense-connections/papers/214","papers":[{"paper":"/paper/symbolic-knowledge-distillation-from-general","slug":"symbolic-knowledge-distillation-from-general","title":"Symbolic Knowledge Distillation: from General Language Models to Commonsense Models","date":"2021-10-14","arxiv_id":"2110.07178","n_code_links":1,"syntology":null},{"paper":"/paper/the-neural-data-router-adaptive-control-flow","slug":"the-neural-data-router-adaptive-control-flow","title":"The Neural Data Router: Adaptive Control Flow in Transformers Improves Systematic Generalization","date":"2021-10-14","arxiv_id":"2110.07732","n_code_links":1,"syntology":{"ran":7,"of":11,"n_ran_checked":6,"n_instrument":1,"unverified":4,"pointer_only":0,"phrase":"7 ran (of which 5 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 1 violated, 5 with no contract checked; 1 where Syntology's instrument failed) · 4 unverified","official":{"repos":["robertcsordas/ndr"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":5,"n_ran_no_instrument_failure":6,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"transformer-for-polyp-detection","title":"Transformer for Polyp Detection","date":"2021-10-14","arxiv_id":"2111.07918","n_code_links":0,"syntology":null},{"paper":null,"slug":"clip4caption-clip-for-video-caption","title":"CLIP4Caption: CLIP for Video Caption","date":"2021-10-13","arxiv_id":"2110.06615","n_code_links":0,"syntology":null},{"paper":null,"slug":"covert-message-passing-over-public-internet","title":"Leveraging Generative Models for Covert Messaging: Challenges and Tradeoffs for \"Dead-Drop\" Deployments","date":"2021-10-13","arxiv_id":"2110.07009","n_code_links":0,"syntology":null},{"paper":"/paper/decoupled-contrastive-learning-1","slug":"decoupled-contrastive-learning-1","title":"Decoupled Contrastive Learning","date":"2021-10-13","arxiv_id":"2110.06848","n_code_links":4,"syntology":{"ran":6,"of":6,"n_ran_checked":6,"n_instrument":0,"unverified":0,"pointer_only":3,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 1 honoured, 2 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":null}},{"paper":"/paper/fake-news-detection-in-spanish-using-deep","slug":"fake-news-detection-in-spanish-using-deep","title":"Fake News Detection in Spanish Using Deep Learning Techniques","date":"2021-10-13","arxiv_id":"2110.06461","n_code_links":1,"syntology":null},{"paper":"/paper/improving-graph-based-sentence-ordering-with","slug":"improving-graph-based-sentence-ordering-with","title":"Improving Graph-based Sentence Ordering with Iteratively Predicted Pairwise Orderings","date":"2021-10-13","arxiv_id":"2110.06446","n_code_links":1,"syntology":{"ran":2,"of":3,"n_ran_checked":2,"n_instrument":0,"unverified":1,"pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["deeplearnxmu/irseg"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"language-modelling-via-learning-to-rank","title":"Language Modelling via Learning to Rank","date":"2021-10-13","arxiv_id":"2110.06961","n_code_links":0,"syntology":null},{"paper":"/paper/leveraging-redundancy-in-attention-with-reuse-1","slug":"leveraging-redundancy-in-attention-with-reuse-1","title":"Leveraging redundancy in attention with Reuse Transformers","date":"2021-10-13","arxiv_id":"2110.06821","n_code_links":1,"syntology":null},{"paper":null,"slug":"maximizing-efficiency-of-language-model-pre","title":"Maximizing Efficiency of Language Model Pre-training for Learning Representation","date":"2021-10-13","arxiv_id":"2110.06620","n_code_links":0,"syntology":null},{"paper":"/paper/mderank-a-masked-document-embedding-rank","slug":"mderank-a-masked-document-embedding-rank","title":"MDERank: A Masked Document Embedding Rank Approach for Unsupervised Keyphrase Extraction","date":"2021-10-13","arxiv_id":"2110.06651","n_code_links":1,"syntology":null},{"paper":null,"slug":"multistage-linguistic-conditioning-of","title":"Multistage linguistic conditioning of convolutional layers for speech emotion recognition","date":"2021-10-13","arxiv_id":"2110.06650","n_code_links":0,"syntology":null},{"paper":null,"slug":"scaling-laws-for-the-few-shot-adaptation-of-1","title":"Scaling Laws for the Few-Shot Adaptation of Pre-trained Image Classifiers","date":"2021-10-13","arxiv_id":"2110.06990","n_code_links":0,"syntology":null},{"paper":null,"slug":"semantics-aware-attention-improves-neural","title":"Semantics-aware Attention Improves Neural Machine Translation","date":"2021-10-13","arxiv_id":"2110.06920","n_code_links":0,"syntology":null},{"paper":"/paper/study-of-positional-encoding-approaches-for","slug":"study-of-positional-encoding-approaches-for","title":"Study of positional encoding approaches for Audio Spectrogram Transformers","date":"2021-10-13","arxiv_id":"2110.06999","n_code_links":1,"syntology":null},{"paper":"/paper/the-dawn-of-quantum-natural-language","slug":"the-dawn-of-quantum-natural-language","title":"The Dawn of Quantum Natural Language Processing","date":"2021-10-13","arxiv_id":"2110.06510","n_code_links":2,"syntology":null},{"paper":"/paper/towards-efficient-nlp-a-standard-evaluation","slug":"towards-efficient-nlp-a-standard-evaluation","title":"Towards Efficient NLP: A Standard Evaluation and A Strong Baseline","date":"2021-10-13","arxiv_id":"2110.07038","n_code_links":1,"syntology":null},{"paper":null,"slug":"transform-and-bitstream-domain-image","title":"Transform and Bitstream Domain Image Classification","date":"2021-10-13","arxiv_id":"2110.06740","n_code_links":0,"syntology":null},{"paper":null,"slug":"understanding-of-emotion-perception-from-art","title":"Understanding of Emotion Perception from Art","date":"2021-10-13","arxiv_id":"2110.06486","n_code_links":0,"syntology":null},{"paper":"/paper/yformer-u-net-inspired-transformer-1","slug":"yformer-u-net-inspired-transformer-1","title":"Yformer: U-Net Inspired Transformer Architecture for Far Horizon Time Series Forecasting","date":"2021-10-13","arxiv_id":"2110.08255","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-large-scale-lexical-and-semantic-analysis","title":"Regionalized models for Spanish language variations based on Twitter","date":"2021-10-12","arxiv_id":"2110.06128","n_code_links":0,"syntology":null},{"paper":null,"slug":"all-dolphins-are-intelligent-and-some-are","title":"ALL Dolphins Are Intelligent and SOME Are Friendly: Probing BERT for Nouns' Semantic Properties and their Prototypicality","date":"2021-10-12","arxiv_id":"2110.06376","n_code_links":0,"syntology":null},{"paper":null,"slug":"attention-guided-generative-models-for","title":"Attention-guided Generative Models for Extractive Question Answering","date":"2021-10-12","arxiv_id":"2110.06393","n_code_links":0,"syntology":null},{"paper":"/paper/bertraffic-a-robust-bert-based-approach-for","slug":"bertraffic-a-robust-bert-based-approach-for","title":"BERTraffic: BERT-based Joint Speaker Role and Speaker Change Detection for Air Traffic Control Communications","date":"2021-10-12","arxiv_id":"2110.05781","n_code_links":2,"syntology":null},{"paper":"/paper/contrastive-learning-for-representation","slug":"contrastive-learning-for-representation","title":"Contrastive Learning for Representation Degeneration Problem in Sequential Recommendation","date":"2021-10-12","arxiv_id":"2110.05730","n_code_links":2,"syntology":{"ran":2,"of":4,"n_ran_checked":2,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["RuihongQiu/DuoRec"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/discodvt-generating-long-text-with-discourse","slug":"discodvt-generating-long-text-with-discourse","title":"DiscoDVT: Generating Long Text with Discourse-Aware Discrete Variational Transformer","date":"2021-10-12","arxiv_id":"2110.05999","n_code_links":1,"syntology":null},{"paper":null,"slug":"dynamic-inference-with-neural-interpreters","title":"Dynamic Inference with Neural Interpreters","date":"2021-10-12","arxiv_id":"2110.06399","n_code_links":0,"syntology":null},{"paper":null,"slug":"extracting-feelings-of-people-regarding-covid","title":"Extracting Feelings of People Regarding COVID-19 by Social Network Mining","date":"2021-10-12","arxiv_id":"2110.06151","n_code_links":0,"syntology":null},{"paper":"/paper/hetformer-heterogeneous-transformer-with","slug":"hetformer-heterogeneous-transformer-with","title":"HETFORMER: Heterogeneous Transformer with Sparse Attention for Long-Text Extractive Summarization","date":"2021-10-12","arxiv_id":"2110.06388","n_code_links":1,"syntology":null},{"paper":"/paper/lightseq-accelerated-training-for-transformer","slug":"lightseq-accelerated-training-for-transformer","title":"LightSeq2: Accelerated Training for Transformer-based Models on GPUs","date":"2021-10-12","arxiv_id":"2110.05722","n_code_links":1,"syntology":null},{"paper":"/paper/list-lite-self-training-makes-efficient-few-1","slug":"list-lite-self-training-makes-efficient-few-1","title":"LiST: Lite Prompted Self-training Makes Parameter-Efficient Few-shot Learners","date":"2021-10-12","arxiv_id":"2110.06274","n_code_links":1,"syntology":{"ran":9,"of":12,"n_ran_checked":7,"n_instrument":2,"unverified":3,"pointer_only":4,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 2 honoured, 0 violated, 5 with no contract checked; 2 where Syntology's instrument failed) · 3 unverified","official":{"repos":["microsoft/list"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":7,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":"/paper/mention-memory-incorporating-textual-1","slug":"mention-memory-incorporating-textual-1","title":"Mention Memory: incorporating textual knowledge into Transformers through entity mention attention","date":"2021-10-12","arxiv_id":"2110.06176","n_code_links":1,"syntology":null},{"paper":"/paper/relative-molecule-self-attention-transformer-1","slug":"relative-molecule-self-attention-transformer-1","title":"Relative Molecule Self-Attention Transformer","date":"2021-10-12","arxiv_id":"2110.05841","n_code_links":1,"syntology":{"ran":4,"of":4,"n_ran_checked":4,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":null}},{"paper":"/paper/rescoring-sequence-to-sequence-models-for","slug":"rescoring-sequence-to-sequence-models-for","title":"Rescoring Sequence-to-Sequence Models for Text Line Recognition with CTC-Prefixes","date":"2021-10-12","arxiv_id":"2110.05909","n_code_links":1,"syntology":null},{"paper":"/paper/satellite-image-semantic-segmentation","slug":"satellite-image-semantic-segmentation","title":"Satellite Image Semantic Segmentation","date":"2021-10-12","arxiv_id":"2110.05812","n_code_links":1,"syntology":null},{"paper":"/paper/starformer-transformer-with-state-action-1","slug":"starformer-transformer-with-state-action-1","title":"StARformer: Transformer with State-Action-Reward Representations for Visual Reinforcement Learning","date":"2021-10-12","arxiv_id":"2110.06206","n_code_links":1,"syntology":{"ran":3,"of":6,"n_ran_checked":1,"n_instrument":2,"unverified":3,"pointer_only":2,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 2 where Syntology's instrument failed) · 3 unverified","official":{"repos":["elicassion/StARformer"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":"/paper/swish-driven-googlenet-for-intelligent-analog","slug":"swish-driven-googlenet-for-intelligent-analog","title":"Swish-Driven GoogleNet for Intelligent Analog Beam Selection in Terahertz Beamspace MIMO","date":"2021-10-12","arxiv_id":"2110.05830","n_code_links":2,"syntology":null},{"paper":"/paper/adaptively-multi-view-and-temporal-fusing","slug":"adaptively-multi-view-and-temporal-fusing","title":"Adaptive Multi-view and Temporal Fusing Transformer for 3D Human Pose Estimation","date":"2021-10-11","arxiv_id":"2110.05092","n_code_links":0,"syntology":null},{"paper":"/paper/bridging-the-gap-between-label-and-reference-1","slug":"bridging-the-gap-between-label-and-reference-1","title":"Bridging the Gap between Label- and Reference-based Synthesis in Multi-attribute Image-to-Image Translation","date":"2021-10-11","arxiv_id":"2110.05055","n_code_links":1,"syntology":null},{"paper":"/paper/certified-patch-robustness-via-smoothed-1","slug":"certified-patch-robustness-via-smoothed-1","title":"Certified Patch Robustness via Smoothed Vision Transformers","date":"2021-10-11","arxiv_id":"2110.07719","n_code_links":1,"syntology":{"ran":6,"of":10,"n_ran_checked":4,"n_instrument":2,"unverified":4,"pointer_only":3,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 2 where Syntology's instrument failed) · 4 unverified","official":{"repos":["madrylab/smoothed-vit"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"improving-gender-fairness-of-pre-trained-1","title":"Improving Gender Fairness of Pre-Trained Language Models without Catastrophic Forgetting","date":"2021-10-11","arxiv_id":"2110.05367","n_code_links":0,"syntology":null},{"paper":"/paper/investigating-transfer-learning-capabilities","slug":"investigating-transfer-learning-capabilities","title":"Investigating Transfer Learning Capabilities of Vision Transformers and CNNs by Fine-Tuning a Single Trainable Block","date":"2021-10-11","arxiv_id":"2110.05270","n_code_links":1,"syntology":null},{"paper":null,"slug":"multi-task-learning-for-situated-multi-domain","title":"Multi-Task Learning for Situated Multi-Domain End-to-End Dialogue Systems","date":"2021-10-11","arxiv_id":"2110.05221","n_code_links":0,"syntology":null},{"paper":null,"slug":"multi-view-self-attention-based-transformer","title":"Multi-View Self-Attention Based Transformer for Speaker Recognition","date":"2021-10-11","arxiv_id":"2110.05036","n_code_links":0,"syntology":null},{"paper":null,"slug":"navigation-in-urban-environments-amongst","title":"Navigation In Urban Environments Amongst Pedestrians Using Multi-Objective Deep Reinforcement Learning","date":"2021-10-11","arxiv_id":"2110.05205","n_code_links":0,"syntology":null},{"paper":null,"slug":"offensive-language-detection-with-bert-based","title":"Offensive Language Detection with BERT-based models, By Customizing Attention Probabilities","date":"2021-10-11","arxiv_id":"2110.05133","n_code_links":0,"syntology":null},{"paper":null,"slug":"partial-variable-training-for-efficient-on","title":"Partial Variable Training for Efficient On-Device Federated Learning","date":"2021-10-11","arxiv_id":"2110.05607","n_code_links":0,"syntology":null},{"paper":null,"slug":"sru-pioneering-fast-recurrence-with-attention","title":"SRU++: Pioneering Fast Recurrence with Attention for Speech Recognition","date":"2021-10-11","arxiv_id":"2110.05571","n_code_links":0,"syntology":null},{"paper":"/paper/topic-modeling-clade-assisted-sentiment","slug":"topic-modeling-clade-assisted-sentiment","title":"Topic Modeling, Clade-assisted Sentiment Analysis, and Vaccine Brand Reputation Analysis of COVID-19 Vaccine-related Facebook Comments in the Philippines","date":"2021-10-11","arxiv_id":"2111.04416","n_code_links":1,"syntology":null},{"paper":"/paper/unsupervised-source-separation-via-bayesian","slug":"unsupervised-source-separation-via-bayesian","title":"Unsupervised Source Separation via Bayesian Inference in the Latent Domain","date":"2021-10-11","arxiv_id":"2110.05313","n_code_links":1,"syntology":null},{"paper":null,"slug":"dct-dynamic-compressive-transformer-for","title":"DCT: Dynamic Compressive Transformer for Modeling Unbounded Sequence","date":"2021-10-10","arxiv_id":"2110.04821","n_code_links":0,"syntology":null},{"paper":null,"slug":"identity-guided-face-generation-with-multi","title":"Identity-guided Face Generation with Multi-modal Contour Conditions","date":"2021-10-10","arxiv_id":"2110.04854","n_code_links":0,"syntology":null},{"paper":null,"slug":"multi-channel-end-to-end-neural-diarization","title":"Multi-Channel End-to-End Neural Diarization with Distributed Microphones","date":"2021-10-10","arxiv_id":"2110.04694","n_code_links":0,"syntology":null},{"paper":"/paper/nvit-vision-transformer-compression-and-1","slug":"nvit-vision-transformer-compression-and-1","title":"Global Vision Transformer Pruning with Hessian-Aware Saliency","date":"2021-10-10","arxiv_id":"2110.04869","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 1 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":null}},{"paper":null,"slug":"on-automatic-text-extractive-summarization","title":"Automatic Text Extractive Summarization Based on Graph and Pre-trained Language Model Attention","date":"2021-10-10","arxiv_id":"2110.04878","n_code_links":0,"syntology":null},{"paper":"/paper/paste-a-tagging-free-decoding-framework-using","slug":"paste-a-tagging-free-decoding-framework-using","title":"PASTE: A Tagging-Free Decoding Framework Using Pointer Networks for Aspect Sentiment Triplet Extraction","date":"2021-10-10","arxiv_id":"2110.04794","n_code_links":1,"syntology":null},{"paper":"/paper/reinforcement-learning-in-two-player-zero-sum","slug":"reinforcement-learning-in-two-player-zero-sum","title":"Reinforcement Learning In Two Player Zero Sum Simultaneous Action Games","date":"2021-10-10","arxiv_id":"2110.04835","n_code_links":1,"syntology":null},{"paper":"/paper/sp-gpt2-semantics-improvement-in-vietnamese","slug":"sp-gpt2-semantics-improvement-in-vietnamese","title":"SP-GPT2: Semantics Improvement in Vietnamese Poetry Generation","date":"2021-10-10","arxiv_id":"2110.15723","n_code_links":2,"syntology":null},{"paper":null,"slug":"supershaper-task-agnostic-super-pre-training","title":"SuperShaper: Task-Agnostic Super Pre-training of BERT Models with Variable Hidden Dimensions","date":"2021-10-10","arxiv_id":"2110.04711","n_code_links":0,"syntology":null},{"paper":"/paper/yuan-1-0-large-scale-pre-trained-language","slug":"yuan-1-0-large-scale-pre-trained-language","title":"Yuan 1.0: Large-Scale Pre-trained Language Model in Zero-Shot and Few-Shot Learning","date":"2021-10-10","arxiv_id":"2110.04725","n_code_links":1,"syntology":null},{"paper":"/paper/an-isotropy-analysis-in-the-multilingual-bert","slug":"an-isotropy-analysis-in-the-multilingual-bert","title":"An Isotropy Analysis in the Multilingual BERT Embedding Space","date":"2021-10-09","arxiv_id":"2110.04504","n_code_links":1,"syntology":null},{"paper":"/paper/end-to-end-keyword-spotting-using-xception-1d","slug":"end-to-end-keyword-spotting-using-xception-1d","title":"End-to-end Keyword Spotting using Xception-1d","date":"2021-10-09","arxiv_id":"2110.07498","n_code_links":1,"syntology":null},{"paper":"/paper/leveraging-recent-advances-in-pre-trained","slug":"leveraging-recent-advances-in-pre-trained","title":"Leveraging recent advances in Pre-Trained Language Models forEye-Tracking Prediction","date":"2021-10-09","arxiv_id":"2110.04475","n_code_links":1,"syntology":null},{"paper":"/paper/vector-quantized-image-modeling-with-improved-1","slug":"vector-quantized-image-modeling-with-improved-1","title":"Vector-quantized Image Modeling with Improved VQGAN","date":"2021-10-09","arxiv_id":"2110.04627","n_code_links":5,"syntology":{"ran":9,"of":9,"n_ran_checked":7,"n_instrument":2,"unverified":0,"pointer_only":2,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":null}},{"paper":"/paper/vision-transformer-based-covid-19-detection","slug":"vision-transformer-based-covid-19-detection","title":"Vision Transformer based COVID-19 Detection using Chest X-rays","date":"2021-10-09","arxiv_id":"2110.04458","n_code_links":0,"syntology":null},{"paper":null,"slug":"2110-06794","title":"The Layout Generation Algorithm of Graphic Design Based on Transformer-CVAE","date":"2021-10-08","arxiv_id":"2110.06794","n_code_links":0,"syntology":null},{"paper":null,"slug":"adversarial-token-attacks-on-vision","title":"Adversarial Token Attacks on Vision Transformers","date":"2021-10-08","arxiv_id":"2110.04337","n_code_links":0,"syntology":null},{"paper":null,"slug":"all-in-one-multi-task-learning-bert-models","title":"ALL-IN-ONE: Multi-Task Learning BERT models for Evaluating Peer Assessments","date":"2021-10-08","arxiv_id":"2110.03895","n_code_links":0,"syntology":null},{"paper":null,"slug":"context-lgm-leveraging-object-context","title":"Context-LGM: Leveraging Object-Context Relation for Context-Aware Object Recognition","date":"2021-10-08","arxiv_id":"2110.04042","n_code_links":0,"syntology":null},{"paper":null,"slug":"contrastive-learning-for-source-code-with","title":"Towards Learning (Dis)-Similarity of Source Code from Program Contrasts","date":"2021-10-08","arxiv_id":"2110.03868","n_code_links":0,"syntology":null},{"paper":null,"slug":"development-of-an-extractive-title-generation","title":"Development of an Extractive Title Generation System Using Titles of Papers of Top Conferences for Intermediate English Students","date":"2021-10-08","arxiv_id":"2110.04204","n_code_links":0,"syntology":null},{"paper":null,"slug":"evaluating-generative-networks-using-gaussian-1","title":"Evaluating generative networks using Gaussian mixtures of image features","date":"2021-10-08","arxiv_id":"2110.05240","n_code_links":0,"syntology":null},{"paper":"/paper/hydrasum-disentangling-stylistic-features-in-1","slug":"hydrasum-disentangling-stylistic-features-in-1","title":"HydraSum: Disentangling Stylistic Features in Text Summarization using Multi-Decoder Models","date":"2021-10-08","arxiv_id":"2110.04400","n_code_links":1,"syntology":{"ran":3,"of":9,"n_ran_checked":3,"n_instrument":0,"unverified":6,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 6 unverified","official":{"repos":["salesforce/hydra-sum"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":6,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"kg-fid-infusing-knowledge-graph-in-fusion-in-1","title":"KG-FiD: Infusing Knowledge Graph in Fusion-in-Decoder for Open-Domain Question Answering","date":"2021-10-08","arxiv_id":"2110.04330","n_code_links":0,"syntology":null},{"paper":"/paper/knowledge-enhanced-hierarchical-graph","slug":"knowledge-enhanced-hierarchical-graph","title":"Knowledge-Enhanced Hierarchical Graph Transformer Network for Multi-Behavior Recommendation","date":"2021-10-08","arxiv_id":"2110.04000","n_code_links":1,"syntology":null},{"paper":"/paper/learning-a-self-expressive-network-for-1","slug":"learning-a-self-expressive-network-for-1","title":"Learning a Self-Expressive Network for Subspace Clustering","date":"2021-10-08","arxiv_id":"2110.04318","n_code_links":1,"syntology":{"ran":4,"of":4,"n_ran_checked":4,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["zhangsz1998/self-expressive-network"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"learning-adaptive-control-flow-in","title":"Learning Adaptive Control Flow in Transformers for Improved Systematic Generalization","date":"2021-10-08","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/local-and-global-context-based-pairwise","slug":"local-and-global-context-based-pairwise","title":"Local and Global Context-Based Pairwise Models for Sentence Ordering","date":"2021-10-08","arxiv_id":"2110.04291","n_code_links":1,"syntology":null},{"paper":null,"slug":"m6-10t-a-sharing-delinking-paradigm-for","title":"M6-10T: A Sharing-Delinking Paradigm for Efficient Multi-Trillion Parameter Pretraining","date":"2021-10-08","arxiv_id":"2110.03888","n_code_links":0,"syntology":null},{"paper":"/paper/multiplex-behavioral-relation-learning-for","slug":"multiplex-behavioral-relation-learning-for","title":"Multiplex Behavioral Relation Learning for Recommendation via Memory Augmented Transformer Network","date":"2021-10-08","arxiv_id":"2110.04002","n_code_links":1,"syntology":null},{"paper":"/paper/rpt-toward-transferable-model-on","slug":"rpt-toward-transferable-model-on","title":"RPT: Toward Transferable Model on Heterogeneous Researcher Data via Pre-Training","date":"2021-10-08","arxiv_id":"2110.07336","n_code_links":1,"syntology":null},{"paper":null,"slug":"speeding-up-deep-model-training-by-sharing","title":"Speeding up Deep Model Training by Sharing Weights and Then Unsharing","date":"2021-10-08","arxiv_id":"2110.03848","n_code_links":0,"syntology":null},{"paper":"/paper/taming-sparsely-activated-transformer-with","slug":"taming-sparsely-activated-transformer-with","title":"Taming Sparsely Activated Transformer with Stochastic Experts","date":"2021-10-08","arxiv_id":"2110.04260","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":2,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["microsoft/stochastic-mixture-of-experts"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"text-analysis-and-deep-learning-a-network","title":"Text analysis and deep learning: A network approach","date":"2021-10-08","arxiv_id":"2110.04151","n_code_links":0,"syntology":null},{"paper":null,"slug":"token-pooling-in-visual-transformers","title":"Token Pooling in Vision Transformers","date":"2021-10-08","arxiv_id":"2110.03860","n_code_links":0,"syntology":null},{"paper":"/paper/uninet-unified-architecture-search-with","slug":"uninet-unified-architecture-search-with","title":"UniNet: Unified Architecture Search with Convolution, Transformer, and MLP","date":"2021-10-08","arxiv_id":"2110.04035","n_code_links":0,"syntology":null},{"paper":"/paper/vidt-an-efficient-and-effective-fully","slug":"vidt-an-efficient-and-effective-fully","title":"ViDT: An Efficient and Effective Fully Transformer-based Object Detector","date":"2021-10-08","arxiv_id":"2110.03921","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-comparative-study-of-transformer-based","title":"A Comparative Study of Transformer-Based Language Models on Extractive Question Answering","date":"2021-10-07","arxiv_id":"2110.03142","n_code_links":0,"syntology":null},{"paper":null,"slug":"attention-is-all-you-need-good-embeddings","title":"Attention is All You Need? Good Embeddings with Statistics are enough:Large Scale Audio Understanding without Transformers/ Convolutions/ BERTs/ Mixers/ Attention/ RNNs or ....","date":"2021-10-07","arxiv_id":"2110.03183","n_code_links":0,"syntology":null},{"paper":"/paper/cross-language-learning-for-entity-matching","slug":"cross-language-learning-for-entity-matching","title":"Cross-Language Learning for Entity Matching","date":"2021-10-07","arxiv_id":"2110.03338","n_code_links":1,"syntology":null},{"paper":"/paper/efficient-large-scale-image-retrieval-with","slug":"efficient-large-scale-image-retrieval-with","title":"Efficient large-scale image retrieval with deep feature orthogonality and Hybrid-Swin-Transformers","date":"2021-10-07","arxiv_id":"2110.03786","n_code_links":1,"syntology":null},{"paper":"/paper/end-to-end-supermask-pruning-learning-to","slug":"end-to-end-supermask-pruning-learning-to","title":"End-to-End Supermask Pruning: Learning to Prune Image Captioning Models","date":"2021-10-07","arxiv_id":"2110.03298","n_code_links":1,"syntology":{"ran":4,"of":6,"n_ran_checked":4,"n_instrument":0,"unverified":2,"pointer_only":1,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 1 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["jiahuei/sparse-image-captioning"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/flow-plugin-network-for-conditional","slug":"flow-plugin-network-for-conditional","title":"Flow Plugin Network for conditional generation","date":"2021-10-07","arxiv_id":"2110.04081","n_code_links":1,"syntology":null},{"paper":null,"slug":"generative-pre-trained-transformer-for","title":"Generative Pre-Trained Transformer for Cardiac Abnormality Detection","date":"2021-10-07","arxiv_id":"2110.04071","n_code_links":0,"syntology":null},{"paper":null,"slug":"investigating-growth-at-risk-using-a-multi","title":"Investigating Growth at Risk Using a Multi-country Non-parametric Quantile Factor Model","date":"2021-10-07","arxiv_id":"2110.03411","n_code_links":0,"syntology":null},{"paper":"/paper/layer-wise-pruning-of-transformer-attention","slug":"layer-wise-pruning-of-transformer-attention","title":"Layer-wise Pruning of Transformer Attention Heads for Efficient Language Modeling","date":"2021-10-07","arxiv_id":"2110.03252","n_code_links":1,"syntology":null},{"paper":null,"slug":"minimum-word-error-training-for-non","title":"Minimum word error training for non-autoregressive Transformer-based code-switching ASR","date":"2021-10-07","arxiv_id":"2110.03573","n_code_links":0,"syntology":null},{"paper":"/paper/mixer-tts-non-autoregressive-fast-and-compact","slug":"mixer-tts-non-autoregressive-fast-and-compact","title":"Mixer-TTS: non-autoregressive, fast and compact text-to-speech model conditioned on language model embeddings","date":"2021-10-07","arxiv_id":"2110.03584","n_code_links":1,"syntology":null},{"paper":null,"slug":"universality-of-deep-neural-network-lottery","title":"Universality of Winning Tickets: A Renormalization Group Perspective","date":"2021-10-07","arxiv_id":"2110.03210","n_code_links":0,"syntology":null}],"record_sha256":"bf328480263e546ec19833679579ad3126d351cc16d9a21de5096349c1dcdecb","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}