{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/transformer/papers/91","list_of":"/method/transformer","method":"Transformer","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":91,"pages_in_order":140,"rows_per_page":100,"rows":[9001,9100],"of":13999,"counts":{"archive_papers_tagged":13999,"with_a_code_link":6572,"where_syntology_ran_a_sample":2248,"not_listed_spam_title":0,"listed":13999,"listed_where_code_ran":2248,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":1919,"every_run_a_failure_of_syntologys_instrument":329,"listed_with_a_run_with_no_instrument_failure":1919,"listed_every_run_a_failure_of_syntologys_instrument":329,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/transformer","prev":"/method/transformer/papers/90","next":"/method/transformer/papers/92","papers":[{"paper":null,"slug":"learning-to-agree-on-vision-attention-for","title":"Learning to Agree on Vision Attention for Visual Commonsense Reasoning","date":"2023-02-04","arxiv_id":"2302.02117","n_code_links":0,"syntology":null},{"paper":"/paper/revisiting-image-deblurring-with-an-efficient","slug":"revisiting-image-deblurring-with-an-efficient","title":"Revisiting Image Deblurring with an Efficient ConvNet","date":"2023-02-04","arxiv_id":"2302.02234","n_code_links":1,"syntology":null},{"paper":null,"slug":"weight-is-attention-all-we-need-aeiuorder","title":"Greedy Ordering of Layer Weight Matrices in Transformers Improves Translation","date":"2023-02-04","arxiv_id":"2302.02123","n_code_links":0,"syntology":null},{"paper":null,"slug":"cfft-gan-cross-domain-feature-fusion","title":"CFFT-GAN: Cross-domain Feature Fusion Transformer for Exemplar-based Image Translation","date":"2023-02-03","arxiv_id":"2302.01608","n_code_links":0,"syntology":null},{"paper":null,"slug":"coinductive-guide-to-inductive-transformer","title":"Coinductive guide to inductive transformer heads","date":"2023-02-03","arxiv_id":"2302.01834","n_code_links":0,"syntology":null},{"paper":null,"slug":"device-depth-and-visual-concepts-aware","title":"DEVICE: DEpth and VIsual ConcEpts Aware Transformer for TextCaps","date":"2023-02-03","arxiv_id":"2302.01540","n_code_links":0,"syntology":null},{"paper":"/paper/dilateformer-multi-scale-dilated-transformer","slug":"dilateformer-multi-scale-dilated-transformer","title":"DilateFormer: Multi-Scale Dilated Transformer for Visual Recognition","date":"2023-02-03","arxiv_id":"2302.01791","n_code_links":1,"syntology":null},{"paper":"/paper/hdformer-high-order-directed-transformer-for","slug":"hdformer-high-order-directed-transformer-for","title":"HDFormer: High-order Directed Transformer for 3D Human Pose Estimation","date":"2023-02-03","arxiv_id":"2302.01825","n_code_links":1,"syntology":{"ran":6,"of":12,"n_ran_checked":4,"n_instrument":2,"unverified":6,"pointer_only":12,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 2 honoured, 1 violated, 1 with no contract checked; 2 where Syntology's instrument failed) · 6 unverified","official":{"repos":["hyer/hdformer"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":6,"ran_from_kinds":["official"]}}},{"paper":"/paper/revisiting-intermediate-layer-distillation","slug":"revisiting-intermediate-layer-distillation","title":"Revisiting Intermediate Layer Distillation for Compressing Language Models: An Overfitting Perspective","date":"2023-02-03","arxiv_id":"2302.01530","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-survey-on-efficient-training-of","title":"A Survey on Efficient Training of Transformers","date":"2023-02-02","arxiv_id":"2302.01107","n_code_links":0,"syntology":null},{"paper":null,"slug":"boosting-low-data-instance-segmentation-by","title":"Boosting Low-Data Instance Segmentation by Unsupervised Pre-training with Saliency Prompt","date":"2023-02-02","arxiv_id":"2302.01171","n_code_links":0,"syntology":null},{"paper":null,"slug":"curriculum-guided-abstractive-summarization-1","title":"Curriculum-Guided Abstractive Summarization","date":"2023-02-02","arxiv_id":"2302.01342","n_code_links":0,"syntology":null},{"paper":"/paper/dual-patchnorm","slug":"dual-patchnorm","title":"Dual PatchNorm","date":"2023-02-02","arxiv_id":"2302.01327","n_code_links":7,"syntology":{"ran":7,"of":9,"n_ran_checked":6,"n_instrument":1,"unverified":2,"pointer_only":2,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 1 honoured, 4 violated, 1 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","official":{"repos":["google-research/big_vision"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/fcb-swinv2-transformer-for-polyp-segmentation","slug":"fcb-swinv2-transformer-for-polyp-segmentation","title":"FCB-SwinV2 Transformer for Polyp Segmentation","date":"2023-02-02","arxiv_id":"2302.01027","n_code_links":0,"syntology":null},{"paper":null,"slug":"history-aware-hierarchical-transformer-for","title":"History-Aware Hierarchical Transformer for Multi-session Open-domain Dialogue System","date":"2023-02-02","arxiv_id":"2302.00907","n_code_links":0,"syntology":null},{"paper":null,"slug":"idt5-indonesian-version-of-multilingual-t5","title":"idT5: Indonesian Version of Multilingual T5 Transformer","date":"2023-02-02","arxiv_id":"2302.00856","n_code_links":0,"syntology":null},{"paper":null,"slug":"molecular-geometry-aware-transformer-for","title":"Molecular Geometry-aware Transformer for accurate 3D Atomic System modeling","date":"2023-02-02","arxiv_id":"2302.00855","n_code_links":0,"syntology":null},{"paper":null,"slug":"vision-transformer-based-feature-extraction","title":"Vision Transformer-based Feature Extraction for Generalized Zero-Shot Learning","date":"2023-02-02","arxiv_id":"2302.00875","n_code_links":0,"syntology":null},{"paper":null,"slug":"clinical-decision-transformer-intended","title":"Clinical Decision Transformer: Intended Treatment Recommendation through Goal Prompting","date":"2023-02-01","arxiv_id":"2302.00612","n_code_links":0,"syntology":null},{"paper":"/paper/feed-forward-blocks-control-contextualization","slug":"feed-forward-blocks-control-contextualization","title":"Analyzing Feed-Forward Blocks in Transformers through the Lens of Attention Maps","date":"2023-02-01","arxiv_id":"2302.00456","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":1,"n_instrument":1,"unverified":0,"pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["gorokoba560/norm-analysis-of-transformer"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/multispectral-pedestrian-detection-via-1","slug":"multispectral-pedestrian-detection-via-1","title":"MS-DETR: Multispectral Pedestrian Detection Transformer with Loosely Coupled Fusion and Modality-Balanced Optimization","date":"2023-02-01","arxiv_id":"2302.00290","n_code_links":1,"syntology":null},{"paper":"/paper/continuous-spatiotemporal-transformers","slug":"continuous-spatiotemporal-transformers","title":"Continuous Spatiotemporal Transformers","date":"2023-01-31","arxiv_id":"2301.13338","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":3,"n_instrument":0,"unverified":0,"pointer_only":2,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 1 honoured, 2 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["vandijklab/cst"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official","unlocated"]}}},{"paper":"/paper/fairness-aware-vision-transformer-via","slug":"fairness-aware-vision-transformer-via","title":"Fairness-aware Vision Transformer via Debiased Self-Attention","date":"2023-01-31","arxiv_id":"2301.13803","n_code_links":1,"syntology":null},{"paper":null,"slug":"skill-decision-transformer","title":"Skill Decision Transformer","date":"2023-01-31","arxiv_id":"2301.13573","n_code_links":0,"syntology":null},{"paper":"/paper/upop-unified-and-progressive-pruning-for","slug":"upop-unified-and-progressive-pruning-for","title":"UPop: Unified and Progressive Pruning for Compressing Vision-Language Transformers","date":"2023-01-31","arxiv_id":"2301.13741","n_code_links":2,"syntology":{"ran":0,"of":2,"n_ran_checked":0,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"0 ran · 2 unverified","official":{"repos":["sdc17/upop"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":[]}}},{"paper":"/paper/blip-2-bootstrapping-language-image-pre","slug":"blip-2-bootstrapping-language-image-pre","title":"BLIP-2: Bootstrapping Language-Image Pre-training with Frozen Image Encoders and Large Language Models","date":"2023-01-30","arxiv_id":"2301.12597","n_code_links":17,"syntology":{"ran":4,"of":8,"n_ran_checked":4,"n_instrument":0,"unverified":4,"pointer_only":1,"phrase":"4 ran (of which 4 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified; every one of the 4 samples that ran constructed an object rather than computing a result","official":{"repos":["salesforce/lavis"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/guiding-online-reinforcement-learning-with","slug":"guiding-online-reinforcement-learning-with","title":"Guiding Online Reinforcement Learning with Action-Free Offline Pretraining","date":"2023-01-30","arxiv_id":"2301.12876","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["vision-cair/af-guide"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"pseudo-3d-perception-transformer-with-multi","title":"Multi-modal Large Language Model Enhanced Pseudo 3D Perception Framework for Visual Commonsense Reasoning","date":"2023-01-30","arxiv_id":"2301.13335","n_code_links":0,"syntology":null},{"paper":null,"slug":"exploring-attention-map-reuse-for-efficient","title":"Exploring Attention Map Reuse for Efficient Transformer Neural Networks","date":"2023-01-29","arxiv_id":"2301.12444","n_code_links":0,"syntology":null},{"paper":"/paper/graph-mixer-networks","slug":"graph-mixer-networks","title":"Graph Mixer Networks","date":"2023-01-29","arxiv_id":"2301.12493","n_code_links":1,"syntology":null},{"paper":"/paper/phavip-phage-virion-protein-classification","slug":"phavip-phage-virion-protein-classification","title":"PhaVIP: Phage VIrion Protein classification based on chaos game representation and Vision Transformer","date":"2023-01-29","arxiv_id":"2301.12422","n_code_links":1,"syntology":null},{"paper":null,"slug":"semantics-enhanced-temporal-graph-networks","title":"Semantics-enhanced Temporal Graph Networks for Content Popularity Prediction","date":"2023-01-29","arxiv_id":"2301.12355","n_code_links":0,"syntology":null},{"paper":null,"slug":"towards-vision-transformer-unrolling-fixed","title":"Towards Vision Transformer Unrolling Fixed-Point Algorithm: a Case Study on Image Restoration","date":"2023-01-29","arxiv_id":"2301.12332","n_code_links":0,"syntology":null},{"paper":"/paper/aerial-image-object-detection-with-vision","slug":"aerial-image-object-detection-with-vision","title":"Aerial Image Object Detection With Vision Transformer Detector (ViTDet)","date":"2023-01-28","arxiv_id":"2301.12058","n_code_links":1,"syntology":null},{"paper":"/paper/multilingual-sentence-transformer-as-a","slug":"multilingual-sentence-transformer-as-a","title":"Multilingual Sentence Transformer as A Multilingual Word Aligner","date":"2023-01-28","arxiv_id":"2301.12140","n_code_links":1,"syntology":null},{"paper":"/paper/predicting-visit-cost-of-obstructive-sleep","slug":"predicting-visit-cost-of-obstructive-sleep","title":"Predicting Visit Cost of Obstructive Sleep Apnea using Electronic Healthcare Records with Transformer","date":"2023-01-28","arxiv_id":"2301.12289","n_code_links":1,"syntology":null},{"paper":null,"slug":"towards-a-single-unified-model-for-effective","title":"CancerUniT: Towards a Single Unified Model for Effective Detection, Segmentation, and Diagnosis of Eight Major Cancers Using a Large Collection of CT Scans","date":"2023-01-28","arxiv_id":"2301.12291","n_code_links":0,"syntology":null},{"paper":null,"slug":"can-we-use-probing-to-better-understand-fine","title":"Can We Use Probing to Better Understand Fine-tuning and Knowledge Distillation of the BERT NLU?","date":"2023-01-27","arxiv_id":"2301.11688","n_code_links":0,"syntology":null},{"paper":"/paper/cross-architectural-positive-pairs-improve","slug":"cross-architectural-positive-pairs-improve","title":"Cross-Architectural Positive Pairs improve the effectiveness of Self-Supervised Learning","date":"2023-01-27","arxiv_id":"2301.12025","n_code_links":1,"syntology":null},{"paper":null,"slug":"large-scale-traffic-data-imputation-with","title":"Large-Scale Traffic Data Imputation with Spatiotemporal Semantic Understanding","date":"2023-01-27","arxiv_id":"2301.11691","n_code_links":0,"syntology":null},{"paper":"/paper/on-the-connection-between-mpnn-and-graph","slug":"on-the-connection-between-mpnn-and-graph","title":"On the Connection Between MPNN and Graph Transformer","date":"2023-01-27","arxiv_id":"2301.11956","n_code_links":1,"syntology":null},{"paper":"/paper/pre-training-for-speech-translation-ctc-meets","slug":"pre-training-for-speech-translation-ctc-meets","title":"Pre-training for Speech Translation: CTC Meets Optimal Transport","date":"2023-01-27","arxiv_id":"2301.11716","n_code_links":1,"syntology":null},{"paper":null,"slug":"skeleton-based-action-recognition-through","title":"Skeleton-based Action Recognition through Contrasting Two-Stream Spatial-Temporal Networks","date":"2023-01-27","arxiv_id":"2301.11495","n_code_links":0,"syntology":null},{"paper":"/paper/swarm-parallelism-training-large-models-can-1","slug":"swarm-parallelism-training-large-models-can-1","title":"SWARM Parallelism: Training Large Models Can Be Surprisingly Communication-Efficient","date":"2023-01-27","arxiv_id":"2301.11913","n_code_links":2,"syntology":null},{"paper":null,"slug":"semantic-segmentation-enhanced-transformer","title":"Semantic Segmentation Enhanced Transformer Model for Human Attention Prediction","date":"2023-01-26","arxiv_id":"2301.11022","n_code_links":0,"syntology":null},{"paper":"/paper/enhancing-medical-image-segmentation-with","slug":"enhancing-medical-image-segmentation-with","title":"Enhancing Medical Image Segmentation with TransCeption: A Multi-Scale Feature Fusion Approach","date":"2023-01-25","arxiv_id":"2301.10847","n_code_links":1,"syntology":null},{"paper":null,"slug":"qualitative-analysis-of-a-graph-transformer","title":"Qualitative Analysis of a Graph Transformer Approach to Addressing Hate Speech: Adapting to Dynamically Changing Content","date":"2023-01-25","arxiv_id":"2301.10871","n_code_links":0,"syntology":null},{"paper":"/paper/transfer-learning-in-deep-learning-models-for","slug":"transfer-learning-in-deep-learning-models-for","title":"Transfer Learning in Deep Learning Models for Building Load Forecasting: Case of Limited Data","date":"2023-01-25","arxiv_id":"2301.10663","n_code_links":2,"syntology":null},{"paper":"/paper/videberta-a-powerful-pre-trained-language","slug":"videberta-a-powerful-pre-trained-language","title":"ViDeBERTa: A powerful pre-trained language model for Vietnamese","date":"2023-01-25","arxiv_id":"2301.10439","n_code_links":1,"syntology":{"ran":2,"of":4,"n_ran_checked":2,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["hysonlab/videberta"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/a-watermark-for-large-language-models","slug":"a-watermark-for-large-language-models","title":"A Watermark for Large Language Models","date":"2023-01-24","arxiv_id":"2301.10226","n_code_links":8,"syntology":{"ran":11,"of":16,"n_ran_checked":10,"n_instrument":1,"unverified":5,"pointer_only":5,"phrase":"11 ran (of which 7 constructed an object rather than computing a result; 10 with no instrument failure: 1 honoured, 1 violated, 8 with no contract checked; 1 where Syntology's instrument failed) · 5 unverified","official":{"repos":["jwkirchenbauer/lm-watermarking"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["listed","official"]}}},{"paper":"/paper/climax-a-foundation-model-for-weather-and","slug":"climax-a-foundation-model-for-weather-and","title":"ClimaX: A foundation model for weather and climate","date":"2023-01-24","arxiv_id":"2301.10343","n_code_links":1,"syntology":null},{"paper":null,"slug":"smart-self-supervised-multi-task-pretraining","title":"SMART: Self-supervised Multi-task pretrAining with contRol Transformers","date":"2023-01-24","arxiv_id":"2301.09816","n_code_links":0,"syntology":null},{"paper":null,"slug":"deep-learning-mental-health-dialogue-system","title":"Deep Learning Mental Health Dialogue System","date":"2023-01-23","arxiv_id":"2301.09412","n_code_links":0,"syntology":null},{"paper":"/paper/efficient-language-model-training-through","slug":"efficient-language-model-training-through","title":"Efficient Language Model Training through Cross-Lingual and Progressive Transfer Learning","date":"2023-01-23","arxiv_id":"2301.09626","n_code_links":1,"syntology":null},{"paper":"/paper/fully-transformer-based-biomarker-prediction","slug":"fully-transformer-based-biomarker-prediction","title":"Fully transformer-based biomarker prediction from colorectal cancer histology: a large-scale multicentric study","date":"2023-01-23","arxiv_id":"2301.09617","n_code_links":2,"syntology":null},{"paper":"/paper/istvt-interpretable-spatial-temporal-video","slug":"istvt-interpretable-spatial-temporal-video","title":"ISTVT: Interpretable Spatial-Temporal Video Transformer for Deepfake Detection","date":"2023-01-23","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"learning-to-view-decision-transformers-for","title":"Learning to View: Decision Transformers for Active Object Detection","date":"2023-01-23","arxiv_id":"2301.09544","n_code_links":0,"syntology":null},{"paper":"/paper/local-window-attention-transformer-for","slug":"local-window-attention-transformer-for","title":"Local Window Attention Transformer for Polarimetric SAR Image Classification","date":"2023-01-23","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"apples-and-oranges-assessing-image-quality","title":"Apples and Oranges? Assessing Image Quality over Content Recognition","date":"2023-01-22","arxiv_id":"2301.09190","n_code_links":0,"syntology":null},{"paper":"/paper/debiasing-the-cloze-task-in-sequential","slug":"debiasing-the-cloze-task-in-sequential","title":"Debiasing the Cloze Task in Sequential Recommendation with Bidirectional Transformers","date":"2023-01-22","arxiv_id":"2301.09210","n_code_links":1,"syntology":null},{"paper":null,"slug":"slice-transformer-and-self-supervised","title":"Slice Transformer and Self-supervised Learning for 6DoF Localization in 3D Point Cloud Maps","date":"2023-01-21","arxiv_id":"2301.08957","n_code_links":0,"syntology":null},{"paper":"/paper/time-conditioned-generative-modeling-of","slug":"time-conditioned-generative-modeling-of","title":"Time-Conditioned Generative Modeling of Object-Centric Representations for Video Decomposition and Prediction","date":"2023-01-21","arxiv_id":"2301.08951","n_code_links":1,"syntology":null},{"paper":null,"slug":"accelerating-multi-agent-planning-using-graph","title":"Accelerating Multi-Agent Planning Using Graph Transformers with Bounded Suboptimality","date":"2023-01-20","arxiv_id":"2301.08451","n_code_links":0,"syntology":null},{"paper":null,"slug":"ontology-pre-training-for-poison-prediction","title":"Ontology Pre-training for Poison Prediction","date":"2023-01-20","arxiv_id":"2301.08577","n_code_links":0,"syntology":null},{"paper":"/paper/diagnose-like-a-pathologist-transformer","slug":"diagnose-like-a-pathologist-transformer","title":"Diagnose Like a Pathologist: Transformer-Enabled Hierarchical Attention-Guided Multiple Instance Learning for Whole Slide Image Classification","date":"2023-01-19","arxiv_id":"2301.08125","n_code_links":1,"syntology":null},{"paper":null,"slug":"fe-tcm-filter-enhanced-transformer-click","title":"FE-TCM: Filter-Enhanced Transformer Click Model for Web Search","date":"2023-01-19","arxiv_id":"2301.07854","n_code_links":0,"syntology":null},{"paper":"/paper/medsegdiff-v2-diffusion-based-medical-image","slug":"medsegdiff-v2-diffusion-based-medical-image","title":"MedSegDiff-V2: Diffusion based Medical Image Segmentation with Transformer","date":"2023-01-19","arxiv_id":"2301.11798","n_code_links":2,"syntology":{"ran":17,"of":21,"n_ran_checked":12,"n_instrument":5,"unverified":4,"pointer_only":10,"phrase":"17 ran (of which 0 constructed an object rather than computing a result; 12 with no instrument failure: 2 honoured, 0 violated, 10 with no contract checked; 5 where Syntology's instrument failed) · 4 unverified","official":{"repos":["kidswithtokens/medsegdiff","wujunde/medsegdiff"],"state":"official (archive's flag): 17 ran","n_ran":17,"n_constructed":0,"n_ran_no_instrument_failure":12,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"cooperation-learning-enhanced-colonic-polyp","title":"Cooperation Learning Enhanced Colonic Polyp Segmentation Based on Transformer-CNN Fusion","date":"2023-01-17","arxiv_id":"2301.06892","n_code_links":0,"syntology":null},{"paper":"/paper/sat-size-aware-transformer-for-3d-point-cloud","slug":"sat-size-aware-transformer-for-3d-point-cloud","title":"SAT: Size-Aware Transformer for 3D Point Cloud Semantic Segmentation","date":"2023-01-17","arxiv_id":"2301.06869","n_code_links":0,"syntology":null},{"paper":"/paper/swindepth-unsupervised-depth-estimation-using","slug":"swindepth-unsupervised-depth-estimation-using","title":"SwinDepth: Unsupervised Depth Estimation using Monocular Sequences via Swin Transformer and Densely Cascaded Network","date":"2023-01-17","arxiv_id":"2301.06715","n_code_links":1,"syntology":null},{"paper":"/paper/tracing-and-manipulating-intermediate-values","slug":"tracing-and-manipulating-intermediate-values","title":"Tracing and Manipulating Intermediate Values in Neural Math Problem Solvers","date":"2023-01-17","arxiv_id":"2301.06758","n_code_links":1,"syntology":null},{"paper":null,"slug":"transformer-based-implementation-for","title":"Transformer Based Implementation for Automatic Book Summarization","date":"2023-01-17","arxiv_id":"2301.07057","n_code_links":0,"syntology":null},{"paper":null,"slug":"bayesspeech-a-bayesian-transformer-network","title":"BayesSpeech: A Bayesian Transformer Network for Automatic Speech Recognition","date":"2023-01-16","arxiv_id":"2301.11276","n_code_links":0,"syntology":null},{"paper":"/paper/tdstf-transformer-based-diffusion","slug":"tdstf-transformer-based-diffusion","title":"A Transformer-based Diffusion Probabilistic Model for Heart Rate and Blood Pressure Forecasting in Intensive Care Unit","date":"2023-01-16","arxiv_id":"2301.06625","n_code_links":1,"syntology":null},{"paper":"/paper/dsvt-dynamic-sparse-voxel-transformer-with","slug":"dsvt-dynamic-sparse-voxel-transformer-with","title":"DSVT: Dynamic Sparse Voxel Transformer with Rotated Sets","date":"2023-01-15","arxiv_id":"2301.06051","n_code_links":4,"syntology":null},{"paper":"/paper/t2m-gpt-generating-human-motion-from-textual","slug":"t2m-gpt-generating-human-motion-from-textual","title":"T2M-GPT: Generating Human Motion from Textual Descriptions with Discrete Representations","date":"2023-01-15","arxiv_id":"2301.06052","n_code_links":1,"syntology":{"ran":5,"of":7,"n_ran_checked":5,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"5 ran (of which 4 constructed an object rather than computing a result; 5 with no instrument failure: 1 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["Mael-zys/T2M-GPT"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":4,"n_ran_no_instrument_failure":5,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"salient-sign-detection-in-safe-autonomous","title":"Salient Sign Detection In Safe Autonomous Driving: AI Which Reasons Over Full Visual Context","date":"2023-01-14","arxiv_id":"2301.05804","n_code_links":0,"syntology":null},{"paper":null,"slug":"automated-speech-and-text-based","title":"Automated speech- and text-based classification of neuropsychiatric conditions in a multidiagnostic setting","date":"2023-01-13","arxiv_id":"2301.06916","n_code_links":0,"syntology":null},{"paper":null,"slug":"text-to-point-cloud-localization-with","title":"Text to Point Cloud Localization with Relation-Enhanced Transformer","date":"2023-01-13","arxiv_id":"2301.05372","n_code_links":0,"syntology":null},{"paper":"/paper/adversarial-adaptation-for-french-named","slug":"adversarial-adaptation-for-french-named","title":"Adversarial Adaptation for French Named Entity Recognition","date":"2023-01-12","arxiv_id":"2301.05220","n_code_links":1,"syntology":null},{"paper":"/paper/vits-for-sits-vision-transformers-for","slug":"vits-for-sits-vision-transformers-for","title":"ViTs for SITS: Vision Transformers for Satellite Image Time Series","date":"2023-01-12","arxiv_id":"2301.04944","n_code_links":3,"syntology":{"ran":11,"of":11,"n_ran_checked":11,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"11 ran (of which 3 constructed an object rather than computing a result; 11 with no instrument failure: 0 honoured, 0 violated, 11 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["michaeltrs/deepsatmodels"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":3,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["listed","official"]}}},{"paper":"/paper/adapointr-diverse-point-cloud-completion-with","slug":"adapointr-diverse-point-cloud-completion-with","title":"AdaPoinTr: Diverse Point Cloud Completion with Adaptive Geometry-Aware Transformers","date":"2023-01-11","arxiv_id":"2301.04545","n_code_links":1,"syntology":{"ran":4,"of":6,"n_ran_checked":2,"n_instrument":2,"unverified":2,"pointer_only":3,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 2 where Syntology's instrument failed) · 2 unverified","official":{"repos":["yuxumin/PoinTr"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"anomalies-representations-and-self","title":"Anomalies, Representations, and Self-Supervision","date":"2023-01-11","arxiv_id":"2301.04660","n_code_links":0,"syntology":null},{"paper":null,"slug":"dynamic-background-reconstruction-via","title":"Dynamic Background Reconstruction via MAE for Infrared Small Target Detection","date":"2023-01-11","arxiv_id":"2301.04497","n_code_links":0,"syntology":null},{"paper":"/paper/head-free-lightweight-semantic-segmentation","slug":"head-free-lightweight-semantic-segmentation","title":"Head-Free Lightweight Semantic Segmentation with Linear Transformer","date":"2023-01-11","arxiv_id":"2301.04648","n_code_links":1,"syntology":null},{"paper":"/paper/predicting-hateful-discussions-on-reddit","slug":"predicting-hateful-discussions-on-reddit","title":"Predicting Hateful Discussions on Reddit using Graph Transformer Networks and Communal Context","date":"2023-01-10","arxiv_id":"2301.04248","n_code_links":1,"syntology":null},{"paper":null,"slug":"streaming-punctuation-a-novel-punctuation","title":"Streaming Punctuation: A Novel Punctuation Technique Leveraging Bidirectional Context for Continuous Speech Recognition","date":"2023-01-10","arxiv_id":"2301.03819","n_code_links":0,"syntology":null},{"paper":"/paper/unsupervised-mandarin-cantonese-machine","slug":"unsupervised-mandarin-cantonese-machine","title":"Unsupervised Mandarin-Cantonese Machine Translation","date":"2023-01-10","arxiv_id":"2301.03971","n_code_links":1,"syntology":null},{"paper":"/paper/a-study-on-the-generality-of-neural-network","slug":"a-study-on-the-generality-of-neural-network","title":"A Study on the Generality of Neural Network Structures for Monocular Depth Estimation","date":"2023-01-09","arxiv_id":"2301.03169","n_code_links":1,"syntology":null},{"paper":"/paper/advances-in-medical-image-analysis-with","slug":"advances-in-medical-image-analysis-with","title":"Advances in Medical Image Analysis with Vision Transformers: A Comprehensive Review","date":"2023-01-09","arxiv_id":"2301.03505","n_code_links":1,"syntology":null},{"paper":"/paper/an-impartial-transformer-for-story","slug":"an-impartial-transformer-for-story","title":"An Impartial Transformer for Story Visualization","date":"2023-01-09","arxiv_id":"2301.03563","n_code_links":0,"syntology":null},{"paper":"/paper/demt-deformable-mixer-transformer-for-multi","slug":"demt-deformable-mixer-transformer-for-multi","title":"DeMT: Deformable Mixer Transformer for Multi-Task Learning of Dense Prediction","date":"2023-01-09","arxiv_id":"2301.03461","n_code_links":1,"syntology":{"ran":3,"of":6,"n_ran_checked":3,"n_instrument":0,"unverified":3,"pointer_only":6,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":{"repos":["yangyangxu0/demt"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"logically-at-factify-2023-a-multi-modal-fact","title":"Logically at Factify 2: A Multi-Modal Fact Checking System Based on Evidence Retrieval techniques and Transformer Encoder Architecture","date":"2023-01-09","arxiv_id":"2301.03127","n_code_links":0,"syntology":null},{"paper":null,"slug":"universal-multimodal-representation-for","title":"Universal Multimodal Representation for Language Understanding","date":"2023-01-09","arxiv_id":"2301.03344","n_code_links":0,"syntology":null},{"paper":"/paper/deepmatcher-a-deep-transformer-based-network","slug":"deepmatcher-a-deep-transformer-based-network","title":"DeepMatcher: A Deep Transformer-based Network for Robust and Accurate Local Feature Matching","date":"2023-01-08","arxiv_id":"2301.02993","n_code_links":1,"syntology":null},{"paper":"/paper/hrtransnet-hrformer-driven-two-modality","slug":"hrtransnet-hrformer-driven-two-modality","title":"HRTransNet: HRFormer-Driven Two-Modality Salient Object Detection","date":"2023-01-08","arxiv_id":"2301.03036","n_code_links":1,"syntology":null},{"paper":"/paper/codetalker-speech-driven-3d-facial-animation","slug":"codetalker-speech-driven-3d-facial-animation","title":"CodeTalker: Speech-Driven 3D Facial Animation with Discrete Motion Prior","date":"2023-01-06","arxiv_id":"2301.02379","n_code_links":1,"syntology":{"ran":8,"of":13,"n_ran_checked":8,"n_instrument":0,"unverified":5,"pointer_only":0,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 1 violated, 7 with no contract checked; 0 where Syntology's instrument failed) · 5 unverified","official":{"repos":["Doubiiu/CodeTalker"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":5,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"does-compressing-activations-help-model","title":"Does compressing activations help model parallel training?","date":"2023-01-06","arxiv_id":"2301.02654","n_code_links":0,"syntology":null},{"paper":"/paper/exploring-efficient-few-shot-adaptation-for","slug":"exploring-efficient-few-shot-adaptation-for","title":"Exploring Efficient Few-shot Adaptation for Vision Transformers","date":"2023-01-06","arxiv_id":"2301.02419","n_code_links":1,"syntology":null},{"paper":null,"slug":"multi-genre-music-transformer-composing-full","title":"Multi-Genre Music Transformer -- Composing Full Length Musical Piece","date":"2023-01-06","arxiv_id":"2301.02385","n_code_links":0,"syntology":null}],"record_sha256":"838b774cbd8304569df7eef61f36b91c3b87c7974c9c543e2058be3872d9f8f7","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}