{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/attention/papers/214","list_of":"/method/attention","method":"Attention","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":214,"pages_in_order":316,"rows_per_page":100,"rows":[21301,21400],"of":31583,"counts":{"archive_papers_tagged":31583,"with_a_code_link":13473,"where_syntology_ran_a_sample":3998,"not_listed_spam_title":0,"listed":31583,"listed_where_code_ran":3998,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":3366,"every_run_a_failure_of_syntologys_instrument":632,"listed_with_a_run_with_no_instrument_failure":3366,"listed_every_run_a_failure_of_syntologys_instrument":632,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/attention","prev":"/method/attention/papers/213","next":"/method/attention/papers/215","papers":[{"paper":null,"slug":"academic-writing-with-gpt-3-5-reflections-on","title":"Academic Writing with GPT-3.5: Reflections on Practices, Efficacy and Transparency","date":"2023-02-12","arxiv_id":"2304.11079","n_code_links":0,"syntology":null},{"paper":"/paper/denoising-and-prompt-tuning-for-multi","slug":"denoising-and-prompt-tuning-for-multi","title":"Denoising and Prompt-Tuning for Multi-Behavior Recommendation","date":"2023-02-12","arxiv_id":"2302.05862","n_code_links":1,"syntology":null},{"paper":"/paper/generalized-few-shot-continual-learning-with","slug":"generalized-few-shot-continual-learning-with","title":"Generalized Few-Shot Continual Learning with Contrastive Mixture of Adapters","date":"2023-02-12","arxiv_id":"2302.05936","n_code_links":1,"syntology":null},{"paper":"/paper/self-supervised-pseudo-colorizing-of-masked","slug":"self-supervised-pseudo-colorizing-of-masked","title":"Self-supervised pseudo-colorizing of masked cells","date":"2023-02-12","arxiv_id":"2302.05968","n_code_links":2,"syntology":null},{"paper":null,"slug":"semantic-communications-with-ordered","title":"Semantic Importance-Aware Communications Using Pre-trained Language Models","date":"2023-02-12","arxiv_id":"2302.07142","n_code_links":0,"syntology":null},{"paper":null,"slug":"transformer-models-an-introduction-and","title":"Transformer models: an introduction and catalog","date":"2023-02-12","arxiv_id":"2302.07730","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-brief-report-on-lawgpt-1-0-a-virtual-legal","title":"A Brief Report on LawGPT 1.0: A Virtual Legal Assistant Based on GPT-3","date":"2023-02-11","arxiv_id":"2302.05729","n_code_links":0,"syntology":null},{"paper":"/paper/differentiable-outlier-detection-enable","slug":"differentiable-outlier-detection-enable","title":"Differentiable Outlier Detection Enable Robust Deep Multimodal Analysis","date":"2023-02-11","arxiv_id":"2302.05608","n_code_links":1,"syntology":null},{"paper":"/paper/docile-benchmark-for-document-information","slug":"docile-benchmark-for-document-information","title":"DocILE Benchmark for Document Information Localization and Extraction","date":"2023-02-11","arxiv_id":"2302.05658","n_code_links":1,"syntology":null},{"paper":null,"slug":"informing-clinical-assessment-by","title":"Informing clinical assessment by contextualizing post-hoc explanations of risk prediction models in type-2 diabetes","date":"2023-02-11","arxiv_id":"2302.05752","n_code_links":0,"syntology":null},{"paper":null,"slug":"rethinking-vision-transformer-and-masked","title":"Rethinking Vision Transformer and Masked Autoencoder in Multimodal Face Anti-Spoofing","date":"2023-02-11","arxiv_id":"2302.05744","n_code_links":0,"syntology":null},{"paper":"/paper/alloprof-a-new-french-question-answer","slug":"alloprof-a-new-french-question-answer","title":"Alloprof: a new French question-answer education dataset and its use in an information retrieval case study","date":"2023-02-10","arxiv_id":"2302.07738","n_code_links":1,"syntology":null},{"paper":null,"slug":"best-bert-pre-training-for-sign-language","title":"BEST: BERT Pre-Training for Sign Language Recognition with Coupling Tokenization","date":"2023-02-10","arxiv_id":"2302.05075","n_code_links":0,"syntology":null},{"paper":"/paper/combat-ai-with-ai-counteract-machine","slug":"combat-ai-with-ai-counteract-machine","title":"Combat AI With AI: Counteract Machine-Generated Fake Restaurant Reviews on Social Media","date":"2023-02-10","arxiv_id":"2302.07731","n_code_links":1,"syntology":null},{"paper":"/paper/dual-memory-units-with-uncertainty-regulation","slug":"dual-memory-units-with-uncertainty-regulation","title":"Dual Memory Units with Uncertainty Regulation for Weakly Supervised Video Anomaly Detection","date":"2023-02-10","arxiv_id":"2302.05160","n_code_links":1,"syntology":{"ran":7,"of":7,"n_ran_checked":7,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"7 ran (of which 6 constructed an object rather than computing a result; 7 with no instrument failure: 1 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["henrryzh1/UR-DMU"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":6,"n_ran_no_instrument_failure":7,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/effective-document-image-enhancement-using","slug":"effective-document-image-enhancement-using","title":"Effective Document Image Enhancement Using tokens-to-token Transformer Network","date":"2023-02-10","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/fairpy-a-toolkit-for-evaluation-of-social","slug":"fairpy-a-toolkit-for-evaluation-of-social","title":"FairPy: A Toolkit for Evaluation of Prediction Biases and their Mitigation in Large Language Models","date":"2023-02-10","arxiv_id":"2302.05508","n_code_links":1,"syntology":null},{"paper":null,"slug":"gtr-ctrl-instrument-and-genre-conditioning","title":"GTR-CTRL: Instrument and Genre Conditioning for Guitar-Focused Music Generation with Transformers","date":"2023-02-10","arxiv_id":"2302.05393","n_code_links":0,"syntology":null},{"paper":null,"slug":"patcorrect-non-autoregressive-phoneme","title":"PATCorrect: Non-autoregressive Phoneme-augmented Transformer for ASR Error Correction","date":"2023-02-10","arxiv_id":"2302.05040","n_code_links":0,"syntology":null},{"paper":"/paper/the-wisdom-of-hindsight-makes-language-models","slug":"the-wisdom-of-hindsight-makes-language-models","title":"The Wisdom of Hindsight Makes Language Models Better Instruction Followers","date":"2023-02-10","arxiv_id":"2302.05206","n_code_links":1,"syntology":null},{"paper":"/paper/translating-natural-language-to-planning","slug":"translating-natural-language-to-planning","title":"Translating Natural Language to Planning Goals with Large-Language Models","date":"2023-02-10","arxiv_id":"2302.05128","n_code_links":1,"syntology":null},{"paper":null,"slug":"better-by-you-better-than-me-chatgpt3-as","title":"Better by you, better than me, chatgpt3 as writing assistance in students essays","date":"2023-02-09","arxiv_id":"2302.04536","n_code_links":0,"syntology":null},{"paper":"/paper/binarized-neural-machine-translation-1","slug":"binarized-neural-machine-translation-1","title":"Binarized Neural Machine Translation","date":"2023-02-09","arxiv_id":"2302.04907","n_code_links":1,"syntology":{"ran":4,"of":6,"n_ran_checked":4,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["google/aqt"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"flexible-model-agnostic-method-for-materials","title":"Flexible, Model-Agnostic Method for Materials Data Extraction from Text Using General Purpose Language Models","date":"2023-02-09","arxiv_id":"2302.04914","n_code_links":0,"syntology":null},{"paper":"/paper/generating-a-structured-summary-of-numerous","slug":"generating-a-structured-summary-of-numerous","title":"Generating a Structured Summary of Numerous Academic Papers: Dataset and Method","date":"2023-02-09","arxiv_id":"2302.04580","n_code_links":1,"syntology":null},{"paper":"/paper/hybrik-transformer","slug":"hybrik-transformer","title":"3D Human Pose and Shape Estimation via HybrIK-Transformer","date":"2023-02-09","arxiv_id":"2302.04774","n_code_links":1,"syntology":null},{"paper":"/paper/reversible-vision-transformers-1","slug":"reversible-vision-transformers-1","title":"Reversible Vision Transformers","date":"2023-02-09","arxiv_id":"2302.04869","n_code_links":4,"syntology":{"ran":24,"of":29,"n_ran_checked":23,"n_instrument":1,"unverified":5,"pointer_only":26,"phrase":"24 ran (of which 5 constructed an object rather than computing a result; 23 with no instrument failure: 0 honoured, 0 violated, 23 with no contract checked; 1 where Syntology's instrument failed) · 5 unverified","official":{"repos":["karttikeya/minrev","facebookresearch/SlowFast","facebookresearch/mvit"],"state":"official (archive's flag): 22 ran","n_ran":22,"n_constructed":5,"n_ran_no_instrument_failure":21,"n_unverified":4,"ran_from_kinds":["listed","official"]}}},{"paper":"/paper/adapting-pre-trained-vision-transformers-from","slug":"adapting-pre-trained-vision-transformers-from","title":"Adapting Pre-trained Vision Transformers from 2D to 3D through Weight Inflation Improves Medical Image Segmentation","date":"2023-02-08","arxiv_id":"2302.04303","n_code_links":1,"syntology":null},{"paper":null,"slug":"an-empirical-study-of-uniform-architecture","title":"An Empirical Study of Uniform-Architecture Knowledge Distillation in Document Ranking","date":"2023-02-08","arxiv_id":"2302.04112","n_code_links":0,"syntology":null},{"paper":"/paper/attending-to-graph-transformers","slug":"attending-to-graph-transformers","title":"Attending to Graph Transformers","date":"2023-02-08","arxiv_id":"2302.04181","n_code_links":1,"syntology":{"ran":5,"of":8,"n_ran_checked":5,"n_instrument":0,"unverified":3,"pointer_only":4,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":{"repos":["luis-mueller/probing-graph-transformers"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"auto-learning-an-adversarial-process-of-two","title":"EvoText: Enhancing Natural Language Generation Models via Self-Escalation Learning for Up-to-Date Knowledge and Improved Performance","date":"2023-02-08","arxiv_id":"2302.03896","n_code_links":0,"syntology":null},{"paper":null,"slug":"crl-a-novel-semi-supervised-deep-active","title":"CRL+: A Novel Semi-Supervised Deep Active Contrastive Representation Learning-Based Text Classification Model for Insurance Data","date":"2023-02-08","arxiv_id":"2302.04343","n_code_links":0,"syntology":null},{"paper":"/paper/dual-interest-factorization-heads-attention","slug":"dual-interest-factorization-heads-attention","title":"Dual-interest Factorization-heads Attention for Sequential Recommendation","date":"2023-02-08","arxiv_id":"2302.03965","n_code_links":1,"syntology":null},{"paper":"/paper/efficient-joint-learning-for-clinical-named","slug":"efficient-joint-learning-for-clinical-named","title":"Efficient Joint Learning for Clinical Named Entity Recognition and Relation Extraction Using Fourier Networks: A Use Case in Adverse Drug Events","date":"2023-02-08","arxiv_id":"2302.04185","n_code_links":1,"syntology":null},{"paper":"/paper/prompting-for-multimodal-hateful-meme","slug":"prompting-for-multimodal-hateful-meme","title":"Prompting for Multimodal Hateful Meme Classification","date":"2023-02-08","arxiv_id":"2302.04156","n_code_links":0,"syntology":null},{"paper":null,"slug":"swincross-cross-modal-swin-transformer-for","title":"SwinCross: Cross-modal Swin Transformer for Head-and-Neck Tumor Segmentation in PET/CT Images","date":"2023-02-08","arxiv_id":"2302.03861","n_code_links":0,"syntology":null},{"paper":null,"slug":"lut-nn-towards-unified-neural-network","title":"LUT-NN: Empower Efficient Neural Network Inference with Centroid Learning and Table Lookup","date":"2023-02-07","arxiv_id":"2302.03213","n_code_links":0,"syntology":null},{"paper":"/paper/osrt-omnidirectional-image-super-resolution","slug":"osrt-omnidirectional-image-super-resolution","title":"OSRT: Omnidirectional Image Super-Resolution with Distortion-aware Transformer","date":"2023-02-07","arxiv_id":"2302.03453","n_code_links":1,"syntology":{"ran":11,"of":12,"n_ran_checked":10,"n_instrument":1,"unverified":1,"pointer_only":1,"phrase":"11 ran (of which 0 constructed an object rather than computing a result; 10 with no instrument failure: 0 honoured, 0 violated, 10 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":{"repos":["fanghua-yu/osrt"],"state":"official (archive's flag): 11 ran","n_ran":11,"n_constructed":0,"n_ran_no_instrument_failure":10,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"reliable-natural-language-understanding-with","title":"Reliable Natural Language Understanding with Large Language Models and Answer Set Programming","date":"2023-02-07","arxiv_id":"2302.03780","n_code_links":0,"syntology":null},{"paper":null,"slug":"what-do-language-models-know-about-word","title":"What do Language Models know about word senses? Zero-Shot WSD with Language Models and Domain Inventories","date":"2023-02-07","arxiv_id":"2302.03353","n_code_links":0,"syntology":null},{"paper":"/paper/what-matters-in-the-structured-pruning-of","slug":"what-matters-in-the-structured-pruning-of","title":"What Matters In The Structured Pruning of Generative Language Models?","date":"2023-02-07","arxiv_id":"2302.03773","n_code_links":1,"syntology":{"ran":2,"of":8,"n_ran_checked":2,"n_instrument":0,"unverified":6,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 6 unverified","official":{"repos":["huggingface/nn_pruning"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":6,"ran_from_kinds":["official"]}}},{"paper":"/paper/aim-adapting-image-models-for-efficient-video","slug":"aim-adapting-image-models-for-efficient-video","title":"AIM: Adapting Image Models for Efficient Video Action Recognition","date":"2023-02-06","arxiv_id":"2302.03024","n_code_links":1,"syntology":null},{"paper":null,"slug":"context-gloss-augmentation-for-improving","title":"Context-Gloss Augmentation for Improving Arabic Target Sense Verification","date":"2023-02-06","arxiv_id":"2302.03126","n_code_links":0,"syntology":null},{"paper":"/paper/controllable-lexical-simplification-for","slug":"controllable-lexical-simplification-for","title":"Controllable Lexical Simplification for English","date":"2023-02-06","arxiv_id":"2302.02900","n_code_links":1,"syntology":null},{"paper":"/paper/gps-reviving-the-art-of-message-passing-for","slug":"gps-reviving-the-art-of-message-passing-for","title":"GPS++: Reviving the Art of Message Passing for Molecular Property Prediction","date":"2023-02-06","arxiv_id":"2302.02947","n_code_links":1,"syntology":null},{"paper":"/paper/v1t-large-scale-mouse-v1-response-prediction","slug":"v1t-large-scale-mouse-v1-response-prediction","title":"V1T: large-scale mouse V1 response prediction using a Vision Transformer","date":"2023-02-06","arxiv_id":"2302.03023","n_code_links":1,"syntology":{"ran":16,"of":17,"n_ran_checked":16,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"16 ran (of which 0 constructed an object rather than computing a result; 16 with no instrument failure: 0 honoured, 0 violated, 16 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["bryanlimy/V1T"],"state":"official (archive's flag): 16 ran","n_ran":16,"n_constructed":0,"n_ran_no_instrument_failure":16,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"cross-modal-fusion-techniques-for-utterance","title":"cross-modal fusion techniques for utterance-level emotion recognition from text and speech","date":"2023-02-05","arxiv_id":"2302.02447","n_code_links":0,"syntology":null},{"paper":"/paper/kdeformer-accelerating-transformers-via","slug":"kdeformer-accelerating-transformers-via","title":"KDEformer: Accelerating Transformers via Kernel Density Estimation","date":"2023-02-05","arxiv_id":"2302.02451","n_code_links":1,"syntology":{"ran":8,"of":12,"n_ran_checked":1,"n_instrument":7,"unverified":4,"pointer_only":0,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 7 where Syntology's instrument failed) · 4 unverified","official":{"repos":["majid-daliri/kdeformer"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"nationality-bias-in-text-generation","title":"Nationality Bias in Text Generation","date":"2023-02-05","arxiv_id":"2302.02463","n_code_links":0,"syntology":null},{"paper":null,"slug":"quantized-distributed-training-of-large","title":"Quantized Distributed Training of Large Models with Convergence Guarantees","date":"2023-02-05","arxiv_id":"2302.02390","n_code_links":0,"syntology":null},{"paper":null,"slug":"vulaste-long-sequence-model-with-abstract","title":"VuLASTE: Long Sequence Model with Abstract Syntax Tree Embedding for vulnerability Detection","date":"2023-02-05","arxiv_id":"2302.02345","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-new-cross-domain-strategy-based-xai-models","title":"A New cross-domain strategy based XAI models for fake news detection","date":"2023-02-04","arxiv_id":"2302.02122","n_code_links":0,"syntology":null},{"paper":null,"slug":"knowledge-graph-completion-method-combined","title":"Knowledge Graph Completion Method Combined With Adaptive Enhanced Semantic Information","date":"2023-02-04","arxiv_id":"2302.02116","n_code_links":0,"syntology":null},{"paper":null,"slug":"learning-to-agree-on-vision-attention-for","title":"Learning to Agree on Vision Attention for Visual Commonsense Reasoning","date":"2023-02-04","arxiv_id":"2302.02117","n_code_links":0,"syntology":null},{"paper":"/paper/realtabformer-generating-realistic-relational","slug":"realtabformer-generating-realistic-relational","title":"REaLTabFormer: Generating Realistic Relational and Tabular Data using Transformers","date":"2023-02-04","arxiv_id":"2302.02041","n_code_links":3,"syntology":{"ran":3,"of":4,"n_ran_checked":3,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["avsolatorio/realtabformer"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/revisiting-image-deblurring-with-an-efficient","slug":"revisiting-image-deblurring-with-an-efficient","title":"Revisiting Image Deblurring with an Efficient ConvNet","date":"2023-02-04","arxiv_id":"2302.02234","n_code_links":1,"syntology":null},{"paper":null,"slug":"theory-of-mind-may-have-spontaneously-emerged","title":"Evaluating Large Language Models in Theory of Mind Tasks","date":"2023-02-04","arxiv_id":"2302.02083","n_code_links":0,"syntology":null},{"paper":null,"slug":"weight-is-attention-all-we-need-aeiuorder","title":"Greedy Ordering of Layer Weight Matrices in Transformers Improves Translation","date":"2023-02-04","arxiv_id":"2302.02123","n_code_links":0,"syntology":null},{"paper":"/paper/antm-an-aligned-neural-topic-model-for","slug":"antm-an-aligned-neural-topic-model-for","title":"ANTM: An Aligned Neural Topic Model for Exploring Evolving Topics","date":"2023-02-03","arxiv_id":"2302.01501","n_code_links":1,"syntology":null},{"paper":"/paper/bioformer-an-efficient-transformer-language","slug":"bioformer-an-efficient-transformer-language","title":"Bioformer: an efficient transformer language model for biomedical text mining","date":"2023-02-03","arxiv_id":"2302.01588","n_code_links":1,"syntology":null},{"paper":null,"slug":"cfft-gan-cross-domain-feature-fusion","title":"CFFT-GAN: Cross-domain Feature Fusion Transformer for Exemplar-based Image Translation","date":"2023-02-03","arxiv_id":"2302.01608","n_code_links":0,"syntology":null},{"paper":null,"slug":"coinductive-guide-to-inductive-transformer","title":"Coinductive guide to inductive transformer heads","date":"2023-02-03","arxiv_id":"2302.01834","n_code_links":0,"syntology":null},{"paper":null,"slug":"detecting-reddit-users-with-depression-using","title":"Detecting Reddit Users with Depression Using a Hybrid Neural Network SBERT-CNN","date":"2023-02-03","arxiv_id":"2302.02759","n_code_links":0,"syntology":null},{"paper":null,"slug":"device-depth-and-visual-concepts-aware","title":"DEVICE: DEpth and VIsual ConcEpts Aware Transformer for TextCaps","date":"2023-02-03","arxiv_id":"2302.01540","n_code_links":0,"syntology":null},{"paper":"/paper/dilateformer-multi-scale-dilated-transformer","slug":"dilateformer-multi-scale-dilated-transformer","title":"DilateFormer: Multi-Scale Dilated Transformer for Visual Recognition","date":"2023-02-03","arxiv_id":"2302.01791","n_code_links":1,"syntology":null},{"paper":"/paper/hdformer-high-order-directed-transformer-for","slug":"hdformer-high-order-directed-transformer-for","title":"HDFormer: High-order Directed Transformer for 3D Human Pose Estimation","date":"2023-02-03","arxiv_id":"2302.01825","n_code_links":1,"syntology":{"ran":6,"of":12,"n_ran_checked":4,"n_instrument":2,"unverified":6,"pointer_only":12,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 2 honoured, 1 violated, 1 with no contract checked; 2 where Syntology's instrument failed) · 6 unverified","official":{"repos":["hyer/hdformer"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":6,"ran_from_kinds":["official"]}}},{"paper":"/paper/revisiting-intermediate-layer-distillation","slug":"revisiting-intermediate-layer-distillation","title":"Revisiting Intermediate Layer Distillation for Compressing Language Models: An Overfitting Perspective","date":"2023-02-03","arxiv_id":"2302.01530","n_code_links":1,"syntology":null},{"paper":null,"slug":"towards-few-shot-identification-of-morality","title":"Towards Few-Shot Identification of Morality Frames using In-Context Learning","date":"2023-02-03","arxiv_id":"2302.02029","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-survey-on-efficient-training-of","title":"A Survey on Efficient Training of Transformers","date":"2023-02-02","arxiv_id":"2302.01107","n_code_links":0,"syntology":null},{"paper":null,"slug":"boosting-low-data-instance-segmentation-by","title":"Boosting Low-Data Instance Segmentation by Unsupervised Pre-training with Saliency Prompt","date":"2023-02-02","arxiv_id":"2302.01171","n_code_links":0,"syntology":null},{"paper":null,"slug":"creating-a-large-language-model-of-a","title":"Creating a Large Language Model of a Philosopher","date":"2023-02-02","arxiv_id":"2302.01339","n_code_links":0,"syntology":null},{"paper":null,"slug":"curriculum-guided-abstractive-summarization-1","title":"Curriculum-Guided Abstractive Summarization","date":"2023-02-02","arxiv_id":"2302.01342","n_code_links":0,"syntology":null},{"paper":"/paper/dual-patchnorm","slug":"dual-patchnorm","title":"Dual PatchNorm","date":"2023-02-02","arxiv_id":"2302.01327","n_code_links":7,"syntology":{"ran":7,"of":9,"n_ran_checked":6,"n_instrument":1,"unverified":2,"pointer_only":2,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 1 honoured, 4 violated, 1 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","official":{"repos":["google-research/big_vision"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/fcb-swinv2-transformer-for-polyp-segmentation","slug":"fcb-swinv2-transformer-for-polyp-segmentation","title":"FCB-SwinV2 Transformer for Polyp Segmentation","date":"2023-02-02","arxiv_id":"2302.01027","n_code_links":0,"syntology":null},{"paper":null,"slug":"history-aware-hierarchical-transformer-for","title":"History-Aware Hierarchical Transformer for Multi-session Open-domain Dialogue System","date":"2023-02-02","arxiv_id":"2302.00907","n_code_links":0,"syntology":null},{"paper":null,"slug":"idt5-indonesian-version-of-multilingual-t5","title":"idT5: Indonesian Version of Multilingual T5 Transformer","date":"2023-02-02","arxiv_id":"2302.00856","n_code_links":0,"syntology":null},{"paper":"/paper/language-quantized-autoencoders-towards-1","slug":"language-quantized-autoencoders-towards-1","title":"Language Quantized AutoEncoders: Towards Unsupervised Text-Image Alignment","date":"2023-02-02","arxiv_id":"2302.00902","n_code_links":1,"syntology":null},{"paper":"/paper/longformer-longitudinal-transformer-for","slug":"longformer-longitudinal-transformer-for","title":"Longformer: Longitudinal Transformer for Alzheimer's Disease Classification with Structural MRIs","date":"2023-02-02","arxiv_id":"2302.00901","n_code_links":1,"syntology":null},{"paper":null,"slug":"mnemosyne-learning-to-train-transformers-with","title":"Mnemosyne: Learning to Train Transformers with Transformers","date":"2023-02-02","arxiv_id":"2302.01128","n_code_links":0,"syntology":null},{"paper":null,"slug":"molecular-geometry-aware-transformer-for","title":"Molecular Geometry-aware Transformer for accurate 3D Atomic System modeling","date":"2023-02-02","arxiv_id":"2302.00855","n_code_links":0,"syntology":null},{"paper":"/paper/resilient-binary-neural-network","slug":"resilient-binary-neural-network","title":"Resilient Binary Neural Network","date":"2023-02-02","arxiv_id":"2302.00956","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified; the one sample that ran constructed an object rather than computing a result","official":{"repos":["stevetsui/rebnn"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/semantic-coherence-markers-for-the-early","slug":"semantic-coherence-markers-for-the-early","title":"Semantic Coherence Markers for the Early Diagnosis of the Alzheimer Disease","date":"2023-02-02","arxiv_id":"2302.01025","n_code_links":1,"syntology":null},{"paper":null,"slug":"vision-transformer-based-feature-extraction","title":"Vision Transformer-based Feature Extraction for Generalized Zero-Shot Learning","date":"2023-02-02","arxiv_id":"2302.00875","n_code_links":0,"syntology":null},{"paper":null,"slug":"what-language-reveals-about-perception","title":"Large language models predict human sensory judgments across six modalities","date":"2023-02-02","arxiv_id":"2302.01308","n_code_links":0,"syntology":null},{"paper":null,"slug":"an-empirical-study-on-the-transferability-of","title":"An Empirical Study on the Transferability of Transformer Modules in Parameter-Efficient Fine-Tuning","date":"2023-02-01","arxiv_id":"2302.00378","n_code_links":0,"syntology":null},{"paper":"/paper/analyzing-leakage-of-personally-identifiable","slug":"analyzing-leakage-of-personally-identifiable","title":"Analyzing Leakage of Personally Identifiable Information in Language Models","date":"2023-02-01","arxiv_id":"2302.00539","n_code_links":1,"syntology":null},{"paper":null,"slug":"clinical-decision-transformer-intended","title":"Clinical Decision Transformer: Intended Treatment Recommendation through Goal Prompting","date":"2023-02-01","arxiv_id":"2302.00612","n_code_links":0,"syntology":null},{"paper":null,"slug":"co-writing-with-opinionated-language-models","title":"Co-Writing with Opinionated Language Models Affects Users' Views","date":"2023-02-01","arxiv_id":"2302.00560","n_code_links":0,"syntology":null},{"paper":"/paper/feed-forward-blocks-control-contextualization","slug":"feed-forward-blocks-control-contextualization","title":"Analyzing Feed-Forward Blocks in Transformers through the Lens of Attention Maps","date":"2023-02-01","arxiv_id":"2302.00456","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":1,"n_instrument":1,"unverified":0,"pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["gorokoba560/norm-analysis-of-transformer"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/hunsum-1-an-abstractive-summarization-dataset","slug":"hunsum-1-an-abstractive-summarization-dataset","title":"HunSum-1: an Abstractive Summarization Dataset for Hungarian","date":"2023-02-01","arxiv_id":"2302.00455","n_code_links":1,"syntology":null},{"paper":null,"slug":"image-based-vehicle-classification-by","title":"Image-Based Vehicle Classification by Synergizing Features from Supervised and Self-Supervised Learning Paradigms","date":"2023-02-01","arxiv_id":"2302.00648","n_code_links":0,"syntology":null},{"paper":"/paper/improving-few-shot-generalization-by-1","slug":"improving-few-shot-generalization-by-1","title":"Improving Few-Shot Generalization by Exploring and Exploiting Auxiliary Data","date":"2023-02-01","arxiv_id":"2302.00674","n_code_links":1,"syntology":null},{"paper":"/paper/multispectral-pedestrian-detection-via-1","slug":"multispectral-pedestrian-detection-via-1","title":"MS-DETR: Multispectral Pedestrian Detection Transformer with Loosely Coupled Fusion and Modality-Balanced Optimization","date":"2023-02-01","arxiv_id":"2302.00290","n_code_links":1,"syntology":null},{"paper":"/paper/turning-the-curse-of-heterogeneity-in","slug":"turning-the-curse-of-heterogeneity-in","title":"Turning the Curse of Heterogeneity in Federated Learning into a Blessing for Out-of-Distribution Detection","date":"2023-02-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"an-comparative-analysis-of-different-pitch","title":"An Comparative Analysis of Different Pitch and Metrical Grid Encoding Methods in the Task of Sequential Music Generation","date":"2023-01-31","arxiv_id":"2301.13383","n_code_links":0,"syntology":null},{"paper":"/paper/continuous-spatiotemporal-transformers","slug":"continuous-spatiotemporal-transformers","title":"Continuous Spatiotemporal Transformers","date":"2023-01-31","arxiv_id":"2301.13338","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":3,"n_instrument":0,"unverified":0,"pointer_only":2,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 1 honoured, 2 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["vandijklab/cst"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official","unlocated"]}}},{"paper":"/paper/fairness-aware-vision-transformer-via","slug":"fairness-aware-vision-transformer-via","title":"Fairness-aware Vision Transformer via Debiased Self-Attention","date":"2023-01-31","arxiv_id":"2301.13803","n_code_links":1,"syntology":null},{"paper":null,"slug":"flame-a-small-language-model-for-spreadsheet","title":"FLAME: A small language model for spreadsheet formulas","date":"2023-01-31","arxiv_id":"2301.13779","n_code_links":0,"syntology":null},{"paper":null,"slug":"numeracy-from-literacy-data-science-as-an","title":"Numeracy from Literacy: Data Science as an Emergent Skill from Large Language Models","date":"2023-01-31","arxiv_id":"2301.13382","n_code_links":0,"syntology":null},{"paper":null,"slug":"skill-decision-transformer","title":"Skill Decision Transformer","date":"2023-01-31","arxiv_id":"2301.13573","n_code_links":0,"syntology":null}],"record_sha256":"decb5a08c54d5066d01d2715d3a279b943c321404852974433687a097b407e89","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}