{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/multi-head-attention/papers/157","list_of":"/method/multi-head-attention","method":"Multi-Head Attention","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":157,"pages_in_order":249,"rows_per_page":100,"rows":[15601,15700],"of":24855,"counts":{"archive_papers_tagged":24855,"with_a_code_link":11214,"where_syntology_ran_a_sample":3454,"not_listed_spam_title":0,"listed":24855,"listed_where_code_ran":3454,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":2916,"every_run_a_failure_of_syntologys_instrument":538,"listed_with_a_run_with_no_instrument_failure":2916,"listed_every_run_a_failure_of_syntologys_instrument":538,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/multi-head-attention","prev":"/method/multi-head-attention/papers/156","next":"/method/multi-head-attention/papers/158","papers":[{"paper":null,"slug":"pair-detr-contrastive-learning-speeds-up-detr","title":"Pair DETR: Contrastive Learning Speeds Up DETR Training","date":"2022-10-29","arxiv_id":"2210.16476","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-long-term-dependent-and-trustworthy","title":"A Long-term Dependent and Trustworthy Approach to Reactor Accident Prognosis based on Temporal Fusion Transformer","date":"2022-10-28","arxiv_id":"2210.17298","n_code_links":0,"syntology":null},{"paper":"/paper/bebert-efficient-and-robust-binary-ensemble","slug":"bebert-efficient-and-robust-binary-ensemble","title":"BEBERT: Efficient and Robust Binary Ensemble BERT","date":"2022-10-28","arxiv_id":"2210.15976","n_code_links":1,"syntology":null},{"paper":"/paper/contextual-learning-in-fourier-complex-field","slug":"contextual-learning-in-fourier-complex-field","title":"Contextual Learning in Fourier Complex Field for VHR Remote Sensing Images","date":"2022-10-28","arxiv_id":"2210.15972","n_code_links":3,"syntology":null},{"paper":null,"slug":"differentially-private-cutmix-for-split","title":"Differentially Private CutMix for Split Learning with Vision Transformer","date":"2022-10-28","arxiv_id":"2210.15986","n_code_links":0,"syntology":null},{"paper":null,"slug":"dimensionality-reduced-antenna-array-for","title":"Dimensionality Reduced Antenna Array for Beamforming/steering","date":"2022-10-28","arxiv_id":"2210.16197","n_code_links":0,"syntology":null},{"paper":"/paper/efficient-speech-translation-with-dynamic","slug":"efficient-speech-translation-with-dynamic","title":"Efficient Speech Translation with Dynamic Latent Perceivers","date":"2022-10-28","arxiv_id":"2210.16264","n_code_links":1,"syntology":null},{"paper":null,"slug":"elastic-weight-consolidation-improves-the","title":"Elastic Weight Consolidation Improves the Robustness of Self-Supervised Learning Methods under Transfer","date":"2022-10-28","arxiv_id":"2210.16365","n_code_links":0,"syntology":null},{"paper":"/paper/exploring-spatial-temporal-features-for","slug":"exploring-spatial-temporal-features-for","title":"Exploring Spatial-Temporal Features for Deepfake Detection and Localization","date":"2022-10-28","arxiv_id":"2210.15872","n_code_links":1,"syntology":null},{"paper":null,"slug":"feature-engineering-vs-bert-on-twitter-data","title":"Feature Engineering vs BERT on Twitter Data","date":"2022-10-28","arxiv_id":"2210.16168","n_code_links":0,"syntology":null},{"paper":null,"slug":"federated-learning-for-chronic-obstructive","title":"Federated Learning for Chronic Obstructive Pulmonary Disease Classification with Partial Personalized Attention Mechanism","date":"2022-10-28","arxiv_id":"2210.16142","n_code_links":0,"syntology":null},{"paper":null,"slug":"grafting-vision-transformers","title":"Grafting Vision Transformers","date":"2022-10-28","arxiv_id":"2210.15943","n_code_links":0,"syntology":null},{"paper":null,"slug":"modeling-structure-building-in-the-brain-with","title":"Modeling structure-building in the brain with CCG parsing and large language models","date":"2022-10-28","arxiv_id":"2210.16147","n_code_links":0,"syntology":null},{"paper":null,"slug":"on-the-use-of-modality-specific-large-scale","title":"On the Use of Modality-Specific Large-Scale Pre-Trained Encoders for Multimodal Sentiment Analysis","date":"2022-10-28","arxiv_id":"2210.15937","n_code_links":0,"syntology":null},{"paper":null,"slug":"parameter-efficient-transfer-learning-of-pre","title":"Parameter-efficient transfer learning of pre-trained Transformer models for speaker verification using adapters","date":"2022-10-28","arxiv_id":"2210.16032","n_code_links":0,"syntology":null},{"paper":"/paper/probing-for-targeted-syntactic-knowledge","slug":"probing-for-targeted-syntactic-knowledge","title":"Probing for targeted syntactic knowledge through grammatical error detection","date":"2022-10-28","arxiv_id":"2210.16228","n_code_links":1,"syntology":null},{"paper":null,"slug":"psformer-point-transformer-for-3d-salient","title":"PSFormer: Point Transformer for 3D Salient Object Detection","date":"2022-10-28","arxiv_id":"2210.15933","n_code_links":0,"syntology":null},{"paper":null,"slug":"softbart-soft-bayesian-additive-regression","title":"SoftBart: Soft Bayesian Additive Regression Trees","date":"2022-10-28","arxiv_id":"2210.16375","n_code_links":0,"syntology":null},{"paper":null,"slug":"upainting-unified-text-to-image-diffusion","title":"UPainting: Unified Text-to-Image Diffusion Generation with Cross-modal Guidance","date":"2022-10-28","arxiv_id":"2210.16031","n_code_links":0,"syntology":null},{"paper":"/paper/vlt-vision-language-transformer-and-query","slug":"vlt-vision-language-transformer-and-query","title":"VLT: Vision-Language Transformer and Query Generation for Referring Segmentation","date":"2022-10-28","arxiv_id":"2210.15871","n_code_links":1,"syntology":{"ran":3,"of":6,"n_ran_checked":3,"n_instrument":0,"unverified":3,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":{"repos":["henghuiding/Vision-Language-Transformer"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"bert-flow-vae-a-weakly-supervised-model-for-1","title":"BERT-Flow-VAE: A Weakly-supervised Model for Multi-Label Text Classification","date":"2022-10-27","arxiv_id":"2210.15225","n_code_links":0,"syntology":null},{"paper":"/paper/coco-dr-combating-distribution-shifts-in-zero","slug":"coco-dr-combating-distribution-shifts-in-zero","title":"COCO-DR: Combating Distribution Shifts in Zero-Shot Dense Retrieval with Contrastive and Distributionally Robust Learning","date":"2022-10-27","arxiv_id":"2210.15212","n_code_links":1,"syntology":{"ran":1,"of":2,"n_ran_checked":1,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"1 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified; the one sample that ran constructed an object rather than computing a result","official":{"repos":["openmatch/coco-dr"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/cost-eff-collaborative-optimization-of","slug":"cost-eff-collaborative-optimization-of","title":"COST-EFF: Collaborative Optimization of Spatial and Temporal Efficiency with Slenderized Multi-exit Language Models","date":"2022-10-27","arxiv_id":"2210.15523","n_code_links":1,"syntology":null},{"paper":"/paper/fast-distilbert-on-cpus","slug":"fast-distilbert-on-cpus","title":"Fast DistilBERT on CPUs","date":"2022-10-27","arxiv_id":"2211.07715","n_code_links":1,"syntology":null},{"paper":"/paper/fctalker-fine-and-coarse-grained-context","slug":"fctalker-fine-and-coarse-grained-context","title":"FCTalker: Fine and Coarse Grained Context Modeling for Expressive Conversational Speech Synthesis","date":"2022-10-27","arxiv_id":"2210.15360","n_code_links":1,"syntology":null},{"paper":"/paper/gaitmixer-skeleton-based-gait-representation","slug":"gaitmixer-skeleton-based-gait-representation","title":"GaitMixer: Skeleton-based Gait Representation Learning via Wide-spectrum Multi-axial Mixer","date":"2022-10-27","arxiv_id":"2210.15491","n_code_links":1,"syntology":null},{"paper":null,"slug":"hydra-hgr-a-hybrid-transformer-based","title":"HYDRA-HGR: A Hybrid Transformer-based Architecture for Fusion of Macroscopic and Microscopic Neural Drive Information","date":"2022-10-27","arxiv_id":"2211.02619","n_code_links":0,"syntology":null},{"paper":null,"slug":"li3detr-a-lidar-based-3d-detection","title":"Li3DeTr: A LiDAR based 3D Detection Transformer","date":"2022-10-27","arxiv_id":"2210.15365","n_code_links":0,"syntology":null},{"paper":null,"slug":"make-more-of-your-data-minimal-effort-data","title":"Make More of Your Data: Minimal Effort Data Augmentation for Automatic Speech Recognition and Translation","date":"2022-10-27","arxiv_id":"2210.15398","n_code_links":0,"syntology":null},{"paper":null,"slug":"masked-transformer-for-image-anomaly","title":"Masked Transformer for image Anomaly Localization","date":"2022-10-27","arxiv_id":"2210.15540","n_code_links":0,"syntology":null},{"paper":"/paper/masked-vision-language-transformer-in-fashion","slug":"masked-vision-language-transformer-in-fashion","title":"Masked Vision-Language Transformer in Fashion","date":"2022-10-27","arxiv_id":"2210.15110","n_code_links":1,"syntology":null},{"paper":null,"slug":"msf3ddetr-multi-sensor-fusion-3d-detection","title":"MSF3DDETR: Multi-Sensor Fusion 3D Detection Transformer for Autonomous Driving","date":"2022-10-27","arxiv_id":"2210.15316","n_code_links":0,"syntology":null},{"paper":"/paper/multimodal-transformer-distillation-for-audio","slug":"multimodal-transformer-distillation-for-audio","title":"Multimodal Transformer Distillation for Audio-Visual Synchronization","date":"2022-10-27","arxiv_id":"2210.15563","n_code_links":2,"syntology":null},{"paper":"/paper/point-voxel-adaptive-feature-abstraction-for","slug":"point-voxel-adaptive-feature-abstraction-for","title":"Point-Voxel Adaptive Feature Abstraction for Robust Point Cloud Classification","date":"2022-10-27","arxiv_id":"2210.15514","n_code_links":1,"syntology":null},{"paper":"/paper/procontext-exploring-progressive-context","slug":"procontext-exploring-progressive-context","title":"ProContEXT: Exploring Progressive Context Transformer for Tracking","date":"2022-10-27","arxiv_id":"2210.15511","n_code_links":4,"syntology":null},{"paper":null,"slug":"spatio-temporal-hybrid-fusion-of-cae-and-swin","title":"Spatio-Temporal Hybrid Fusion of CAE and SWIn Transformers for Lung Cancer Malignancy Prediction","date":"2022-10-27","arxiv_id":"2210.15297","n_code_links":0,"syntology":null},{"paper":"/paper/the-1st-place-solution-for-eccv-2022-multiple","slug":"the-1st-place-solution-for-eccv-2022-multiple","title":"The 1st-place Solution for ECCV 2022 Multiple People Tracking in Group Dance Challenge","date":"2022-10-27","arxiv_id":"2210.15281","n_code_links":3,"syntology":null},{"paper":"/paper/transformers-meet-stochastic-block-models","slug":"transformers-meet-stochastic-block-models","title":"Transformers meet Stochastic Block Models: Attention with Data-Adaptive Sparsity and Cost","date":"2022-10-27","arxiv_id":"2210.15541","n_code_links":1,"syntology":null},{"paper":null,"slug":"trscore-a-novel-gpt-based-readability-scorer","title":"TRScore: A Novel GPT-based Readability Scorer for ASR Segmentation and Punctuation model evaluation and selection","date":"2022-10-27","arxiv_id":"2210.15104","n_code_links":0,"syntology":null},{"paper":"/paper/unsupervised-boundary-aware-language-model","slug":"unsupervised-boundary-aware-language-model","title":"Unsupervised Boundary-Aware Language Model Pretraining for Chinese Sequence Labeling","date":"2022-10-27","arxiv_id":"2210.15231","n_code_links":2,"syntology":null},{"paper":"/paper/what-language-model-to-train-if-you-have-one","slug":"what-language-model-to-train-if-you-have-one","title":"What Language Model to Train if You Have One Million GPU Hours?","date":"2022-10-27","arxiv_id":"2210.15424","n_code_links":1,"syntology":null},{"paper":"/paper/working-alliance-transformer-for","slug":"working-alliance-transformer-for","title":"Working Alliance Transformer for Psychotherapy Dialogue Classification","date":"2022-10-27","arxiv_id":"2210.15603","n_code_links":1,"syntology":null},{"paper":null,"slug":"automatic-diagnosis-of-myocarditis-disease-in","title":"Automatic Diagnosis of Myocarditis Disease in Cardiac MRI Modality using Deep Transformers and Explainable Artificial Intelligence","date":"2022-10-26","arxiv_id":"2210.14611","n_code_links":0,"syntology":null},{"paper":"/paper/automatic-extraction-of-materials-and","slug":"automatic-extraction-of-materials-and","title":"Automatic extraction of materials and properties from superconductors scientific literature","date":"2022-10-26","arxiv_id":"2210.15600","n_code_links":2,"syntology":null},{"paper":null,"slug":"beyond-english-centric-bitexts-for-better","title":"Beyond English-Centric Bitexts for Better Multilingual Language Representation Learning","date":"2022-10-26","arxiv_id":"2210.14867","n_code_links":0,"syntology":null},{"paper":null,"slug":"bi-link-bridging-inductive-link-predictions","title":"Bi-Link: Bridging Inductive Link Predictions from Text via Contrastive Learning of Transformers and Prompts","date":"2022-10-26","arxiv_id":"2210.14463","n_code_links":0,"syntology":null},{"paper":"/paper/disentangling-past-future-modeling-in","slug":"disentangling-past-future-modeling-in","title":"Disentangling Past-Future Modeling in Sequential Recommendation via Dual Networks","date":"2022-10-26","arxiv_id":"2210.14577","n_code_links":1,"syntology":null},{"paper":null,"slug":"don-t-prompt-search-mining-based-zero-shot","title":"Don't Prompt, Search! Mining-based Zero-Shot Learning with Language Models","date":"2022-10-26","arxiv_id":"2210.14803","n_code_links":0,"syntology":null},{"paper":"/paper/eeny-meeny-miny-moe-how-to-choose-data-for","slug":"eeny-meeny-miny-moe-how-to-choose-data-for","title":"Eeny, meeny, miny, moe. How to choose data for morphological inflection","date":"2022-10-26","arxiv_id":"2210.14465","n_code_links":1,"syntology":null},{"paper":null,"slug":"exploring-robustness-of-prefix-tuning-in","title":"Exploring Robustness of Prefix Tuning in Noisy Data: A Case Study in Financial Sentiment Analysis","date":"2022-10-26","arxiv_id":"2211.05584","n_code_links":0,"syntology":null},{"paper":null,"slug":"learning-a-task-specific-descriptor-for","title":"Learning a Task-specific Descriptor for Robust Matching of 3D Point Clouds","date":"2022-10-26","arxiv_id":"2210.14899","n_code_links":0,"syntology":null},{"paper":"/paper/leveraging-affirmative-interpretations-from","slug":"leveraging-affirmative-interpretations-from","title":"Leveraging Affirmative Interpretations from Negation Improves Natural Language Understanding","date":"2022-10-26","arxiv_id":"2210.14486","n_code_links":1,"syntology":null},{"paper":"/paper/leveraging-demonstrations-with-latent-space","slug":"leveraging-demonstrations-with-latent-space","title":"Leveraging Demonstrations with Latent Space Priors","date":"2022-10-26","arxiv_id":"2210.14685","n_code_links":1,"syntology":{"ran":2,"of":3,"n_ran_checked":2,"n_instrument":0,"unverified":1,"pointer_only":3,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 1 honoured, 1 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["facebookresearch/latent-space-priors"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/m-3-vit-mixture-of-experts-vision-transformer","slug":"m-3-vit-mixture-of-experts-vision-transformer","title":"M$^3$ViT: Mixture-of-Experts Vision Transformer for Efficient Multi-task Learning with Model-Accelerator Co-design","date":"2022-10-26","arxiv_id":"2210.14793","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":2,"n_instrument":1,"unverified":0,"pointer_only":1,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 2 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["vita-group/m3vit"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official","unlocated"]}}},{"paper":null,"slug":"multilevel-transformer-for-multimodal-emotion","title":"Multilevel Transformer For Multimodal Emotion Recognition","date":"2022-10-26","arxiv_id":"2211.07711","n_code_links":0,"syntology":null},{"paper":"/paper/pretrained-audio-neural-networks-for-speech","slug":"pretrained-audio-neural-networks-for-speech","title":"Pretrained audio neural networks for Speech emotion recognition in Portuguese","date":"2022-10-26","arxiv_id":"2210.14716","n_code_links":1,"syntology":null},{"paper":"/paper/semformer-semantic-guided-activation","slug":"semformer-semantic-guided-activation","title":"SemFormer: Semantic Guided Activation Transformer for Weakly Supervised Semantic Segmentation","date":"2022-10-26","arxiv_id":"2210.14618","n_code_links":1,"syntology":null},{"paper":"/paper/xiaoicesing-2-a-high-fidelity-singing-voice","slug":"xiaoicesing-2-a-high-fidelity-singing-voice","title":"Xiaoicesing 2: A High-Fidelity Singing Voice Synthesizer Based on Generative Adversarial Network","date":"2022-10-26","arxiv_id":"2210.14666","n_code_links":1,"syntology":null},{"paper":"/paper/audio-mfcc-gram-transformers-for-respiratory","slug":"audio-mfcc-gram-transformers-for-respiratory","title":"Audio MFCC-gram Transformers for respiratory insufficiency detection in COVID-19","date":"2022-10-25","arxiv_id":"2210.14085","n_code_links":1,"syntology":null},{"paper":null,"slug":"dynamic-survival-transformers-for-causal","title":"Dynamic Survival Transformers for Causal Inference with Electronic Health Records","date":"2022-10-25","arxiv_id":"2210.15417","n_code_links":0,"syntology":null},{"paper":null,"slug":"end-to-end-transformer-for-compressed-video","title":"End-to-end Transformer for Compressed Video Quality Enhancement","date":"2022-10-25","arxiv_id":"2210.13827","n_code_links":0,"syntology":null},{"paper":"/paper/explicitly-increasing-input-information","slug":"explicitly-increasing-input-information","title":"Explicitly Increasing Input Information Density for Vision Transformers on Small Datasets","date":"2022-10-25","arxiv_id":"2210.14319","n_code_links":1,"syntology":{"ran":4,"of":7,"n_ran_checked":2,"n_instrument":2,"unverified":3,"pointer_only":7,"phrase":"4 ran (of which 2 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 2 where Syntology's instrument failed) · 3 unverified","official":{"repos":["xiangyu8/densevt"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":2,"n_ran_no_instrument_failure":2,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":"/paper/how-long-is-enough-exploring-the-optimal","slug":"how-long-is-enough-exploring-the-optimal","title":"How Long Is Enough? Exploring the Optimal Intervals of Long-Range Clinical Note Language Modeling","date":"2022-10-25","arxiv_id":"2211.07713","n_code_links":1,"syntology":null},{"paper":null,"slug":"ielm-an-open-information-extraction-benchmark","title":"IELM: An Open Information Extraction Benchmark for Pre-Trained Language Models","date":"2022-10-25","arxiv_id":"2210.14128","n_code_links":0,"syntology":null},{"paper":null,"slug":"knowgl-knowledge-generation-and-linking-from","title":"KnowGL: Knowledge Generation and Linking from Text","date":"2022-10-25","arxiv_id":"2210.13952","n_code_links":0,"syntology":null},{"paper":"/paper/mew-unet-multi-axis-representation-learning","slug":"mew-unet-multi-axis-representation-learning","title":"MEW-UNet: Multi-axis representation learning in frequency domain for medical image segmentation","date":"2022-10-25","arxiv_id":"2210.14007","n_code_links":1,"syntology":null},{"paper":null,"slug":"minutiae-guided-fingerprint-embeddings-via","title":"Minutiae-Guided Fingerprint Embeddings via Vision Transformers","date":"2022-10-25","arxiv_id":"2210.13994","n_code_links":0,"syntology":null},{"paper":"/paper/moformer-self-supervised-transformer-model","slug":"moformer-self-supervised-transformer-model","title":"MOFormer: Self-Supervised Transformer model for Metal-Organic Framework Property Prediction","date":"2022-10-25","arxiv_id":"2210.14188","n_code_links":1,"syntology":null},{"paper":"/paper/thor-net-end-to-end-graformer-based-realistic","slug":"thor-net-end-to-end-graformer-based-realistic","title":"THOR-Net: End-to-end Graformer-based Realistic Two Hands and Object Reconstruction with Self-supervision","date":"2022-10-25","arxiv_id":"2210.13853","n_code_links":1,"syntology":null},{"paper":null,"slug":"xricl-cross-lingual-retrieval-augmented-in","title":"XRICL: Cross-lingual Retrieval-Augmented In-Context Learning for Cross-lingual Text-to-SQL Semantic Parsing","date":"2022-10-25","arxiv_id":"2210.13693","n_code_links":0,"syntology":null},{"paper":"/paper/abductive-action-inference","slug":"abductive-action-inference","title":"Inferring Past Human Actions in Homes with Abductive Reasoning","date":"2022-10-24","arxiv_id":"2210.13984","n_code_links":1,"syntology":null},{"paper":null,"slug":"effective-pre-training-objectives-for","title":"Effective Pre-Training Objectives for Transformer-based Autoencoders","date":"2022-10-24","arxiv_id":"2210.13536","n_code_links":0,"syntology":null},{"paper":"/paper/elmer-a-non-autoregressive-pre-trained","slug":"elmer-a-non-autoregressive-pre-trained","title":"ELMER: A Non-Autoregressive Pre-trained Language Model for Efficient and Effective Text Generation","date":"2022-10-24","arxiv_id":"2210.13304","n_code_links":1,"syntology":null},{"paper":"/paper/emergent-world-representations-exploring-a","slug":"emergent-world-representations-exploring-a","title":"Emergent World Representations: Exploring a Sequence Model Trained on a Synthetic Task","date":"2022-10-24","arxiv_id":"2210.13382","n_code_links":4,"syntology":{"ran":6,"of":6,"n_ran_checked":5,"n_instrument":1,"unverified":0,"pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["likenneth/othello_world"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"entity-level-sentiment-analysis-in-contact","title":"Entity-level Sentiment Analysis in Contact Center Telephone Conversations","date":"2022-10-24","arxiv_id":"2210.13401","n_code_links":0,"syntology":null},{"paper":null,"slug":"explaining-translationese-why-are-neural","title":"Explaining Translationese: why are Neural Classifiers Better and what do they Learn?","date":"2022-10-24","arxiv_id":"2210.13391","n_code_links":0,"syntology":null},{"paper":"/paper/exploring-euphemism-detection-in-few-shot-and","slug":"exploring-euphemism-detection-in-few-shot-and","title":"Exploring Euphemism Detection in Few-Shot and Zero-Shot Settings","date":"2022-10-24","arxiv_id":"2210.12926","n_code_links":1,"syntology":null},{"paper":"/paper/foreground-guidance-and-multi-layer-feature","slug":"foreground-guidance-and-multi-layer-feature","title":"Foreground Guidance and Multi-Layer Feature Fusion for Unsupervised Object Discovery with Transformers","date":"2022-10-24","arxiv_id":"2210.13053","n_code_links":1,"syntology":null},{"paper":"/paper/high-fidelity-neural-audio-compression","slug":"high-fidelity-neural-audio-compression","title":"High Fidelity Neural Audio Compression","date":"2022-10-24","arxiv_id":"2210.13438","n_code_links":6,"syntology":null},{"paper":"/paper/metaformer-baselines-for-vision","slug":"metaformer-baselines-for-vision","title":"MetaFormer Baselines for Vision","date":"2022-10-24","arxiv_id":"2210.13452","n_code_links":8,"syntology":{"ran":1,"of":4,"n_ran_checked":1,"n_instrument":0,"unverified":3,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":{"repos":["rwightman/pytorch-image-models","sail-sg/metaformer"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/perfectly-secure-steganography-using-minimum","slug":"perfectly-secure-steganography-using-minimum","title":"Perfectly Secure Steganography Using Minimum Entropy Coupling","date":"2022-10-24","arxiv_id":"2210.14889","n_code_links":2,"syntology":{"ran":1,"of":4,"n_ran_checked":1,"n_instrument":0,"unverified":3,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":{"repos":["schroederdewitt/perfectly-secure-steganography"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/sequential-recommendation-with-auxiliary-item","slug":"sequential-recommendation-with-auxiliary-item","title":"Sequential Recommendation with Auxiliary Item Relationships via Multi-Relational Transformer","date":"2022-10-24","arxiv_id":"2210.13572","n_code_links":1,"syntology":null},{"paper":null,"slug":"the-better-your-syntax-the-better-your","title":"The Better Your Syntax, the Better Your Semantics? Probing Pretrained Language Models for the English Comparative Correlative","date":"2022-10-24","arxiv_id":"2210.13181","n_code_links":0,"syntology":null},{"paper":"/paper/video-based-object-6d-pose-estimation-using","slug":"video-based-object-6d-pose-estimation-using","title":"Video based Object 6D Pose Estimation using Transformers","date":"2022-10-24","arxiv_id":"2210.13540","n_code_links":1,"syntology":null},{"paper":"/paper/vlc-bert-visual-question-answering-with","slug":"vlc-bert-visual-question-answering-with","title":"VLC-BERT: Visual Question Answering with Contextualized Commonsense Knowledge","date":"2022-10-24","arxiv_id":"2210.13626","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":2,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["aditya10/vlc-bert"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"a-bert-based-deep-learning-approach-for","title":"A BERT-based Deep Learning Approach for Reputation Analysis in Social Media","date":"2022-10-23","arxiv_id":"2211.01954","n_code_links":0,"syntology":null},{"paper":"/paper/anticipative-feature-fusion-transformer-for","slug":"anticipative-feature-fusion-transformer-for","title":"Anticipative Feature Fusion Transformer for Multi-Modal Action Anticipation","date":"2022-10-23","arxiv_id":"2210.12649","n_code_links":1,"syntology":null},{"paper":null,"slug":"automated-essay-scoring-using-transformers","title":"Data Augmentation for Automated Essay Scoring using Transformer Models","date":"2022-10-23","arxiv_id":"2210.12809","n_code_links":0,"syntology":null},{"paper":"/paper/delving-into-masked-autoencoders-for-multi","slug":"delving-into-masked-autoencoders-for-multi","title":"Delving into Masked Autoencoders for Multi-Label Thorax Disease Classification","date":"2022-10-23","arxiv_id":"2210.12843","n_code_links":1,"syntology":{"ran":1,"of":2,"n_ran_checked":1,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["lambert-x/medical_mae"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"discriminative-language-model-as-semantic","title":"Discriminative Language Model as Semantic Consistency Scorer for Prompt-based Few-Shot Text Classification","date":"2022-10-23","arxiv_id":"2210.12763","n_code_links":0,"syntology":null},{"paper":"/paper/holistic-interaction-transformer-network-for","slug":"holistic-interaction-transformer-network-for","title":"Holistic Interaction Transformer Network for Action Detection","date":"2022-10-23","arxiv_id":"2210.12686","n_code_links":1,"syntology":null},{"paper":"/paper/on-cross-domain-pre-trained-language-models","slug":"on-cross-domain-pre-trained-language-models","title":"Exploring the Value of Pre-trained Language Models for Clinical Named Entity Recognition","date":"2022-10-23","arxiv_id":"2210.12770","n_code_links":2,"syntology":null},{"paper":null,"slug":"uia-vit-unsupervised-inconsistency-aware","title":"UIA-ViT: Unsupervised Inconsistency-Aware Method based on Vision Transformer for Face Forgery Detection","date":"2022-10-23","arxiv_id":"2210.12752","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-comprehensive-comparison-of-neural-networks","title":"A Comprehensive Comparison of Neural Networks as Cognitive Models of Inflection","date":"2022-10-22","arxiv_id":"2210.12321","n_code_links":0,"syntology":null},{"paper":"/paper/ham-hierarchical-attention-model-with-high","slug":"ham-hierarchical-attention-model-with-high","title":"Learning Point-Language Hierarchical Alignment for 3D Visual Grounding","date":"2022-10-22","arxiv_id":"2210.12513","n_code_links":1,"syntology":{"ran":1,"of":2,"n_ran_checked":0,"n_instrument":1,"unverified":1,"pointer_only":2,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":{"repos":["ppjmchen/ham"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/leveraging-large-language-models-for-multiple","slug":"leveraging-large-language-models-for-multiple","title":"Leveraging Large Language Models for Multiple Choice Question Answering","date":"2022-10-22","arxiv_id":"2210.12353","n_code_links":1,"syntology":{"ran":6,"of":7,"n_ran_checked":3,"n_instrument":3,"unverified":1,"pointer_only":0,"phrase":"6 ran (of which 3 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 3 where Syntology's instrument failed) · 1 unverified","official":{"repos":["byu-pccl/leveraging-llms-for-mcqa"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":3,"n_ran_no_instrument_failure":3,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"meta-learning-pathologies-from-radiology","title":"Meta-learning Pathologies from Radiology Reports using Variance Aware Prototypical Networks","date":"2022-10-22","arxiv_id":"2210.13979","n_code_links":0,"syntology":null},{"paper":null,"slug":"ms-dc-unext-an-mlp-based-multi-scale-feature","title":"MS-DCANet: A Novel Segmentation Network For Multi-Modality COVID-19 Medical Images","date":"2022-10-22","arxiv_id":"2210.12361","n_code_links":0,"syntology":null},{"paper":null,"slug":"recurrence-boosts-diversity-revisiting","title":"Recurrence Boosts Diversity! Revisiting Recurrent Latent Variable in Transformer-Based Variational AutoEncoder for Diverse Text Generation","date":"2022-10-22","arxiv_id":"2210.12409","n_code_links":0,"syntology":null},{"paper":"/paper/s2wat-image-style-transfer-via-hierarchical","slug":"s2wat-image-style-transfer-via-hierarchical","title":"S2WAT: Image Style Transfer via Hierarchical Vision Transformer using Strips Window Attention","date":"2022-10-22","arxiv_id":"2210.12381","n_code_links":1,"syntology":null}],"record_sha256":"1883f2732104f440db90872d9e13c6d570ee57a3ed23a7b5c2a22f9c0c9a6f51","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}