{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/multi-head-attention/papers/69","list_of":"/method/multi-head-attention","method":"Multi-Head Attention","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":69,"pages_in_order":249,"rows_per_page":100,"rows":[6801,6900],"of":24855,"counts":{"archive_papers_tagged":24855,"with_a_code_link":11214,"where_syntology_ran_a_sample":3454,"not_listed_spam_title":0,"listed":24855,"listed_where_code_ran":3454,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":2916,"every_run_a_failure_of_syntologys_instrument":538,"listed_with_a_run_with_no_instrument_failure":2916,"listed_every_run_a_failure_of_syntologys_instrument":538,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/multi-head-attention","prev":"/method/multi-head-attention/papers/68","next":"/method/multi-head-attention/papers/70","papers":[{"paper":"/paper/llms-achieve-adult-human-performance-on","slug":"llms-achieve-adult-human-performance-on","title":"LLMs achieve adult human performance on higher-order theory of mind tasks","date":"2024-05-29","arxiv_id":"2405.18870","n_code_links":0,"syntology":null},{"paper":null,"slug":"lmo-dp-optimizing-the-randomization-mechanism","title":"LMO-DP: Optimizing the Randomization Mechanism for Differentially Private Fine-Tuning (Large) Language Models","date":"2024-05-29","arxiv_id":"2405.18776","n_code_links":0,"syntology":null},{"paper":"/paper/map-neo-highly-capable-and-transparent","slug":"map-neo-highly-capable-and-transparent","title":"MAP-Neo: Highly Capable and Transparent Bilingual Large Language Model Series","date":"2024-05-29","arxiv_id":"2405.19327","n_code_links":1,"syntology":{"ran":12,"of":14,"n_ran_checked":12,"n_instrument":0,"unverified":2,"pointer_only":14,"phrase":"12 ran (of which 0 constructed an object rather than computing a result; 12 with no instrument failure: 0 honoured, 0 violated, 12 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["multimodal-art-projection/map-neo"],"state":"official (archive's flag): 12 ran","n_ran":12,"n_constructed":0,"n_ran_no_instrument_failure":12,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/matryoshka-query-transformer-for-large-vision","slug":"matryoshka-query-transformer-for-large-vision","title":"Matryoshka Query Transformer for Large Vision-Language Models","date":"2024-05-29","arxiv_id":"2405.19315","n_code_links":1,"syntology":{"ran":9,"of":9,"n_ran_checked":3,"n_instrument":6,"unverified":0,"pointer_only":2,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 1 violated, 2 with no contract checked; 6 where Syntology's instrument failed) · 0 unverified","official":{"repos":["gordonhu608/mqt-llava"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official","unlocated"]}}},{"paper":"/paper/mds-vitnet-improving-saliency-prediction-for","slug":"mds-vitnet-improving-saliency-prediction-for","title":"MDS-ViTNet: Improving saliency prediction for Eye-Tracking with Vision Transformer","date":"2024-05-29","arxiv_id":"2405.19501","n_code_links":1,"syntology":null},{"paper":null,"slug":"mindsemantix-deciphering-brain-visual","title":"MindSemantix: Deciphering Brain Visual Experiences with a Brain-Language Model","date":"2024-05-29","arxiv_id":"2405.18812","n_code_links":0,"syntology":null},{"paper":null,"slug":"multi-channel-multi-step-spectrum-prediction","title":"Multi-Channel Multi-Step Spectrum Prediction Using Transformer and Stacked Bi-LSTM","date":"2024-05-29","arxiv_id":"2405.19138","n_code_links":0,"syntology":null},{"paper":null,"slug":"offline-regularised-reinforcement-learning","title":"Offline Regularised Reinforcement Learning for Large Language Models Alignment","date":"2024-05-29","arxiv_id":"2405.19107","n_code_links":0,"syntology":null},{"paper":"/paper/pediatricsgpt-large-language-models-as","slug":"pediatricsgpt-large-language-models-as","title":"PediatricsGPT: Large Language Models as Chinese Medical Assistants for Pediatric Applications","date":"2024-05-29","arxiv_id":"2405.19266","n_code_links":1,"syntology":{"ran":1,"of":2,"n_ran_checked":1,"n_instrument":0,"unverified":1,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["ydk122024/pediatricsgpt"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/reverse-image-retrieval-cues-parametric","slug":"reverse-image-retrieval-cues-parametric","title":"Reverse Image Retrieval Cues Parametric Memory in Multimodal LLMs","date":"2024-05-29","arxiv_id":"2405.18740","n_code_links":1,"syntology":null},{"paper":null,"slug":"stat-shrinking-transformers-after-training","title":"STAT: Shrinking Transformers After Training","date":"2024-05-29","arxiv_id":"2406.00061","n_code_links":0,"syntology":null},{"paper":"/paper/toward-conversational-agents-with-context-and","slug":"toward-conversational-agents-with-context-and","title":"Toward Conversational Agents with Context and Time Sensitive Long-term Memory","date":"2024-05-29","arxiv_id":"2406.00057","n_code_links":1,"syntology":null},{"paper":"/paper/transcending-fusion-a-multi-scale-alignment","slug":"transcending-fusion-a-multi-scale-alignment","title":"Transcending Fusion: A Multi-Scale Alignment Method for Remote Sensing Image-Text Retrieval","date":"2024-05-29","arxiv_id":"2405.18959","n_code_links":1,"syntology":null},{"paper":null,"slug":"two-layer-retrieval-augmented-generation","title":"Two-Layer Retrieval-Augmented Generation Framework for Low-Resource Medical Question Answering Using Reddit Data: Proof-of-Concept Study","date":"2024-05-29","arxiv_id":"2405.19519","n_code_links":0,"syntology":null},{"paper":"/paper/adapting-pre-trained-vision-models-for-novel","slug":"adapting-pre-trained-vision-models-for-novel","title":"Adapting Pre-Trained Vision Models for Novel Instance Detection and Segmentation","date":"2024-05-28","arxiv_id":"2405.17859","n_code_links":1,"syntology":{"ran":7,"of":7,"n_ran_checked":7,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["youngsean/nids-net"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":7,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/aligning-to-thousands-of-preferences-via","slug":"aligning-to-thousands-of-preferences-via","title":"Aligning to Thousands of Preferences via System Message Generalization","date":"2024-05-28","arxiv_id":"2405.17977","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":1,"n_instrument":2,"unverified":0,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 1 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":{"repos":["kaistAI/Janus"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["community","official"]}}},{"paper":"/paper/an-empirical-analysis-on-large-language","slug":"an-empirical-analysis-on-large-language","title":"An Empirical Analysis on Large Language Models in Debate Evaluation","date":"2024-05-28","arxiv_id":"2406.00050","n_code_links":1,"syntology":{"ran":1,"of":2,"n_ran_checked":1,"n_instrument":0,"unverified":1,"pointer_only":2,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["xinyiliu0227/llm_debate_bias"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"are-ppo-ed-language-models-hackable","title":"Are PPO-ed Language Models Hackable?","date":"2024-05-28","arxiv_id":"2406.02577","n_code_links":0,"syntology":null},{"paper":"/paper/atm-adversarial-tuning-multi-agent-system","slug":"atm-adversarial-tuning-multi-agent-system","title":"ATM: Adversarial Tuning Multi-agent System Makes a Robust Retrieval-Augmented Generator","date":"2024-05-28","arxiv_id":"2405.18111","n_code_links":1,"syntology":{"ran":6,"of":10,"n_ran_checked":6,"n_instrument":0,"unverified":4,"pointer_only":10,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","official":{"repos":["chuhac/atm-rag"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"attention-based-sequential-recommendation","title":"Attention-based sequential recommendation system using multimodal data","date":"2024-05-28","arxiv_id":"2405.17959","n_code_links":0,"syntology":null},{"paper":null,"slug":"benchmark-underestimates-the-readiness-of","title":"Benchmarks Underestimate the Readiness of Multi-lingual Dialogue Agents","date":"2024-05-28","arxiv_id":"2405.17840","n_code_links":0,"syntology":null},{"paper":null,"slug":"delving-into-differentially-private","title":"Delving into Differentially Private Transformer","date":"2024-05-28","arxiv_id":"2405.18194","n_code_links":0,"syntology":null},{"paper":null,"slug":"don-t-forget-to-connect-improving-rag-with","title":"Don't Forget to Connect! Improving RAG with Graph-based Reranking","date":"2024-05-28","arxiv_id":"2405.18414","n_code_links":0,"syntology":null},{"paper":null,"slug":"dual-path-multi-scale-transformer-for-high","title":"Dual-Path Multi-Scale Transformer for High-Quality Image Deraining","date":"2024-05-28","arxiv_id":"2405.18124","n_code_links":0,"syntology":null},{"paper":null,"slug":"edinburgh-clinical-nlp-at-mediqa-corr-2024","title":"Edinburgh Clinical NLP at MEDIQA-CORR 2024: Guiding Large Language Models with Hints","date":"2024-05-28","arxiv_id":"2405.18028","n_code_links":0,"syntology":null},{"paper":null,"slug":"efficient-time-series-processing-for","title":"Efficient Time Series Processing for Transformers and State-Space Models through Token Merging","date":"2024-05-28","arxiv_id":"2405.17951","n_code_links":0,"syntology":null},{"paper":"/paper/fastopic-a-fast-adaptive-stable-and","slug":"fastopic-a-fast-adaptive-stable-and","title":"FASTopic: Pretrained Transformer is a Fast, Adaptive, Stable, and Transferable Topic Model","date":"2024-05-28","arxiv_id":"2405.17978","n_code_links":2,"syntology":{"ran":3,"of":3,"n_ran_checked":3,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["bobxwu/fastopic","bobxwu/topmost"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/forecastgrapher-redefining-multivariate-time","slug":"forecastgrapher-redefining-multivariate-time","title":"ForecastGrapher: Redefining Multivariate Time Series Forecasting with Graph Neural Networks","date":"2024-05-28","arxiv_id":"2405.18036","n_code_links":0,"syntology":null},{"paper":"/paper/harmodt-harmony-multi-task-decision","slug":"harmodt-harmony-multi-task-decision","title":"HarmoDT: Harmony Multi-Task Decision Transformer for Offline Reinforcement Learning","date":"2024-05-28","arxiv_id":"2405.18080","n_code_links":1,"syntology":null},{"paper":null,"slug":"iapt-instruction-aware-prompt-tuning-for","title":"IAPT: Instruction-Aware Prompt Tuning for Large Language Models","date":"2024-05-28","arxiv_id":"2405.18203","n_code_links":0,"syntology":null},{"paper":"/paper/ldmol-text-conditioned-molecule-diffusion","slug":"ldmol-text-conditioned-molecule-diffusion","title":"LDMol: Text-to-Molecule Diffusion Model with Structurally Informative Latent Space","date":"2024-05-28","arxiv_id":"2405.17829","n_code_links":1,"syntology":{"ran":6,"of":9,"n_ran_checked":5,"n_instrument":1,"unverified":3,"pointer_only":1,"phrase":"6 ran (of which 2 constructed an object rather than computing a result; 5 with no instrument failure: 1 honoured, 0 violated, 4 with no contract checked; 1 where Syntology's instrument failed) · 3 unverified","official":{"repos":["jinhojsk515/ldmol"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":2,"n_ran_no_instrument_failure":5,"n_unverified":3,"ran_from_kinds":["official","unlocated"]}}},{"paper":"/paper/llms-and-memorization-on-quality-and","slug":"llms-and-memorization-on-quality-and","title":"LLMs and Memorization: On Quality and Specificity of Copyright Compliance","date":"2024-05-28","arxiv_id":"2405.18492","n_code_links":1,"syntology":null},{"paper":null,"slug":"lstm-cox-model-a-concise-and-efficient-deep","title":"Modeling Long Sequences in Bladder Cancer Recurrence: A Comparative Evaluation of LSTM,Transformer,and Mamba","date":"2024-05-28","arxiv_id":"2405.18518","n_code_links":0,"syntology":null},{"paper":null,"slug":"mindformer-a-transformer-architecture-for","title":"MindFormer: Semantic Alignment of Multi-Subject fMRI for Brain Decoding","date":"2024-05-28","arxiv_id":"2405.17720","n_code_links":0,"syntology":null},{"paper":null,"slug":"more-than-catastrophic-forgetting-integrating","title":"More Than Catastrophic Forgetting: Integrating General Capabilities For Domain-Specific LLMs","date":"2024-05-28","arxiv_id":"2405.17830","n_code_links":0,"syntology":null},{"paper":null,"slug":"multi-objective-representation-for-numbers-in","title":"Multi-objective Representation for Numbers in Clinical Narratives: A CamemBERT-Bio-Based Alternative to Large-Scale LLMs","date":"2024-05-28","arxiv_id":"2405.18448","n_code_links":0,"syntology":null},{"paper":null,"slug":"near-infrared-and-low-rank-adaptation-of","title":"Near-Infrared and Low-Rank Adaptation of Vision Transformers in Remote Sensing","date":"2024-05-28","arxiv_id":"2405.17901","n_code_links":0,"syntology":null},{"paper":null,"slug":"notes-on-applicability-of-gpt-4-to-document","title":"Notes on Applicability of GPT-4 to Document Understanding","date":"2024-05-28","arxiv_id":"2405.18433","n_code_links":0,"syntology":null},{"paper":"/paper/orlm-training-large-language-models-for","slug":"orlm-training-large-language-models-for","title":"ORLM: A Customizable Framework in Training Large Models for Automated Optimization Modeling","date":"2024-05-28","arxiv_id":"2405.17743","n_code_links":1,"syntology":{"ran":1,"of":3,"n_ran_checked":0,"n_instrument":1,"unverified":2,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","official":{"repos":["cardinal-operations/orlm"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/ov-dquo-open-vocabulary-detr-with-denoising","slug":"ov-dquo-open-vocabulary-detr-with-denoising","title":"OV-DQUO: Open-Vocabulary DETR with Denoising Text Query Training and Open-World Unknown Objects Supervision","date":"2024-05-28","arxiv_id":"2405.17913","n_code_links":1,"syntology":{"ran":9,"of":11,"n_ran_checked":7,"n_instrument":2,"unverified":2,"pointer_only":2,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 1 honoured, 0 violated, 6 with no contract checked; 2 where Syntology's instrument failed) · 2 unverified","official":{"repos":["xiaomoguhz/ov-dquo"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":7,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/peering-into-the-mind-of-language-models-an","slug":"peering-into-the-mind-of-language-models-an","title":"Peering into the Mind of Language Models: An Approach for Attribution in Contextual Question Answering","date":"2024-05-28","arxiv_id":"2405.17980","n_code_links":1,"syntology":null},{"paper":null,"slug":"proof-of-quality-a-costless-paradigm-for","title":"Proof of Quality: A Costless Paradigm for Trustless Generative AI Model Inference on Blockchains","date":"2024-05-28","arxiv_id":"2405.17934","n_code_links":0,"syntology":null},{"paper":null,"slug":"realitysummary-on-demand-mixed-reality","title":"RealitySummary: Exploring On-Demand Mixed Reality Text Summarization and Question Answering using Large Language Models","date":"2024-05-28","arxiv_id":"2405.18620","n_code_links":0,"syntology":null},{"paper":"/paper/thai-winograd-schemas-a-benchmark-for-thai","slug":"thai-winograd-schemas-a-benchmark-for-thai","title":"Thai Winograd Schemas: A Benchmark for Thai Commonsense Reasoning","date":"2024-05-28","arxiv_id":"2405.18375","n_code_links":1,"syntology":null},{"paper":null,"slug":"the-battle-of-llms-a-comparative-study-in","title":"The Battle of LLMs: A Comparative Study in Conversational QA Tasks","date":"2024-05-28","arxiv_id":"2405.18344","n_code_links":0,"syntology":null},{"paper":null,"slug":"understanding-intrinsic-socioeconomic-biases","title":"Understanding Intrinsic Socioeconomic Biases in Large Language Models","date":"2024-05-28","arxiv_id":"2405.18662","n_code_links":0,"syntology":null},{"paper":"/paper/vig-linear-complexity-visual-sequence","slug":"vig-linear-complexity-visual-sequence","title":"ViG: Linear-complexity Visual Sequence Learning with Gated Linear Attention","date":"2024-05-28","arxiv_id":"2405.18425","n_code_links":1,"syntology":null},{"paper":"/paper/visual-anchors-are-strong-information","slug":"visual-anchors-are-strong-information","title":"Visual Anchors Are Strong Information Aggregators For Multimodal Large Language Model","date":"2024-05-28","arxiv_id":"2405.17815","n_code_links":1,"syntology":null},{"paper":null,"slug":"visualizing-the-loss-landscape-of-self","title":"Visualizing the loss landscape of Self-supervised Vision Transformer","date":"2024-05-28","arxiv_id":"2405.18042","n_code_links":0,"syntology":null},{"paper":null,"slug":"viton-dit-learning-in-the-wild-video-try-on","title":"VITON-DiT: Learning In-the-Wild Video Try-On from Human Dance Videos via Diffusion Transformers","date":"2024-05-28","arxiv_id":"2405.18326","n_code_links":0,"syntology":null},{"paper":null,"slug":"wavelet-based-image-tokenizer-for-vision","title":"Wavelet-Based Image Tokenizer for Vision Transformers","date":"2024-05-28","arxiv_id":"2405.18616","n_code_links":0,"syntology":null},{"paper":null,"slug":"widin-wording-image-for-domain-invariant","title":"WIDIn: Wording Image for Domain-Invariant Representation in Single-Source Domain Generalization","date":"2024-05-28","arxiv_id":"2405.18405","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-one-layer-decoder-only-transformer-is-a-two","title":"A One-Layer Decoder-Only Transformer is a Two-Layer RNN: With an Application to Certified Robustness","date":"2024-05-27","arxiv_id":"2405.17361","n_code_links":0,"syntology":null},{"paper":"/paper/advanced-language-model-based-translator-for","slug":"advanced-language-model-based-translator-for","title":"Advanced Language Model-based Translator for English-Vietnamese Translation","date":"2024-05-27","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/are-self-attentions-effective-for-time-series","slug":"are-self-attentions-effective-for-time-series","title":"Are Self-Attentions Effective for Time Series Forecasting?","date":"2024-05-27","arxiv_id":"2405.16877","n_code_links":1,"syntology":{"ran":5,"of":13,"n_ran_checked":5,"n_instrument":0,"unverified":8,"pointer_only":3,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 8 unverified","official":{"repos":["dongbeank/cats"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":8,"ran_from_kinds":["official"]}}},{"paper":"/paper/assessing-llms-suitability-for-knowledge","slug":"assessing-llms-suitability-for-knowledge","title":"Assessing LLMs Suitability for Knowledge Graph Completion","date":"2024-05-27","arxiv_id":"2405.17249","n_code_links":1,"syntology":null},{"paper":null,"slug":"augmenting-textual-generation-via-topology","title":"Augmenting Textual Generation via Topology Aware Retrieval","date":"2024-05-27","arxiv_id":"2405.17602","n_code_links":0,"syntology":null},{"paper":"/paper/autoformalizing-euclidean-geometry","slug":"autoformalizing-euclidean-geometry","title":"Autoformalizing Euclidean Geometry","date":"2024-05-27","arxiv_id":"2405.17216","n_code_links":1,"syntology":{"ran":9,"of":13,"n_ran_checked":4,"n_instrument":5,"unverified":4,"pointer_only":0,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 1 honoured, 2 violated, 1 with no contract checked; 5 where Syntology's instrument failed) · 4 unverified","official":{"repos":["loganrjmurphy/leaneuclid"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"automatic-domain-adaptation-by-transformers","title":"Automatic Domain Adaptation by Transformers in In-Context Learning","date":"2024-05-27","arxiv_id":"2405.16819","n_code_links":0,"syntology":null},{"paper":null,"slug":"behaviorgpt-smart-agent-simulation-for","title":"BehaviorGPT: Smart Agent Simulation for Autonomous Driving with Next-Patch Prediction","date":"2024-05-27","arxiv_id":"2405.17372","n_code_links":0,"syntology":null},{"paper":"/paper/chess-contextual-harnessing-for-efficient-sql","slug":"chess-contextual-harnessing-for-efficient-sql","title":"CHESS: Contextual Harnessing for Efficient SQL Synthesis","date":"2024-05-27","arxiv_id":"2405.16755","n_code_links":2,"syntology":{"ran":6,"of":6,"n_ran_checked":6,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["shayantalaei/chess"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"cost-efficient-knowledge-based-question","title":"Cost-efficient Knowledge-based Question Answering with Large Language Models","date":"2024-05-27","arxiv_id":"2405.17337","n_code_links":0,"syntology":null},{"paper":"/paper/deciphering-movement-unified-trajectory","slug":"deciphering-movement-unified-trajectory","title":"Deciphering Movement: Unified Trajectory Generation Model for Multi-Agent","date":"2024-05-27","arxiv_id":"2405.17680","n_code_links":1,"syntology":{"ran":8,"of":10,"n_ran_checked":8,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 0 violated, 8 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["colorfulfuture/unitraj-pytorch"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/deeperimpact-optimizing-sparse-learned-index","slug":"deeperimpact-optimizing-sparse-learned-index","title":"DeeperImpact: Optimizing Sparse Learned Index Structures","date":"2024-05-27","arxiv_id":"2405.17093","n_code_links":1,"syntology":null},{"paper":null,"slug":"detecting-deceptive-dark-patterns-in-e","title":"Detecting Deceptive Dark Patterns in E-commerce Platforms","date":"2024-05-27","arxiv_id":"2406.01608","n_code_links":0,"syntology":null},{"paper":null,"slug":"exploiting-the-layered-intrinsic","title":"Exploiting the Layered Intrinsic Dimensionality of Deep Models for Practical Adversarial Training","date":"2024-05-27","arxiv_id":"2405.17130","n_code_links":0,"syntology":null},{"paper":null,"slug":"galaxy-a-resource-efficient-collaborative","title":"Galaxy: A Resource-Efficient Collaborative Edge AI System for In-situ Transformer Inference","date":"2024-05-27","arxiv_id":"2405.17245","n_code_links":0,"syntology":null},{"paper":"/paper/generation-and-human-expert-evaluation-of","slug":"generation-and-human-expert-evaluation-of","title":"Interesting Scientific Idea Generation using Knowledge Graphs and LLMs: Evaluations with 100 Research Group Leaders","date":"2024-05-27","arxiv_id":"2405.17044","n_code_links":1,"syntology":null},{"paper":null,"slug":"how-does-perfect-fitting-affect","title":"How Do the Architecture and Optimizer Affect Representation Learning? On the Training Dynamics of Representations in Deep Neural Networks","date":"2024-05-27","arxiv_id":"2405.17377","n_code_links":0,"syntology":null},{"paper":"/paper/inversionview-a-general-purpose-method-for","slug":"inversionview-a-general-purpose-method-for","title":"InversionView: A General-Purpose Method for Reading Information from Neural Activations","date":"2024-05-27","arxiv_id":"2405.17653","n_code_links":1,"syntology":{"ran":1,"of":6,"n_ran_checked":1,"n_instrument":0,"unverified":5,"pointer_only":6,"phrase":"1 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 5 unverified; the one sample that ran constructed an object rather than computing a result","official":{"repos":["huangxt39/inversionview"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":5,"ran_from_kinds":["official"]}}},{"paper":"/paper/lcm-locally-constrained-compact-point-cloud","slug":"lcm-locally-constrained-compact-point-cloud","title":"LCM: Locally Constrained Compact Point Cloud Model for Masked Point Modeling","date":"2024-05-27","arxiv_id":"2405.17149","n_code_links":1,"syntology":{"ran":12,"of":12,"n_ran_checked":5,"n_instrument":7,"unverified":0,"pointer_only":10,"phrase":"12 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 3 honoured, 0 violated, 2 with no contract checked; 7 where Syntology's instrument failed) · 0 unverified","official":{"repos":["zyh16143998882/lcm"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":0,"ran_from_kinds":["community","official"]}}},{"paper":null,"slug":"llm-based-cooperative-agents-using","title":"REVECA: Adaptive Planning and Trajectory-based Validation in Cooperative Language Agents using Information Relevance and Relative Proximity","date":"2024-05-27","arxiv_id":"2405.16751","n_code_links":0,"syntology":null},{"paper":"/paper/loretrack-efficient-and-accurate-low","slug":"loretrack-efficient-and-accurate-low","title":"LoReTrack: Efficient and Accurate Low-Resolution Transformer Tracking","date":"2024-05-27","arxiv_id":"2405.17660","n_code_links":1,"syntology":null},{"paper":null,"slug":"masked-face-recognition-with-generative-to","title":"Masked Face Recognition with Generative-to-Discriminative Representations","date":"2024-05-27","arxiv_id":"2405.16761","n_code_links":0,"syntology":null},{"paper":"/paper/motionllm-multimodal-motion-language-learning","slug":"motionllm-multimodal-motion-language-learning","title":"Motion-Agent: A Conversational Framework for Human Motion Generation with LLMs","date":"2024-05-27","arxiv_id":"2405.17013","n_code_links":1,"syntology":{"ran":6,"of":6,"n_ran_checked":6,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["szqwu/Motion-Agent"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"novel-approaches-for-ml-assisted-particle","title":"Novel Approaches for ML-Assisted Particle Track Reconstruction and Hit Clustering","date":"2024-05-27","arxiv_id":"2405.17325","n_code_links":0,"syntology":null},{"paper":null,"slug":"nv-embed-improved-techniques-for-training","title":"NV-Embed: Improved Techniques for Training LLMs as Generalist Embedding Models","date":"2024-05-27","arxiv_id":"2405.17428","n_code_links":0,"syntology":null},{"paper":null,"slug":"pae-llm-based-product-attribute-extraction","title":"PAE: LLM-based Product Attribute Extraction for E-Commerce Fashion Trends","date":"2024-05-27","arxiv_id":"2405.17533","n_code_links":0,"syntology":null},{"paper":null,"slug":"performance-evaluation-of-reddit-comments","title":"Performance evaluation of Reddit Comments using Machine Learning and Natural Language Processing methods in Sentiment Analysis","date":"2024-05-27","arxiv_id":"2405.16810","n_code_links":0,"syntology":null},{"paper":null,"slug":"pivotmesh-generic-3d-mesh-generation-via","title":"PivotMesh: Generic 3D Mesh Generation via Pivot Vertices Guidance","date":"2024-05-27","arxiv_id":"2405.16890","n_code_links":0,"syntology":null},{"paper":"/paper/q-value-regularized-transformer-for-offline","slug":"q-value-regularized-transformer-for-offline","title":"Q-value Regularized Transformer for Offline Reinforcement Learning","date":"2024-05-27","arxiv_id":"2405.17098","n_code_links":2,"syntology":{"ran":6,"of":8,"n_ran_checked":5,"n_instrument":1,"unverified":2,"pointer_only":3,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","official":null}},{"paper":null,"slug":"qub-cirdan-at-discharge-me-zero-shot","title":"QUB-Cirdan at \"Discharge Me!\": Zero shot discharge letter generation by open-source LLM","date":"2024-05-27","arxiv_id":"2406.00041","n_code_links":0,"syntology":null},{"paper":"/paper/reflectioncoder-learning-from-reflection","slug":"reflectioncoder-learning-from-reflection","title":"ReflectionCoder: Learning from Reflection Sequence for Enhanced One-off Code Generation","date":"2024-05-27","arxiv_id":"2405.17057","n_code_links":1,"syntology":{"ran":0,"of":1,"n_ran_checked":0,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"0 ran · 1 unverified","official":{"repos":["sensellm/reflectioncoder"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":[]}}},{"paper":"/paper/rethinking-transformers-in-solving-pomdps","slug":"rethinking-transformers-in-solving-pomdps","title":"Rethinking Transformers in Solving POMDPs","date":"2024-05-27","arxiv_id":"2405.17358","n_code_links":1,"syntology":{"ran":0,"of":1,"n_ran_checked":0,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"0 ran · 1 unverified","official":{"repos":["ctp314/tfporl"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":[]}}},{"paper":"/paper/rtl-repo-a-benchmark-for-evaluating-llms-on","slug":"rtl-repo-a-benchmark-for-evaluating-llms-on","title":"RTL-Repo: A Benchmark for Evaluating LLMs on Large-Scale RTL Design Projects","date":"2024-05-27","arxiv_id":"2405.17378","n_code_links":1,"syntology":null},{"paper":null,"slug":"sa-gs-semantic-aware-gaussian-splatting-for","title":"SA-GS: Semantic-Aware Gaussian Splatting for Large Scene Reconstruction with Geometry Constrain","date":"2024-05-27","arxiv_id":"2405.16923","n_code_links":0,"syntology":null},{"paper":"/paper/safe-lora-the-silver-lining-of-reducing","slug":"safe-lora-the-silver-lining-of-reducing","title":"Safe LoRA: the Silver Lining of Reducing Safety Risks when Fine-tuning Large Language Models","date":"2024-05-27","arxiv_id":"2405.16833","n_code_links":1,"syntology":{"ran":0,"of":1,"n_ran_checked":0,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"0 ran · 1 unverified","official":{"repos":["ibm/safelora"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":[]}}},{"paper":null,"slug":"supervised-batch-normalization","title":"Supervised Batch Normalization","date":"2024-05-27","arxiv_id":"2405.17027","n_code_links":0,"syntology":null},{"paper":null,"slug":"swat-scalable-and-efficient-window-attention","title":"SWAT: Scalable and Efficient Window Attention-based Transformers Acceleration on FPGAs","date":"2024-05-27","arxiv_id":"2405.17025","n_code_links":0,"syntology":null},{"paper":null,"slug":"the-scaling-law-in-stellar-light-curves","title":"The Scaling Law in Stellar Light Curves","date":"2024-05-27","arxiv_id":"2405.17156","n_code_links":0,"syntology":null},{"paper":"/paper/thread-thinking-deeper-with-recursive","slug":"thread-thinking-deeper-with-recursive","title":"THREAD: Thinking Deeper with Recursive Spawning","date":"2024-05-27","arxiv_id":"2405.17402","n_code_links":1,"syntology":{"ran":7,"of":11,"n_ran_checked":6,"n_instrument":1,"unverified":4,"pointer_only":11,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 1 where Syntology's instrument failed) · 4 unverified","official":{"repos":["philipmit/thread"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"transformer-in-context-learning-for","title":"On Understanding Attention-Based In-Context Learning for Categorical Data","date":"2024-05-27","arxiv_id":"2405.17248","n_code_links":0,"syntology":null},{"paper":null,"slug":"uit-darkcow-team-at-imageclefmedical-caption","title":"UIT-DarkCow team at ImageCLEFmedical Caption 2024: Diagnostic Captioning for Radiology Images Efficiency with Transformer Models","date":"2024-05-27","arxiv_id":"2405.17002","n_code_links":0,"syntology":null},{"paper":"/paper/unisolver-pde-conditional-transformers-are","slug":"unisolver-pde-conditional-transformers-are","title":"Unisolver: PDE-Conditional Transformers Are Universal PDE Solvers","date":"2024-05-27","arxiv_id":"2405.17527","n_code_links":1,"syntology":{"ran":5,"of":6,"n_ran_checked":5,"n_instrument":0,"unverified":1,"pointer_only":1,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 1 honoured, 1 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":null}},{"paper":"/paper/video-enriched-retrieval-augmented-generation","slug":"video-enriched-retrieval-augmented-generation","title":"Video Enriched Retrieval Augmented Generation Using Aligned Video Captions","date":"2024-05-27","arxiv_id":"2405.17706","n_code_links":1,"syntology":null},{"paper":null,"slug":"vision-and-language-navigation-generative","title":"Vision-and-Language Navigation Generative Pretrained Transformer","date":"2024-05-27","arxiv_id":"2405.16994","n_code_links":0,"syntology":null},{"paper":"/paper/a-unified-implicit-attention-formulation-for","slug":"a-unified-implicit-attention-formulation-for","title":"Explaining Modern Gated-Linear RNNs via a Unified Implicit Attention Formulation","date":"2024-05-26","arxiv_id":"2405.16504","n_code_links":1,"syntology":null},{"paper":null,"slug":"ai-generated-text-detection-and","title":"AI-Generated Text Detection and Classification Based on BERT Deep Learning Algorithm","date":"2024-05-26","arxiv_id":"2405.16422","n_code_links":0,"syntology":null},{"paper":"/paper/darijabanking-a-new-resource-for-overcoming","slug":"darijabanking-a-new-resource-for-overcoming","title":"DarijaBanking: A New Resource for Overcoming Language Barriers in Banking Intent Detection for Moroccan Arabic Speakers","date":"2024-05-26","arxiv_id":"2405.16482","n_code_links":1,"syntology":null},{"paper":"/paper/demystify-mamba-in-vision-a-linear-attention","slug":"demystify-mamba-in-vision-a-linear-attention","title":"Demystify Mamba in Vision: A Linear Attention Perspective","date":"2024-05-26","arxiv_id":"2405.16605","n_code_links":1,"syntology":{"ran":6,"of":8,"n_ran_checked":6,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"6 ran (of which 6 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified; every one of the 6 samples that ran constructed an object rather than computing a result","official":{"repos":["LeapLabTHU/MLLA"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":6,"n_ran_no_instrument_failure":6,"n_unverified":2,"ran_from_kinds":["official"]}}}],"record_sha256":"4b054635484718854bbb2ae6ac29a2108d73e7d313584a9fbbc0c70ccb9ad461","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}