{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/attention-dropout/papers/29","list_of":"/method/attention-dropout","method":"Attention Dropout","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":29,"pages_in_order":109,"rows_per_page":100,"rows":[2801,2900],"of":10892,"counts":{"archive_papers_tagged":10892,"with_a_code_link":4634,"where_syntology_ran_a_sample":1270,"not_listed_spam_title":0,"listed":10892,"listed_where_code_ran":1270,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":1043,"every_run_a_failure_of_syntologys_instrument":227,"listed_with_a_run_with_no_instrument_failure":1043,"listed_every_run_a_failure_of_syntologys_instrument":227,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/attention-dropout","prev":"/method/attention-dropout/papers/28","next":"/method/attention-dropout/papers/30","papers":[{"paper":null,"slug":"focus-on-the-core-efficient-attention-via","title":"Focus on the Core: Efficient Attention via Pruned Token Compression for Document Classification","date":"2024-06-03","arxiv_id":"2406.01283","n_code_links":0,"syntology":null},{"paper":null,"slug":"in-context-learning-of-physical-properties","title":"In-Context Learning of Physical Properties: Few-Shot Adaptation to Out-of-Distribution Molecular Graphs","date":"2024-06-03","arxiv_id":"2406.01808","n_code_links":0,"syntology":null},{"paper":null,"slug":"luna-an-evaluation-foundation-model-to-catch","title":"Luna: An Evaluation Foundation Model to Catch Language Model Hallucinations with High Accuracy and Low Cost","date":"2024-06-03","arxiv_id":"2406.00975","n_code_links":0,"syntology":null},{"paper":null,"slug":"rag-enabled-conversations-about-household","title":"Natural Language Interaction with a Household Electricity Knowledge-based Digital Twin","date":"2024-06-03","arxiv_id":"2406.06566","n_code_links":0,"syntology":null},{"paper":"/paper/semcoder-training-code-language-models-with","slug":"semcoder-training-code-language-models-with","title":"SemCoder: Training Code Language Models with Comprehensive Semantics Reasoning","date":"2024-06-03","arxiv_id":"2406.01006","n_code_links":1,"syntology":{"ran":12,"of":15,"n_ran_checked":10,"n_instrument":2,"unverified":3,"pointer_only":0,"phrase":"12 ran (of which 0 constructed an object rather than computing a result; 10 with no instrument failure: 1 honoured, 0 violated, 9 with no contract checked; 2 where Syntology's instrument failed) · 3 unverified","official":{"repos":["arise-lab/semcoder"],"state":"official (archive's flag): 12 ran","n_ran":12,"n_constructed":0,"n_ran_no_instrument_failure":10,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":"/paper/soccerrag-multimodal-soccer-information","slug":"soccerrag-multimodal-soccer-information","title":"SoccerRAG: Multimodal Soccer Information Retrieval via Natural Queries","date":"2024-06-03","arxiv_id":"2406.01273","n_code_links":1,"syntology":null},{"paper":"/paper/spatialrgpt-grounded-spatial-reasoning-in","slug":"spatialrgpt-grounded-spatial-reasoning-in","title":"SpatialRGPT: Grounded Spatial Reasoning in Vision Language Models","date":"2024-06-03","arxiv_id":"2406.01584","n_code_links":1,"syntology":null},{"paper":null,"slug":"superhuman-performance-in-urology-board","title":"Superhuman performance in urology board questions by an explainable large language model enabled for context integration of the European Association of Urology guidelines: the UroBot study","date":"2024-06-03","arxiv_id":"2406.01428","n_code_links":0,"syntology":null},{"paper":null,"slug":"unsupervised-distractor-generation-via-large","title":"Unsupervised Distractor Generation via Large Language Model Distilling and Counterfactual Contrastive Decoding","date":"2024-06-03","arxiv_id":"2406.01306","n_code_links":0,"syntology":null},{"paper":null,"slug":"unveil-the-duality-of-retrieval-augmented","title":"A Theory for Token-Level Harmonization in Retrieval-Augmented Generation","date":"2024-06-03","arxiv_id":"2406.00944","n_code_links":0,"syntology":null},{"paper":null,"slug":"applying-fine-tuned-llms-for-reducing-data","title":"Applying Fine-Tuned LLMs for Reducing Data Needs in Load Profile Analysis","date":"2024-06-02","arxiv_id":"2406.02479","n_code_links":0,"syntology":null},{"paper":"/paper/evaluating-mathematical-reasoning-of-large","slug":"evaluating-mathematical-reasoning-of-large","title":"Evaluating Mathematical Reasoning of Large Language Models: A Focus on Error Identification and Correction","date":"2024-06-02","arxiv_id":"2406.00755","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":3,"n_instrument":0,"unverified":0,"pointer_only":3,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["littlecirc1e/eic"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"focus-forging-originality-through-contrastive","title":"FOCUS: Forging Originality through Contrastive Use in Self-Plagiarism for Language Models","date":"2024-06-02","arxiv_id":"2406.00839","n_code_links":0,"syntology":null},{"paper":null,"slug":"formality-style-transfer-in-persian","title":"Formality Style Transfer in Persian","date":"2024-06-02","arxiv_id":"2406.00867","n_code_links":0,"syntology":null},{"paper":null,"slug":"pretrained-hybrids-with-mad-skills","title":"Pretrained Hybrids with MAD Skills","date":"2024-06-02","arxiv_id":"2406.00894","n_code_links":0,"syntology":null},{"paper":null,"slug":"2406-07572","title":"Domain-specific ReAct for physics-integrated iterative modeling: A case study of LLM agents for gas path analysis of gas turbines","date":"2024-06-01","arxiv_id":"2406.07572","n_code_links":0,"syntology":null},{"paper":null,"slug":"an-evaluation-benchmark-for-autoformalization","title":"An Evaluation Benchmark for Autoformalization in Lean4","date":"2024-06-01","arxiv_id":"2406.06555","n_code_links":0,"syntology":null},{"paper":null,"slug":"beyond-metrics-evaluating-llms-effectiveness","title":"Beyond Metrics: Evaluating LLMs' Effectiveness in Culturally Nuanced, Low-Resource Real-World Scenarios","date":"2024-06-01","arxiv_id":"2406.00343","n_code_links":0,"syntology":null},{"paper":"/paper/case-curricular-data-pre-training-for","slug":"case-curricular-data-pre-training-for","title":"CASE: Efficient Curricular Data Pre-training for Building Assistive Psychology Expert Models","date":"2024-06-01","arxiv_id":"2406.00314","n_code_links":1,"syntology":null},{"paper":"/paper/mix-of-granularity-optimize-the-chunking","slug":"mix-of-granularity-optimize-the-chunking","title":"Mix-of-Granularity: Optimize the Chunking Granularity for Retrieval-Augmented Generation","date":"2024-06-01","arxiv_id":"2406.00456","n_code_links":1,"syntology":null},{"paper":null,"slug":"pseudo-label-based-domain-adaptation-for-zero","title":"Pseudo-label Based Domain Adaptation for Zero-Shot Text Steganalysis","date":"2024-06-01","arxiv_id":"2406.18565","n_code_links":0,"syntology":null},{"paper":"/paper/roberta-bilstm-a-context-aware-hybrid-model","slug":"roberta-bilstm-a-context-aware-hybrid-model","title":"RoBERTa-BiLSTM: A Context-Aware Hybrid Model for Sentiment Analysis","date":"2024-06-01","arxiv_id":"2406.00367","n_code_links":1,"syntology":null},{"paper":"/paper/a-comparison-of-correspondence-analysis-with","slug":"a-comparison-of-correspondence-analysis-with","title":"A comparison of correspondence analysis with PMI-based word embedding methods","date":"2024-05-31","arxiv_id":"2405.20895","n_code_links":1,"syntology":null},{"paper":null,"slug":"bi-directional-transformers-vs-word2vec","title":"Bi-Directional Transformers vs. word2vec: Discovering Vulnerabilities in Lifted Compiled Code","date":"2024-05-31","arxiv_id":"2405.20611","n_code_links":0,"syntology":null},{"paper":"/paper/enhancing-noise-robustness-of-retrieval","slug":"enhancing-noise-robustness-of-retrieval","title":"Enhancing Noise Robustness of Retrieval-Augmented Language Models with Adaptive Adversarial Training","date":"2024-05-31","arxiv_id":"2405.20978","n_code_links":1,"syntology":{"ran":9,"of":9,"n_ran_checked":9,"n_instrument":0,"unverified":0,"pointer_only":9,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 9 with no instrument failure: 0 honoured, 0 violated, 9 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["calubkk/raat"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":9,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/generative-ai-voting-fair-collective-choice","slug":"generative-ai-voting-fair-collective-choice","title":"Generative AI Voting: Fair Collective Choice is Resilient to LLM Biases and Inconsistencies","date":"2024-05-31","arxiv_id":"2406.11871","n_code_links":1,"syntology":null},{"paper":"/paper/hard-cases-detection-in-motion-prediction-by","slug":"hard-cases-detection-in-motion-prediction-by","title":"Hard Cases Detection in Motion Prediction by Vision-Language Foundation Models","date":"2024-05-31","arxiv_id":"2405.20991","n_code_links":1,"syntology":null},{"paper":"/paper/large-language-models-are-zero-shot-next","slug":"large-language-models-are-zero-shot-next","title":"Large Language Models are Zero-Shot Next Location Predictors","date":"2024-05-31","arxiv_id":"2405.20962","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":3,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["ssai-trento/llm-zero-shot-nl"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"lolameme-logic-language-memory-mechanistic","title":"LOLAMEME: Logic, Language, Memory, Mechanistic Framework","date":"2024-05-31","arxiv_id":"2406.02592","n_code_links":0,"syntology":null},{"paper":"/paper/multilingual-text-style-transfer-datasets","slug":"multilingual-text-style-transfer-datasets","title":"Multilingual Text Style Transfer: Datasets & Models for Indian Languages","date":"2024-05-31","arxiv_id":"2405.20805","n_code_links":2,"syntology":null},{"paper":null,"slug":"rag-does-not-work-for-enterprises","title":"RAG Does Not Work for Enterprises","date":"2024-05-31","arxiv_id":"2406.04369","n_code_links":0,"syntology":null},{"paper":null,"slug":"retrieval-meets-reasoning-even-high-school","title":"Retrieval Meets Reasoning: Even High-school Textbook Knowledge Benefits Multimodal Reasoning","date":"2024-05-31","arxiv_id":"2405.20834","n_code_links":0,"syntology":null},{"paper":null,"slug":"the-point-of-view-of-a-sentiment-towards","title":"The Point of View of a Sentiment: Towards Clinician Bias Detection in Psychiatric Notes","date":"2024-05-31","arxiv_id":"2405.20582","n_code_links":0,"syntology":null},{"paper":"/paper/anah-analytical-annotation-of-hallucinations","slug":"anah-analytical-annotation-of-hallucinations","title":"ANAH: Analytical Annotation of Hallucinations in Large Language Models","date":"2024-05-30","arxiv_id":"2405.20315","n_code_links":1,"syntology":{"ran":9,"of":9,"n_ran_checked":9,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 9 with no instrument failure: 0 honoured, 0 violated, 9 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["open-compass/anah"],"state":"official: no sample here; runs from other or unrecorded repositories","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["found_in_text"]}}},{"paper":null,"slug":"autobreach-universal-and-adaptive","title":"AutoBreach: Universal and Adaptive Jailbreaking with Efficient Wordplay-Guided Optimization","date":"2024-05-30","arxiv_id":"2405.19668","n_code_links":0,"syntology":null},{"paper":null,"slug":"divide-and-conquer-meets-consensus-unleashing","title":"Divide-and-Conquer Meets Consensus: Unleashing the Power of Functions in Code Generation","date":"2024-05-30","arxiv_id":"2405.20092","n_code_links":0,"syntology":null},{"paper":null,"slug":"ensemble-model-with-bert-roberta-and-xlnet","title":"Ensemble Model With Bert,Roberta and Xlnet For Molecular property prediction","date":"2024-05-30","arxiv_id":"2406.06553","n_code_links":0,"syntology":null},{"paper":"/paper/gnn-rag-graph-neural-retrieval-for-large","slug":"gnn-rag-graph-neural-retrieval-for-large","title":"GNN-RAG: Graph Neural Retrieval for Large Language Model Reasoning","date":"2024-05-30","arxiv_id":"2405.20139","n_code_links":1,"syntology":{"ran":6,"of":7,"n_ran_checked":6,"n_instrument":0,"unverified":1,"pointer_only":7,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["cmavro/gnn-rag"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/heidelberg-boston-sigtyp-2024-shared-task","slug":"heidelberg-boston-sigtyp-2024-shared-task","title":"Heidelberg-Boston @ SIGTYP 2024 Shared Task: Enhancing Low-Resource Language Analysis With Character-Aware Hierarchical Transformers","date":"2024-05-30","arxiv_id":"2405.20145","n_code_links":1,"syntology":null},{"paper":null,"slug":"is-my-data-in-your-retrieval-database","title":"Is My Data in Your Retrieval Database? Membership Inference Attacks Against Retrieval Augmented Generation","date":"2024-05-30","arxiv_id":"2405.20446","n_code_links":0,"syntology":null},{"paper":null,"slug":"kerascv-and-kerasnlp-vision-and-language","title":"KerasCV and KerasNLP: Vision and Language Power-Ups","date":"2024-05-30","arxiv_id":"2405.20247","n_code_links":0,"syntology":null},{"paper":null,"slug":"knowledge-graph-tuning-real-time-large","title":"Knowledge Graph Tuning: Real-time Large Language Model Personalization based on Human Feedback","date":"2024-05-30","arxiv_id":"2405.19686","n_code_links":0,"syntology":null},{"paper":"/paper/llamea-a-large-language-model-evolutionary","slug":"llamea-a-large-language-model-evolutionary","title":"LLaMEA: A Large Language Model Evolutionary Algorithm for Automatically Generating Metaheuristics","date":"2024-05-30","arxiv_id":"2405.20132","n_code_links":2,"syntology":{"ran":6,"of":6,"n_ran_checked":6,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["nikivanstein/LLaMEA"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["listed","official"]}}},{"paper":"/paper/one-token-can-help-learning-scalable-and","slug":"one-token-can-help-learning-scalable-and","title":"One Token Can Help! Learning Scalable and Pluggable Virtual Tokens for Retrieval-Augmented Large Language Models","date":"2024-05-30","arxiv_id":"2405.19670","n_code_links":2,"syntology":null},{"paper":null,"slug":"phantom-general-trigger-attacks-on-retrieval","title":"Phantom: General Trigger Attacks on Retrieval Augmented Language Generation","date":"2024-05-30","arxiv_id":"2405.20485","n_code_links":0,"syntology":null},{"paper":null,"slug":"robo-instruct-simulator-augmented-instruction","title":"Robo-Instruct: Simulator-Augmented Instruction Alignment For Finetuning Code LLMs","date":"2024-05-30","arxiv_id":"2405.20179","n_code_links":0,"syntology":null},{"paper":null,"slug":"significance-of-chain-of-thought-in-gender","title":"Significance of Chain of Thought in Gender Bias Mitigation for English-Dravidian Machine Translation","date":"2024-05-30","arxiv_id":"2405.19701","n_code_links":0,"syntology":null},{"paper":"/paper/student-answer-forecasting-transformer-driven","slug":"student-answer-forecasting-transformer-driven","title":"Student Answer Forecasting: Transformer-Driven Answer Choice Prediction for Language Learning","date":"2024-05-30","arxiv_id":"2405.20079","n_code_links":1,"syntology":null},{"paper":"/paper/towards-ontology-enhanced-representation","slug":"towards-ontology-enhanced-representation","title":"Towards Ontology-Enhanced Representation Learning for Large Language Models","date":"2024-05-30","arxiv_id":"2405.20527","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-multi-source-retrieval-question-answering","title":"A Multi-Source Retrieval Question Answering Framework Based on RAG","date":"2024-05-29","arxiv_id":"2405.19207","n_code_links":0,"syntology":null},{"paper":"/paper/beyond-agreement-diagnosing-the-rationale","slug":"beyond-agreement-diagnosing-the-rationale","title":"Beyond Agreement: Diagnosing the Rationale Alignment of Automated Essay Scoring Methods based on Linguistically-informed Counterfactuals","date":"2024-05-29","arxiv_id":"2405.19433","n_code_links":1,"syntology":null},{"paper":null,"slug":"can-gpt-redefine-medical-understanding","title":"Can GPT Redefine Medical Understanding? Evaluating GPT on Biomedical Machine Reading Comprehension","date":"2024-05-29","arxiv_id":"2405.18682","n_code_links":0,"syntology":null},{"paper":"/paper/ctrla-adaptive-retrieval-augmented-generation","slug":"ctrla-adaptive-retrieval-augmented-generation","title":"CtrlA: Adaptive Retrieval-Augmented Generation via Inherent Control","date":"2024-05-29","arxiv_id":"2405.18727","n_code_links":1,"syntology":{"ran":7,"of":12,"n_ran_checked":7,"n_instrument":0,"unverified":5,"pointer_only":0,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 0 where Syntology's instrument failed) · 5 unverified","official":{"repos":["hsliu-initial/ctrla"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":7,"n_unverified":5,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"efficient-model-agnostic-alignment-via","title":"Efficient Model-agnostic Alignment via Bayesian Persuasion","date":"2024-05-29","arxiv_id":"2405.18718","n_code_links":0,"syntology":null},{"paper":null,"slug":"faster-cascades-via-speculative-decoding","title":"Faster Cascades via Speculative Decoding","date":"2024-05-29","arxiv_id":"2405.19261","n_code_links":0,"syntology":null},{"paper":null,"slug":"lmo-dp-optimizing-the-randomization-mechanism","title":"LMO-DP: Optimizing the Randomization Mechanism for Differentially Private Fine-Tuning (Large) Language Models","date":"2024-05-29","arxiv_id":"2405.18776","n_code_links":0,"syntology":null},{"paper":"/paper/map-neo-highly-capable-and-transparent","slug":"map-neo-highly-capable-and-transparent","title":"MAP-Neo: Highly Capable and Transparent Bilingual Large Language Model Series","date":"2024-05-29","arxiv_id":"2405.19327","n_code_links":1,"syntology":{"ran":12,"of":14,"n_ran_checked":12,"n_instrument":0,"unverified":2,"pointer_only":14,"phrase":"12 ran (of which 0 constructed an object rather than computing a result; 12 with no instrument failure: 0 honoured, 0 violated, 12 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["multimodal-art-projection/map-neo"],"state":"official (archive's flag): 12 ran","n_ran":12,"n_constructed":0,"n_ran_no_instrument_failure":12,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"offline-regularised-reinforcement-learning","title":"Offline Regularised Reinforcement Learning for Large Language Models Alignment","date":"2024-05-29","arxiv_id":"2405.19107","n_code_links":0,"syntology":null},{"paper":null,"slug":"stat-shrinking-transformers-after-training","title":"STAT: Shrinking Transformers After Training","date":"2024-05-29","arxiv_id":"2406.00061","n_code_links":0,"syntology":null},{"paper":"/paper/toward-conversational-agents-with-context-and","slug":"toward-conversational-agents-with-context-and","title":"Toward Conversational Agents with Context and Time Sensitive Long-term Memory","date":"2024-05-29","arxiv_id":"2406.00057","n_code_links":1,"syntology":null},{"paper":null,"slug":"two-layer-retrieval-augmented-generation","title":"Two-Layer Retrieval-Augmented Generation Framework for Low-Resource Medical Question Answering Using Reddit Data: Proof-of-Concept Study","date":"2024-05-29","arxiv_id":"2405.19519","n_code_links":0,"syntology":null},{"paper":"/paper/aligning-to-thousands-of-preferences-via","slug":"aligning-to-thousands-of-preferences-via","title":"Aligning to Thousands of Preferences via System Message Generalization","date":"2024-05-28","arxiv_id":"2405.17977","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":1,"n_instrument":2,"unverified":0,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 1 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":{"repos":["kaistAI/Janus"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["community","official"]}}},{"paper":"/paper/an-empirical-analysis-on-large-language","slug":"an-empirical-analysis-on-large-language","title":"An Empirical Analysis on Large Language Models in Debate Evaluation","date":"2024-05-28","arxiv_id":"2406.00050","n_code_links":1,"syntology":{"ran":1,"of":2,"n_ran_checked":1,"n_instrument":0,"unverified":1,"pointer_only":2,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["xinyiliu0227/llm_debate_bias"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"are-ppo-ed-language-models-hackable","title":"Are PPO-ed Language Models Hackable?","date":"2024-05-28","arxiv_id":"2406.02577","n_code_links":0,"syntology":null},{"paper":"/paper/atm-adversarial-tuning-multi-agent-system","slug":"atm-adversarial-tuning-multi-agent-system","title":"ATM: Adversarial Tuning Multi-agent System Makes a Robust Retrieval-Augmented Generator","date":"2024-05-28","arxiv_id":"2405.18111","n_code_links":1,"syntology":{"ran":6,"of":10,"n_ran_checked":6,"n_instrument":0,"unverified":4,"pointer_only":10,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","official":{"repos":["chuhac/atm-rag"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"attention-based-sequential-recommendation","title":"Attention-based sequential recommendation system using multimodal data","date":"2024-05-28","arxiv_id":"2405.17959","n_code_links":0,"syntology":null},{"paper":null,"slug":"don-t-forget-to-connect-improving-rag-with","title":"Don't Forget to Connect! Improving RAG with Graph-based Reranking","date":"2024-05-28","arxiv_id":"2405.18414","n_code_links":0,"syntology":null},{"paper":null,"slug":"edinburgh-clinical-nlp-at-mediqa-corr-2024","title":"Edinburgh Clinical NLP at MEDIQA-CORR 2024: Guiding Large Language Models with Hints","date":"2024-05-28","arxiv_id":"2405.18028","n_code_links":0,"syntology":null},{"paper":"/paper/llms-and-memorization-on-quality-and","slug":"llms-and-memorization-on-quality-and","title":"LLMs and Memorization: On Quality and Specificity of Copyright Compliance","date":"2024-05-28","arxiv_id":"2405.18492","n_code_links":1,"syntology":null},{"paper":null,"slug":"understanding-intrinsic-socioeconomic-biases","title":"Understanding Intrinsic Socioeconomic Biases in Large Language Models","date":"2024-05-28","arxiv_id":"2405.18662","n_code_links":0,"syntology":null},{"paper":null,"slug":"widin-wording-image-for-domain-invariant","title":"WIDIn: Wording Image for Domain-Invariant Representation in Single-Source Domain Generalization","date":"2024-05-28","arxiv_id":"2405.18405","n_code_links":0,"syntology":null},{"paper":"/paper/assessing-llms-suitability-for-knowledge","slug":"assessing-llms-suitability-for-knowledge","title":"Assessing LLMs Suitability for Knowledge Graph Completion","date":"2024-05-27","arxiv_id":"2405.17249","n_code_links":1,"syntology":null},{"paper":null,"slug":"augmenting-textual-generation-via-topology","title":"Augmenting Textual Generation via Topology Aware Retrieval","date":"2024-05-27","arxiv_id":"2405.17602","n_code_links":0,"syntology":null},{"paper":"/paper/deeperimpact-optimizing-sparse-learned-index","slug":"deeperimpact-optimizing-sparse-learned-index","title":"DeeperImpact: Optimizing Sparse Learned Index Structures","date":"2024-05-27","arxiv_id":"2405.17093","n_code_links":1,"syntology":null},{"paper":null,"slug":"detecting-deceptive-dark-patterns-in-e","title":"Detecting Deceptive Dark Patterns in E-commerce Platforms","date":"2024-05-27","arxiv_id":"2406.01608","n_code_links":0,"syntology":null},{"paper":null,"slug":"exploiting-the-layered-intrinsic","title":"Exploiting the Layered Intrinsic Dimensionality of Deep Models for Practical Adversarial Training","date":"2024-05-27","arxiv_id":"2405.17130","n_code_links":0,"syntology":null},{"paper":"/paper/inversionview-a-general-purpose-method-for","slug":"inversionview-a-general-purpose-method-for","title":"InversionView: A General-Purpose Method for Reading Information from Neural Activations","date":"2024-05-27","arxiv_id":"2405.17653","n_code_links":1,"syntology":{"ran":1,"of":6,"n_ran_checked":1,"n_instrument":0,"unverified":5,"pointer_only":6,"phrase":"1 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 5 unverified; the one sample that ran constructed an object rather than computing a result","official":{"repos":["huangxt39/inversionview"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":5,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"llm-based-cooperative-agents-using","title":"REVECA: Adaptive Planning and Trajectory-based Validation in Cooperative Language Agents using Information Relevance and Relative Proximity","date":"2024-05-27","arxiv_id":"2405.16751","n_code_links":0,"syntology":null},{"paper":null,"slug":"nv-embed-improved-techniques-for-training","title":"NV-Embed: Improved Techniques for Training LLMs as Generalist Embedding Models","date":"2024-05-27","arxiv_id":"2405.17428","n_code_links":0,"syntology":null},{"paper":null,"slug":"pae-llm-based-product-attribute-extraction","title":"PAE: LLM-based Product Attribute Extraction for E-Commerce Fashion Trends","date":"2024-05-27","arxiv_id":"2405.17533","n_code_links":0,"syntology":null},{"paper":null,"slug":"performance-evaluation-of-reddit-comments","title":"Performance evaluation of Reddit Comments using Machine Learning and Natural Language Processing methods in Sentiment Analysis","date":"2024-05-27","arxiv_id":"2405.16810","n_code_links":0,"syntology":null},{"paper":null,"slug":"qub-cirdan-at-discharge-me-zero-shot","title":"QUB-Cirdan at \"Discharge Me!\": Zero shot discharge letter generation by open-source LLM","date":"2024-05-27","arxiv_id":"2406.00041","n_code_links":0,"syntology":null},{"paper":"/paper/reflectioncoder-learning-from-reflection","slug":"reflectioncoder-learning-from-reflection","title":"ReflectionCoder: Learning from Reflection Sequence for Enhanced One-off Code Generation","date":"2024-05-27","arxiv_id":"2405.17057","n_code_links":1,"syntology":{"ran":0,"of":1,"n_ran_checked":0,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"0 ran · 1 unverified","official":{"repos":["sensellm/reflectioncoder"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":[]}}},{"paper":"/paper/rtl-repo-a-benchmark-for-evaluating-llms-on","slug":"rtl-repo-a-benchmark-for-evaluating-llms-on","title":"RTL-Repo: A Benchmark for Evaluating LLMs on Large-Scale RTL Design Projects","date":"2024-05-27","arxiv_id":"2405.17378","n_code_links":1,"syntology":null},{"paper":null,"slug":"the-scaling-law-in-stellar-light-curves","title":"The Scaling Law in Stellar Light Curves","date":"2024-05-27","arxiv_id":"2405.17156","n_code_links":0,"syntology":null},{"paper":"/paper/thread-thinking-deeper-with-recursive","slug":"thread-thinking-deeper-with-recursive","title":"THREAD: Thinking Deeper with Recursive Spawning","date":"2024-05-27","arxiv_id":"2405.17402","n_code_links":1,"syntology":{"ran":7,"of":11,"n_ran_checked":6,"n_instrument":1,"unverified":4,"pointer_only":11,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 1 where Syntology's instrument failed) · 4 unverified","official":{"repos":["philipmit/thread"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":"/paper/video-enriched-retrieval-augmented-generation","slug":"video-enriched-retrieval-augmented-generation","title":"Video Enriched Retrieval Augmented Generation Using Aligned Video Captions","date":"2024-05-27","arxiv_id":"2405.17706","n_code_links":1,"syntology":null},{"paper":null,"slug":"ai-generated-text-detection-and","title":"AI-Generated Text Detection and Classification Based on BERT Deep Learning Algorithm","date":"2024-05-26","arxiv_id":"2405.16422","n_code_links":0,"syntology":null},{"paper":"/paper/grag-graph-retrieval-augmented-generation","slug":"grag-graph-retrieval-augmented-generation","title":"GRAG: Graph Retrieval-Augmented Generation","date":"2024-05-26","arxiv_id":"2405.16506","n_code_links":1,"syntology":null},{"paper":null,"slug":"m-rag-reinforcing-large-language-model","title":"M-RAG: Reinforcing Large Language Model Performance through Retrieval-Augmented Generation with Multiple Partitions","date":"2024-05-26","arxiv_id":"2405.16420","n_code_links":0,"syntology":null},{"paper":null,"slug":"accelerating-inference-of-retrieval-augmented","title":"Accelerating Inference of Retrieval-Augmented Generation via Sparse Context Selection","date":"2024-05-25","arxiv_id":"2405.16178","n_code_links":0,"syntology":null},{"paper":"/paper/accelerating-transformers-with-spectrum-1","slug":"accelerating-transformers-with-spectrum-1","title":"Accelerating Transformers with Spectrum-Preserving Token Merging","date":"2024-05-25","arxiv_id":"2405.16148","n_code_links":1,"syntology":{"ran":6,"of":7,"n_ran_checked":2,"n_instrument":4,"unverified":1,"pointer_only":7,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 2 honoured, 0 violated, 0 with no contract checked; 4 where Syntology's instrument failed) · 1 unverified","official":{"repos":["hchautran/PiToMe"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/automanual-generating-instruction-manuals-by","slug":"automanual-generating-instruction-manuals-by","title":"AutoManual: Constructing Instruction Manuals by LLM Agents via Interactive Environmental Learning","date":"2024-05-25","arxiv_id":"2405.16247","n_code_links":1,"syntology":{"ran":4,"of":4,"n_ran_checked":4,"n_instrument":0,"unverified":0,"pointer_only":4,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["minghchen/automanual"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"incremental-comprehension-of-garden-path","title":"Incremental Comprehension of Garden-Path Sentences by Large Language Models: Semantic Interpretation, Syntactic Re-Analysis, and Attention","date":"2024-05-25","arxiv_id":"2405.16042","n_code_links":0,"syntology":null},{"paper":null,"slug":"mindstar-enhancing-math-reasoning-in-pre","title":"MindStar: Enhancing Math Reasoning in Pre-trained LLMs at Inference Time","date":"2024-05-25","arxiv_id":"2405.16265","n_code_links":0,"syntology":null},{"paper":null,"slug":"towards-unlocking-insights-from-logbooks","title":"Towards Unlocking Insights from Logbooks Using AI","date":"2024-05-25","arxiv_id":"2406.12881","n_code_links":0,"syntology":null},{"paper":null,"slug":"an-evaluation-of-estimative-uncertainty-in","title":"An Evaluation of Estimative Uncertainty in Large Language Models","date":"2024-05-24","arxiv_id":"2405.15185","n_code_links":0,"syntology":null},{"paper":null,"slug":"benchmarking-pre-trained-large-language","title":"Benchmarking the Performance of Pre-trained LLMs across Urdu NLP Tasks","date":"2024-05-24","arxiv_id":"2405.15453","n_code_links":0,"syntology":null},{"paper":"/paper/culturepark-boosting-cross-cultural","slug":"culturepark-boosting-cross-cultural","title":"CulturePark: Boosting Cross-cultural Understanding in Large Language Models","date":"2024-05-24","arxiv_id":"2405.15145","n_code_links":1,"syntology":{"ran":0,"of":7,"n_ran_checked":0,"n_instrument":0,"unverified":7,"pointer_only":7,"phrase":"0 ran · 7 unverified","official":{"repos":["scarelette/culturepark"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":7,"ran_from_kinds":[]}}},{"paper":null,"slug":"enhancing-augmentative-and-alternative","title":"Enhancing Augmentative and Alternative Communication with Card Prediction and Colourful Semantics","date":"2024-05-24","arxiv_id":"2405.15896","n_code_links":0,"syntology":null}],"record_sha256":"80da3c2163732696245460368feb2f39e5834666f8714a867b319a94f0e12dca","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}