{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/weight-decay/papers/21","list_of":"/method/weight-decay","method":"Weight Decay","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":21,"pages_in_order":108,"rows_per_page":100,"rows":[2001,2100],"of":10713,"counts":{"archive_papers_tagged":10713,"with_a_code_link":4533,"where_syntology_ran_a_sample":1291,"not_listed_spam_title":0,"listed":10713,"listed_where_code_ran":1291,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":1064,"every_run_a_failure_of_syntologys_instrument":227,"listed_with_a_run_with_no_instrument_failure":1064,"listed_every_run_a_failure_of_syntologys_instrument":227,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/weight-decay","prev":"/method/weight-decay/papers/20","next":"/method/weight-decay/papers/22","papers":[{"paper":null,"slug":"hypa-rag-a-hybrid-parameter-adaptive","title":"HyPA-RAG: A Hybrid Parameter Adaptive Retrieval-Augmented Generation System for AI Legal and Policy Applications","date":"2024-08-29","arxiv_id":"2409.09046","n_code_links":0,"syntology":null},{"paper":"/paper/llava-chef-a-multi-modal-generative-model-for","slug":"llava-chef-a-multi-modal-generative-model-for","title":"LLaVA-Chef: A Multi-modal Generative Model for Food Recipes","date":"2024-08-29","arxiv_id":"2408.16889","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-simple-baseline-with-single-encoder-for","title":"A Simple Baseline with Single-encoder for Referring Image Segmentation","date":"2024-08-28","arxiv_id":"2408.15521","n_code_links":0,"syntology":null},{"paper":null,"slug":"an-extremely-data-efficient-and-generative","title":"An Extremely Data-efficient and Generative LLM-based Reinforcement Learning Agent for Recommenders","date":"2024-08-28","arxiv_id":"2408.16032","n_code_links":0,"syntology":null},{"paper":"/paper/cbf-llm-safe-control-for-llm-alignment","slug":"cbf-llm-safe-control-for-llm-alignment","title":"CBF-LLM: Safe Control for LLM Alignment","date":"2024-08-28","arxiv_id":"2408.15625","n_code_links":1,"syntology":null},{"paper":null,"slug":"conan-embedding-general-text-embedding-with","title":"Conan-embedding: General Text Embedding with More and Better Negative Samples","date":"2024-08-28","arxiv_id":"2408.15710","n_code_links":0,"syntology":null},{"paper":null,"slug":"fractured-sorry-bench-framework-for-revealing","title":"FRACTURED-SORRY-Bench: Framework for Revealing Attacks in Conversational Turns Undermining Refusal Efficacy and Defenses over SORRY-Bench (Automated Multi-shot Jailbreaks)","date":"2024-08-28","arxiv_id":"2408.16163","n_code_links":0,"syntology":null},{"paper":null,"slug":"is-personality-prediction-possible-based-on","title":"Is Personality Prediction Possible Based on Reddit Comments?","date":"2024-08-28","arxiv_id":"2408.16089","n_code_links":0,"syntology":null},{"paper":null,"slug":"leveraging-large-language-models-for-wireless","title":"Leveraging Large Language Models for Wireless Symbol Detection via In-Context Learning","date":"2024-08-28","arxiv_id":"2409.00124","n_code_links":0,"syntology":null},{"paper":"/paper/lrp4rag-detecting-hallucinations-in-retrieval","slug":"lrp4rag-detecting-hallucinations-in-retrieval","title":"LRP4RAG: Detecting Hallucinations in Retrieval-Augmented Generation via Layer-wise Relevance Propagation","date":"2024-08-28","arxiv_id":"2408.15533","n_code_links":3,"syntology":null},{"paper":"/paper/unleashing-the-temporal-spatial-reasoning","slug":"unleashing-the-temporal-spatial-reasoning","title":"Unleashing the Temporal-Spatial Reasoning Capacity of GPT for Training-Free Audio and Language Referenced Video Object Segmentation","date":"2024-08-28","arxiv_id":"2408.15876","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-survey-of-large-language-models-for-3","title":"A Survey of Large Language Models for European Languages","date":"2024-08-27","arxiv_id":"2408.15040","n_code_links":0,"syntology":null},{"paper":"/paper/into-the-unknown-unknowns-engaged-human","slug":"into-the-unknown-unknowns-engaged-human","title":"Into the Unknown Unknowns: Engaged Human Learning through Participation in Language Model Agent Conversations","date":"2024-08-27","arxiv_id":"2408.15232","n_code_links":1,"syntology":null},{"paper":null,"slug":"strategic-optimization-and-challenges-of","title":"Strategic Optimization and Challenges of Large Language Models in Object-Oriented Programming","date":"2024-08-27","arxiv_id":"2408.14834","n_code_links":0,"syntology":null},{"paper":null,"slug":"zero-shot-visual-reasoning-by-vision-language","title":"Zero-Shot Visual Reasoning by Vision-Language Models: Benchmarking and Analysis","date":"2024-08-27","arxiv_id":"2409.00106","n_code_links":0,"syntology":null},{"paper":"/paper/chartom-a-visual-theory-of-mind-benchmark-for","slug":"chartom-a-visual-theory-of-mind-benchmark-for","title":"CHARTOM: A Visual Theory-of-Mind Benchmark for Multimodal Large Language Models","date":"2024-08-26","arxiv_id":"2408.14419","n_code_links":1,"syntology":null},{"paper":"/paper/question-answering-system-of-bridge-design","slug":"question-answering-system-of-bridge-design","title":"Question answering system of bridge design specification based on large language model","date":"2024-08-26","arxiv_id":"2408.13282","n_code_links":1,"syntology":null},{"paper":null,"slug":"bidirectional-awareness-induction-in","title":"Bidirectional Awareness Induction in Autoregressive Seq2Seq Models","date":"2024-08-25","arxiv_id":"2408.13959","n_code_links":0,"syntology":null},{"paper":"/paper/codegraph-enhancing-graph-reasoning-of-llms","slug":"codegraph-enhancing-graph-reasoning-of-llms","title":"CodeGraph: Enhancing Graph Reasoning of LLMs with Code","date":"2024-08-25","arxiv_id":"2408.13863","n_code_links":1,"syntology":null},{"paper":null,"slug":"lowclip-adapting-the-clip-model-architecture","title":"LowCLIP: Adapting the CLIP Model Architecture for Low-Resource Languages in Multimodal Image Retrieval Task","date":"2024-08-25","arxiv_id":"2408.13909","n_code_links":0,"syntology":null},{"paper":"/paper/vision-language-and-large-language-model","slug":"vision-language-and-large-language-model","title":"Vision-Language and Large Language Model Performance in Gastroenterology: GPT, Claude, Llama, Phi, Mistral, Gemma, and Quantized Models","date":"2024-08-25","arxiv_id":"2409.00084","n_code_links":1,"syntology":null},{"paper":"/paper/pandora-s-box-or-aladdin-s-lamp-a","slug":"pandora-s-box-or-aladdin-s-lamp-a","title":"Pandora's Box or Aladdin's Lamp: A Comprehensive Analysis Revealing the Role of RAG Noise in Large Language Models","date":"2024-08-24","arxiv_id":"2408.13533","n_code_links":1,"syntology":{"ran":2,"of":3,"n_ran_checked":1,"n_instrument":1,"unverified":1,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":{"repos":["jinyangwu/NoiserBench"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"data-exposure-from-llm-apps-an-in-depth","title":"An In-Depth Investigation of Data Collection in LLM App Ecosystems","date":"2024-08-23","arxiv_id":"2408.13247","n_code_links":0,"syntology":null},{"paper":null,"slug":"enhancing-llm-based-automated-program-repair","title":"Enhancing Automated Program Repair with Solution Design","date":"2024-08-22","arxiv_id":"2408.12056","n_code_links":0,"syntology":null},{"paper":null,"slug":"enhancing-multi-hop-reasoning-through","title":"Enhancing Multi-hop Reasoning through Knowledge Erasure in Large Language Model Editing","date":"2024-08-22","arxiv_id":"2408.12456","n_code_links":0,"syntology":null},{"paper":"/paper/graph-retrieval-augmented-trustworthiness","slug":"graph-retrieval-augmented-trustworthiness","title":"GRATR: Zero-Shot Evidence Graph Retrieval-Augmented Trustworthiness Reasoning","date":"2024-08-22","arxiv_id":"2408.12333","n_code_links":1,"syntology":null},{"paper":null,"slug":"large-language-models-as-foundations-for-next","title":"Large Language Models as Foundations for Next-Gen Dense Retrieval: A Comprehensive Empirical Assessment","date":"2024-08-22","arxiv_id":"2408.12194","n_code_links":0,"syntology":null},{"paper":null,"slug":"llms-are-not-zero-shot-reasoners-for","title":"LLMs are not Zero-Shot Reasoners for Biomedical Information Extraction","date":"2024-08-22","arxiv_id":"2408.12249","n_code_links":0,"syntology":null},{"paper":null,"slug":"optimizing-performance-how-compact-models","title":"Optimizing Performance: How Compact Models Match or Exceed GPT's Classification Capabilities through Fine-Tuning","date":"2024-08-22","arxiv_id":"2409.11408","n_code_links":0,"syntology":null},{"paper":null,"slug":"unlearning-trojans-in-large-language-models-a","title":"Unlearning Trojans in Large Language Models: A Comparison Between Natural Language and Source Code","date":"2024-08-22","arxiv_id":"2408.12416","n_code_links":0,"syntology":null},{"paper":"/paper/a-quick-trustworthy-spectral-detection-q-a","slug":"a-quick-trustworthy-spectral-detection-q-a","title":"A Quick, trustworthy spectral knowledge Q&A system leveraging retrieval-augmented generation on LLM","date":"2024-08-21","arxiv_id":"2408.11557","n_code_links":1,"syntology":null},{"paper":"/paper/ancient-wisdom-modern-tools-exploring","slug":"ancient-wisdom-modern-tools-exploring","title":"Ancient Wisdom, Modern Tools: Exploring Retrieval-Augmented LLMs for Ancient Indian Philosophy","date":"2024-08-21","arxiv_id":"2408.11903","n_code_links":1,"syntology":null},{"paper":null,"slug":"applying-and-evaluating-large-language-models","title":"Applying and Evaluating Large Language Models in Mental Health Care: A Scoping Review of Human-Assessed Generative Tasks","date":"2024-08-21","arxiv_id":"2408.11288","n_code_links":0,"syntology":null},{"paper":"/paper/approaching-deep-learning-through-the","slug":"approaching-deep-learning-through-the","title":"Approaching Deep Learning through the Spectral Dynamics of Weights","date":"2024-08-21","arxiv_id":"2408.11804","n_code_links":1,"syntology":{"ran":9,"of":14,"n_ran_checked":9,"n_instrument":0,"unverified":5,"pointer_only":0,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 9 with no instrument failure: 0 honoured, 0 violated, 9 with no contract checked; 0 where Syntology's instrument failed) · 5 unverified","official":{"repos":["dyunis/spectral_dynamics"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":9,"n_unverified":5,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"d-rmgpt-robot-assisted-collaborative-tasks","title":"D-RMGPT: Robot-assisted collaborative tasks driven by large multimodal models","date":"2024-08-21","arxiv_id":"2408.11761","n_code_links":0,"syntology":null},{"paper":null,"slug":"mixed-sparsity-training-achieving-4-times","title":"Mixed Sparsity Training: Achieving 4$\\times$ FLOP Reduction for Transformer Pretraining","date":"2024-08-21","arxiv_id":"2408.11746","n_code_links":0,"syntology":null},{"paper":null,"slug":"permitqa-a-benchmark-for-retrieval-augmented","title":"WeQA: A Benchmark for Retrieval Augmented Generation in Wind Energy Domain","date":"2024-08-21","arxiv_id":"2408.11800","n_code_links":0,"syntology":null},{"paper":null,"slug":"rag-optimized-tibetan-tourism-llms-enhancing","title":"RAG-Optimized Tibetan Tourism LLMs: Enhancing Accuracy and Personalization","date":"2024-08-21","arxiv_id":"2408.12003","n_code_links":0,"syntology":null},{"paper":"/paper/raglab-a-modular-and-research-oriented","slug":"raglab-a-modular-and-research-oriented","title":"RAGLAB: A Modular and Research-Oriented Unified Framework for Retrieval-Augmented Generation","date":"2024-08-21","arxiv_id":"2408.11381","n_code_links":1,"syntology":null},{"paper":null,"slug":"the-self-contained-negation-test-set","title":"The Self-Contained Negation Test Set","date":"2024-08-21","arxiv_id":"2408.11469","n_code_links":0,"syntology":null},{"paper":"/paper/unlocking-adversarial-suffix-optimization","slug":"unlocking-adversarial-suffix-optimization","title":"Unlocking Adversarial Suffix Optimization Without Affirmative Phrases: Efficient Black-box Jailbreaking via LLM as Optimizer","date":"2024-08-21","arxiv_id":"2408.11313","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":1,"n_instrument":1,"unverified":0,"pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 1 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["lenijwp/eclipse"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"ctp-llm-clinical-trial-phase-transition","title":"CTP-LLM: Clinical Trial Phase Transition Prediction Using Large Language Models","date":"2024-08-20","arxiv_id":"2408.10995","n_code_links":0,"syntology":null},{"paper":null,"slug":"how-well-do-large-language-models-serve-as","title":"How Well Do Large Language Models Serve as End-to-End Secure Code Agents for Python?","date":"2024-08-20","arxiv_id":"2408.10495","n_code_links":0,"syntology":null},{"paper":"/paper/language-modeling-on-tabular-data-a-survey-of","slug":"language-modeling-on-tabular-data-a-survey-of","title":"Language Modeling on Tabular Data: A Survey of Foundations, Techniques and Evolution","date":"2024-08-20","arxiv_id":"2408.10548","n_code_links":1,"syntology":null},{"paper":null,"slug":"reading-with-intent","title":"Reading with Intent","date":"2024-08-20","arxiv_id":"2408.11189","n_code_links":0,"syntology":null},{"paper":null,"slug":"reconciling-methodological-paradigms","title":"Reconciling Methodological Paradigms: Employing Large Language Models as Novice Qualitative Research Assistants in Talent Management Research","date":"2024-08-20","arxiv_id":"2408.11043","n_code_links":0,"syntology":null},{"paper":"/paper/soda-eval-open-domain-dialogue-evaluation-in","slug":"soda-eval-open-domain-dialogue-evaluation-in","title":"Soda-Eval: Open-Domain Dialogue Evaluation in the age of LLMs","date":"2024-08-20","arxiv_id":"2408.10902","n_code_links":1,"syntology":null},{"paper":null,"slug":"towardseffective-teaching-assistants-from","title":"Towardseffective teaching assistants: From intent-based chatbots to LLM-poweredteachingassistants","date":"2024-08-20","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"tracing-privacy-leakage-of-language-models-to","title":"Tracing Privacy Leakage of Language Models to Training Data via Adjusted Influence Functions","date":"2024-08-20","arxiv_id":"2408.10468","n_code_links":0,"syntology":null},{"paper":"/paper/while-github-copilot-excels-at-coding-does-it","slug":"while-github-copilot-excels-at-coding-does-it","title":"Security Attacks on LLM-based Code Completion Tools","date":"2024-08-20","arxiv_id":"2408.11006","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-strategy-to-combine-1stgen-transformers-and","title":"A Strategy to Combine 1stGen Transformers and Open LLMs for Automatic Text Classification","date":"2024-08-19","arxiv_id":"2408.09629","n_code_links":0,"syntology":null},{"paper":"/paper/acquiring-bidirectionality-via-large-and","slug":"acquiring-bidirectionality-via-large-and","title":"Acquiring Bidirectionality via Large and Small Language Models","date":"2024-08-19","arxiv_id":"2408.09640","n_code_links":1,"syntology":null},{"paper":null,"slug":"active-learning-for-identifying-disaster","title":"Active Learning for Identifying Disaster-Related Tweets: A Comparison with Keyword Filtering and Generic Fine-Tuning","date":"2024-08-19","arxiv_id":"2408.09914","n_code_links":0,"syntology":null},{"paper":null,"slug":"carbon-footprint-accounting-driven-by-large","title":"Carbon Footprint Accounting Driven by Large Language Models and Retrieval-augmented Generation","date":"2024-08-19","arxiv_id":"2408.09713","n_code_links":0,"syntology":null},{"paper":"/paper/enhance-lifelong-model-editing-with","slug":"enhance-lifelong-model-editing-with","title":"ELDER: Enhancing Lifelong Model Editing with Mixture-of-LoRA","date":"2024-08-19","arxiv_id":"2408.11869","n_code_links":1,"syntology":null},{"paper":null,"slug":"enhanced-document-retrieval-with-topic","title":"Enhanced document retrieval with topic embeddings","date":"2024-08-19","arxiv_id":"2408.10435","n_code_links":0,"syntology":null},{"paper":null,"slug":"gpt-augmented-reinforcement-learning-with","title":"GARLIC: GPT-Augmented Reinforcement Learning with Intelligent Control for Vehicle Dispatching","date":"2024-08-19","arxiv_id":"2408.10286","n_code_links":0,"syntology":null},{"paper":"/paper/legalbench-rag-a-benchmark-for-retrieval","slug":"legalbench-rag-a-benchmark-for-retrieval","title":"LegalBench-RAG: A Benchmark for Retrieval-Augmented Generation in the Legal Domain","date":"2024-08-19","arxiv_id":"2408.10343","n_code_links":1,"syntology":{"ran":7,"of":8,"n_ran_checked":7,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["zeroentropy-cc/legalbenchrag"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":7,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"rhyme-aware-chinese-lyric-generator-based-on","title":"Rhyme-aware Chinese lyric generator based on GPT","date":"2024-08-19","arxiv_id":"2408.10130","n_code_links":0,"syntology":null},{"paper":null,"slug":"tba-faster-large-language-model-training","title":"SSDTrain: An Activation Offloading Framework to SSDs for Faster Large Language Model Training","date":"2024-08-19","arxiv_id":"2408.10013","n_code_links":0,"syntology":null},{"paper":null,"slug":"agentic-retrieval-augmented-generation-for","title":"Agentic Retrieval-Augmented Generation for Time Series Analysis","date":"2024-08-18","arxiv_id":"2408.14484","n_code_links":0,"syntology":null},{"paper":null,"slug":"clustering-and-alignment-understanding-the","title":"Clustering and Alignment: Understanding the Training Dynamics in Modular Addition","date":"2024-08-18","arxiv_id":"2408.09414","n_code_links":0,"syntology":null},{"paper":null,"slug":"conversum-a-contrastive-learning-based","title":"ConVerSum: A Contrastive Learning-based Approach for Data-Scarce Solution of Cross-Lingual Summarization Beyond Direct Equivalents","date":"2024-08-17","arxiv_id":"2408.09273","n_code_links":0,"syntology":null},{"paper":null,"slug":"sentiment-analysis-of-preservice-teachers","title":"Sentiment analysis of preservice teachers' reflections using a large language model","date":"2024-08-17","arxiv_id":"2408.11862","n_code_links":0,"syntology":null},{"paper":null,"slug":"tablebench-a-comprehensive-and-complex","title":"TableBench: A Comprehensive and Complex Benchmark for Table Question Answering","date":"2024-08-17","arxiv_id":"2408.09174","n_code_links":0,"syntology":null},{"paper":"/paper/tc-rag-turing-complete-rag-s-case-study-on","slug":"tc-rag-turing-complete-rag-s-case-study-on","title":"TC-RAG:Turing-Complete RAG's Case study on Medical LLM Systems","date":"2024-08-17","arxiv_id":"2408.09199","n_code_links":2,"syntology":null},{"paper":null,"slug":"a-mean-field-ansatz-for-zero-shot-weight","title":"A Mean Field Ansatz for Zero-Shot Weight Transfer","date":"2024-08-16","arxiv_id":"2408.08681","n_code_links":0,"syntology":null},{"paper":null,"slug":"cikmar-a-dual-encoder-approach-to-prompt","title":"CIKMar: A Dual-Encoder Approach to Prompt-Based Reranking in Educational Dialogue Systems","date":"2024-08-16","arxiv_id":"2408.08805","n_code_links":0,"syntology":null},{"paper":"/paper/communitykg-rag-leveraging-community","slug":"communitykg-rag-leveraging-community","title":"CommunityKG-RAG: Leveraging Community Structures in Knowledge Graphs for Advanced Retrieval-Augmented Generation in Fact-Checking","date":"2024-08-16","arxiv_id":"2408.08535","n_code_links":1,"syntology":null},{"paper":"/paper/fine-tuning-llms-for-autonomous-spacecraft","slug":"fine-tuning-llms-for-autonomous-spacecraft","title":"Fine-tuning LLMs for Autonomous Spacecraft Control: A Case Study Using Kerbal Space Program","date":"2024-08-16","arxiv_id":"2408.08676","n_code_links":1,"syntology":null},{"paper":null,"slug":"improving-vte-identification-through-language","title":"Improving VTE Identification through Language Models from Radiology Reports: A Comparative Study of Mamba, Phi-3 Mini, and BERT","date":"2024-08-16","arxiv_id":"2408.09043","n_code_links":0,"syntology":null},{"paper":null,"slug":"information-theoretic-progress-measures","title":"Information-Theoretic Progress Measures reveal Grokking is an Emergent Phase Transition","date":"2024-08-16","arxiv_id":"2408.08944","n_code_links":0,"syntology":null},{"paper":null,"slug":"meta-knowledge-for-retrieval-augmented-large","title":"Meta Knowledge for Retrieval Augmented Large Language Models","date":"2024-08-16","arxiv_id":"2408.09017","n_code_links":0,"syntology":null},{"paper":null,"slug":"quantifying-the-effectiveness-of-student","title":"Quantifying the Effectiveness of Student Organization Activities using Natural Language Processing","date":"2024-08-16","arxiv_id":"2408.08694","n_code_links":0,"syntology":null},{"paper":"/paper/the-fellowship-of-the-llms-multi-agent","slug":"the-fellowship-of-the-llms-multi-agent","title":"The Fellowship of the LLMs: Multi-Agent Workflows for Synthetic Preference Optimization Dataset Generation","date":"2024-08-16","arxiv_id":"2408.08688","n_code_links":1,"syntology":null},{"paper":null,"slug":"vera-validation-and-evaluation-of-retrieval","title":"VERA: Validation and Evaluation of Retrieval-Augmented Systems","date":"2024-08-16","arxiv_id":"2409.03759","n_code_links":0,"syntology":null},{"paper":null,"slug":"analytical-uncertainty-based-loss-weighting","title":"Analytical Uncertainty-Based Loss Weighting in Multi-Task Learning","date":"2024-08-15","arxiv_id":"2408.07985","n_code_links":0,"syntology":null},{"paper":"/paper/fusechat-knowledge-fusion-of-chat-models-1","slug":"fusechat-knowledge-fusion-of-chat-models-1","title":"FuseChat: Knowledge Fusion of Chat Models","date":"2024-08-15","arxiv_id":"2408.07990","n_code_links":3,"syntology":null},{"paper":"/paper/graph-retrieval-augmented-generation-a-survey","slug":"graph-retrieval-augmented-generation-a-survey","title":"Graph Retrieval-Augmented Generation: A Survey","date":"2024-08-15","arxiv_id":"2408.08921","n_code_links":1,"syntology":null},{"paper":"/paper/leveraging-web-crawled-data-for-high-quality","slug":"leveraging-web-crawled-data-for-high-quality","title":"Leveraging Web-Crawled Data for High-Quality Fine-Tuning","date":"2024-08-15","arxiv_id":"2408.08003","n_code_links":1,"syntology":null},{"paper":null,"slug":"plan-with-code-comparing-approaches-for","title":"Plan with Code: Comparing approaches for robust NL to DSL generation","date":"2024-08-15","arxiv_id":"2408.08335","n_code_links":0,"syntology":null},{"paper":null,"slug":"predicting-lung-cancer-patient-prognosis-with","title":"Predicting Lung Cancer Patient Prognosis with Large Language Models","date":"2024-08-15","arxiv_id":"2408.07971","n_code_links":0,"syntology":null},{"paper":"/paper/ragchecker-a-fine-grained-framework-for","slug":"ragchecker-a-fine-grained-framework-for","title":"RAGChecker: A Fine-grained Framework for Diagnosing Retrieval-Augmented Generation","date":"2024-08-15","arxiv_id":"2408.08067","n_code_links":1,"syntology":{"ran":3,"of":5,"n_ran_checked":3,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["amazon-science/ragchecker"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"codemirage-hallucinations-in-code-generated","title":"CodeMirage: Hallucinations in Code Generated by Large Language Models","date":"2024-08-14","arxiv_id":"2408.08333","n_code_links":0,"syntology":null},{"paper":"/paper/datavist5-a-pre-trained-language-model-for","slug":"datavist5-a-pre-trained-language-model-for","title":"DataVisT5: A Pre-trained Language Model for Jointly Understanding Text and Data Visualization","date":"2024-08-14","arxiv_id":"2408.07401","n_code_links":1,"syntology":null},{"paper":null,"slug":"enhancing-visual-question-answering-through-1","title":"Enhancing Visual Question Answering through Ranking-Based Hybrid Training and Multimodal Fusion","date":"2024-08-14","arxiv_id":"2408.07303","n_code_links":0,"syntology":null},{"paper":"/paper/exploring-retrieval-augmented-generation-in","slug":"exploring-retrieval-augmented-generation-in","title":"Exploring Retrieval Augmented Generation in Arabic","date":"2024-08-14","arxiv_id":"2408.07425","n_code_links":1,"syntology":null},{"paper":"/paper/lipcot-linear-predictive-coding-based","slug":"lipcot-linear-predictive-coding-based","title":"LiPCoT: Linear Predictive Coding based Tokenizer for Self-supervised Learning of Time Series Data via Language Models","date":"2024-08-14","arxiv_id":"2408.07292","n_code_links":1,"syntology":null},{"paper":null,"slug":"sage-rt-synthetic-alignment-data-generation","title":"SAGE-RT: Synthetic Alignment data Generation for Safety Evaluation and Red Teaming","date":"2024-08-14","arxiv_id":"2408.11851","n_code_links":0,"syntology":null},{"paper":null,"slug":"transformers-and-large-language-models-for-1","title":"Transformers and Large Language Models for Efficient Intrusion Detection Systems: A Comprehensive Survey","date":"2024-08-14","arxiv_id":"2408.07583","n_code_links":0,"syntology":null},{"paper":null,"slug":"bert-s-conceptual-cartography-mapping-the","title":"BERT's Conceptual Cartography: Mapping the Landscapes of Meaning","date":"2024-08-13","arxiv_id":"2408.07190","n_code_links":0,"syntology":null},{"paper":"/paper/evaluating-cultural-adaptability-of-a-large","slug":"evaluating-cultural-adaptability-of-a-large","title":"Evaluating Cultural Adaptability of a Large Language Model via Simulation of Synthetic Personas","date":"2024-08-13","arxiv_id":"2408.06929","n_code_links":1,"syntology":null},{"paper":null,"slug":"generative-ai-for-automatic-topic-labelling","title":"Generative AI for automatic topic labelling","date":"2024-08-13","arxiv_id":"2408.07003","n_code_links":0,"syntology":null},{"paper":null,"slug":"pragmatic-inference-of-scalar-implicature-by","title":"Pragmatic inference of scalar implicature by LLMs","date":"2024-08-13","arxiv_id":"2408.06673","n_code_links":0,"syntology":null},{"paper":null,"slug":"tableguard-securing-structured-unstructured","title":"TableGuard -- Securing Structured & Unstructured Data","date":"2024-08-13","arxiv_id":"2408.07045","n_code_links":0,"syntology":null},{"paper":null,"slug":"bayesian-inference-to-improve-quality-of","title":"Bayesian inference to improve quality of Retrieval Augmented Generation","date":"2024-08-12","arxiv_id":"2408.08901","n_code_links":0,"syntology":null},{"paper":null,"slug":"improving-structural-diversity-of-blackbox","title":"Improving Structural Diversity of Blackbox LLMs via Chain-of-Specification Prompting","date":"2024-08-12","arxiv_id":"2408.06186","n_code_links":0,"syntology":null},{"paper":null,"slug":"lolgorithm-integrating-semantic-syntactic-and","title":"LOLgorithm: Integrating Semantic,Syntactic and Contextual Elements for Humor Classification","date":"2024-08-12","arxiv_id":"2408.06335","n_code_links":0,"syntology":null},{"paper":null,"slug":"optimizing-rag-techniques-for-automotive","title":"Optimizing RAG Techniques for Automotive Industry PDF Chatbots: A Case Study with Locally Deployed Ollama Models","date":"2024-08-12","arxiv_id":"2408.05933","n_code_links":0,"syntology":null},{"paper":null,"slug":"the-language-of-trauma-modeling-traumatic","title":"The Language of Trauma: Modeling Traumatic Event Descriptions Across Domains with Explainable AI","date":"2024-08-12","arxiv_id":"2408.05977","n_code_links":0,"syntology":null}],"record_sha256":"188a82065c06ccdac29e54c09d1690e96b678d998b7a2d40bdd76b6760fcde21","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}