{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/weight-decay/papers/28","list_of":"/method/weight-decay","method":"Weight Decay","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":28,"pages_in_order":108,"rows_per_page":100,"rows":[2701,2800],"of":10713,"counts":{"archive_papers_tagged":10713,"with_a_code_link":4533,"where_syntology_ran_a_sample":1291,"not_listed_spam_title":0,"listed":10713,"listed_where_code_ran":1291,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":1064,"every_run_a_failure_of_syntologys_instrument":227,"listed_with_a_run_with_no_instrument_failure":1064,"listed_every_run_a_failure_of_syntologys_instrument":227,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/weight-decay","prev":"/method/weight-decay/papers/27","next":"/method/weight-decay/papers/29","papers":[{"paper":null,"slug":"are-ppo-ed-language-models-hackable","title":"Are PPO-ed Language Models Hackable?","date":"2024-05-28","arxiv_id":"2406.02577","n_code_links":0,"syntology":null},{"paper":"/paper/atm-adversarial-tuning-multi-agent-system","slug":"atm-adversarial-tuning-multi-agent-system","title":"ATM: Adversarial Tuning Multi-agent System Makes a Robust Retrieval-Augmented Generator","date":"2024-05-28","arxiv_id":"2405.18111","n_code_links":1,"syntology":{"ran":6,"of":10,"n_ran_checked":6,"n_instrument":0,"unverified":4,"pointer_only":10,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","official":{"repos":["chuhac/atm-rag"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"attention-based-sequential-recommendation","title":"Attention-based sequential recommendation system using multimodal data","date":"2024-05-28","arxiv_id":"2405.17959","n_code_links":0,"syntology":null},{"paper":null,"slug":"don-t-forget-to-connect-improving-rag-with","title":"Don't Forget to Connect! Improving RAG with Graph-based Reranking","date":"2024-05-28","arxiv_id":"2405.18414","n_code_links":0,"syntology":null},{"paper":null,"slug":"edinburgh-clinical-nlp-at-mediqa-corr-2024","title":"Edinburgh Clinical NLP at MEDIQA-CORR 2024: Guiding Large Language Models with Hints","date":"2024-05-28","arxiv_id":"2405.18028","n_code_links":0,"syntology":null},{"paper":"/paper/llms-and-memorization-on-quality-and","slug":"llms-and-memorization-on-quality-and","title":"LLMs and Memorization: On Quality and Specificity of Copyright Compliance","date":"2024-05-28","arxiv_id":"2405.18492","n_code_links":1,"syntology":null},{"paper":null,"slug":"towards-a-sampling-theory-for-implicit-neural","title":"Towards a Sampling Theory for Implicit Neural Representations","date":"2024-05-28","arxiv_id":"2405.18410","n_code_links":0,"syntology":null},{"paper":null,"slug":"understanding-intrinsic-socioeconomic-biases","title":"Understanding Intrinsic Socioeconomic Biases in Large Language Models","date":"2024-05-28","arxiv_id":"2405.18662","n_code_links":0,"syntology":null},{"paper":null,"slug":"widin-wording-image-for-domain-invariant","title":"WIDIn: Wording Image for Domain-Invariant Representation in Single-Source Domain Generalization","date":"2024-05-28","arxiv_id":"2405.18405","n_code_links":0,"syntology":null},{"paper":"/paper/assessing-llms-suitability-for-knowledge","slug":"assessing-llms-suitability-for-knowledge","title":"Assessing LLMs Suitability for Knowledge Graph Completion","date":"2024-05-27","arxiv_id":"2405.17249","n_code_links":1,"syntology":null},{"paper":null,"slug":"augmenting-textual-generation-via-topology","title":"Augmenting Textual Generation via Topology Aware Retrieval","date":"2024-05-27","arxiv_id":"2405.17602","n_code_links":0,"syntology":null},{"paper":"/paper/deeperimpact-optimizing-sparse-learned-index","slug":"deeperimpact-optimizing-sparse-learned-index","title":"DeeperImpact: Optimizing Sparse Learned Index Structures","date":"2024-05-27","arxiv_id":"2405.17093","n_code_links":1,"syntology":null},{"paper":null,"slug":"detecting-deceptive-dark-patterns-in-e","title":"Detecting Deceptive Dark Patterns in E-commerce Platforms","date":"2024-05-27","arxiv_id":"2406.01608","n_code_links":0,"syntology":null},{"paper":null,"slug":"exploiting-the-layered-intrinsic","title":"Exploiting the Layered Intrinsic Dimensionality of Deep Models for Practical Adversarial Training","date":"2024-05-27","arxiv_id":"2405.17130","n_code_links":0,"syntology":null},{"paper":"/paper/inversionview-a-general-purpose-method-for","slug":"inversionview-a-general-purpose-method-for","title":"InversionView: A General-Purpose Method for Reading Information from Neural Activations","date":"2024-05-27","arxiv_id":"2405.17653","n_code_links":1,"syntology":{"ran":1,"of":6,"n_ran_checked":1,"n_instrument":0,"unverified":5,"pointer_only":6,"phrase":"1 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 5 unverified; the one sample that ran constructed an object rather than computing a result","official":{"repos":["huangxt39/inversionview"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":5,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"llm-based-cooperative-agents-using","title":"REVECA: Adaptive Planning and Trajectory-based Validation in Cooperative Language Agents using Information Relevance and Relative Proximity","date":"2024-05-27","arxiv_id":"2405.16751","n_code_links":0,"syntology":null},{"paper":null,"slug":"nv-embed-improved-techniques-for-training","title":"NV-Embed: Improved Techniques for Training LLMs as Generalist Embedding Models","date":"2024-05-27","arxiv_id":"2405.17428","n_code_links":0,"syntology":null},{"paper":null,"slug":"pae-llm-based-product-attribute-extraction","title":"PAE: LLM-based Product Attribute Extraction for E-Commerce Fashion Trends","date":"2024-05-27","arxiv_id":"2405.17533","n_code_links":0,"syntology":null},{"paper":null,"slug":"performance-evaluation-of-reddit-comments","title":"Performance evaluation of Reddit Comments using Machine Learning and Natural Language Processing methods in Sentiment Analysis","date":"2024-05-27","arxiv_id":"2405.16810","n_code_links":0,"syntology":null},{"paper":null,"slug":"qub-cirdan-at-discharge-me-zero-shot","title":"QUB-Cirdan at \"Discharge Me!\": Zero shot discharge letter generation by open-source LLM","date":"2024-05-27","arxiv_id":"2406.00041","n_code_links":0,"syntology":null},{"paper":"/paper/reflectioncoder-learning-from-reflection","slug":"reflectioncoder-learning-from-reflection","title":"ReflectionCoder: Learning from Reflection Sequence for Enhanced One-off Code Generation","date":"2024-05-27","arxiv_id":"2405.17057","n_code_links":1,"syntology":{"ran":0,"of":1,"n_ran_checked":0,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"0 ran · 1 unverified","official":{"repos":["sensellm/reflectioncoder"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":[]}}},{"paper":"/paper/rtl-repo-a-benchmark-for-evaluating-llms-on","slug":"rtl-repo-a-benchmark-for-evaluating-llms-on","title":"RTL-Repo: A Benchmark for Evaluating LLMs on Large-Scale RTL Design Projects","date":"2024-05-27","arxiv_id":"2405.17378","n_code_links":1,"syntology":null},{"paper":null,"slug":"the-scaling-law-in-stellar-light-curves","title":"The Scaling Law in Stellar Light Curves","date":"2024-05-27","arxiv_id":"2405.17156","n_code_links":0,"syntology":null},{"paper":"/paper/thread-thinking-deeper-with-recursive","slug":"thread-thinking-deeper-with-recursive","title":"THREAD: Thinking Deeper with Recursive Spawning","date":"2024-05-27","arxiv_id":"2405.17402","n_code_links":1,"syntology":{"ran":7,"of":11,"n_ran_checked":6,"n_instrument":1,"unverified":4,"pointer_only":11,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 1 where Syntology's instrument failed) · 4 unverified","official":{"repos":["philipmit/thread"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":"/paper/video-enriched-retrieval-augmented-generation","slug":"video-enriched-retrieval-augmented-generation","title":"Video Enriched Retrieval Augmented Generation Using Aligned Video Captions","date":"2024-05-27","arxiv_id":"2405.17706","n_code_links":1,"syntology":null},{"paper":null,"slug":"ai-generated-text-detection-and","title":"AI-Generated Text Detection and Classification Based on BERT Deep Learning Algorithm","date":"2024-05-26","arxiv_id":"2405.16422","n_code_links":0,"syntology":null},{"paper":"/paper/grag-graph-retrieval-augmented-generation","slug":"grag-graph-retrieval-augmented-generation","title":"GRAG: Graph Retrieval-Augmented Generation","date":"2024-05-26","arxiv_id":"2405.16506","n_code_links":1,"syntology":null},{"paper":null,"slug":"m-rag-reinforcing-large-language-model","title":"M-RAG: Reinforcing Large Language Model Performance through Retrieval-Augmented Generation with Multiple Partitions","date":"2024-05-26","arxiv_id":"2405.16420","n_code_links":0,"syntology":null},{"paper":null,"slug":"accelerating-inference-of-retrieval-augmented","title":"Accelerating Inference of Retrieval-Augmented Generation via Sparse Context Selection","date":"2024-05-25","arxiv_id":"2405.16178","n_code_links":0,"syntology":null},{"paper":"/paper/accelerating-transformers-with-spectrum-1","slug":"accelerating-transformers-with-spectrum-1","title":"Accelerating Transformers with Spectrum-Preserving Token Merging","date":"2024-05-25","arxiv_id":"2405.16148","n_code_links":1,"syntology":{"ran":6,"of":7,"n_ran_checked":2,"n_instrument":4,"unverified":1,"pointer_only":7,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 2 honoured, 0 violated, 0 with no contract checked; 4 where Syntology's instrument failed) · 1 unverified","official":{"repos":["hchautran/PiToMe"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/automanual-generating-instruction-manuals-by","slug":"automanual-generating-instruction-manuals-by","title":"AutoManual: Constructing Instruction Manuals by LLM Agents via Interactive Environmental Learning","date":"2024-05-25","arxiv_id":"2405.16247","n_code_links":1,"syntology":{"ran":4,"of":4,"n_ran_checked":4,"n_instrument":0,"unverified":0,"pointer_only":4,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["minghchen/automanual"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"incremental-comprehension-of-garden-path","title":"Incremental Comprehension of Garden-Path Sentences by Large Language Models: Semantic Interpretation, Syntactic Re-Analysis, and Attention","date":"2024-05-25","arxiv_id":"2405.16042","n_code_links":0,"syntology":null},{"paper":null,"slug":"mindstar-enhancing-math-reasoning-in-pre","title":"MindStar: Enhancing Math Reasoning in Pre-trained LLMs at Inference Time","date":"2024-05-25","arxiv_id":"2405.16265","n_code_links":0,"syntology":null},{"paper":null,"slug":"towards-unlocking-insights-from-logbooks","title":"Towards Unlocking Insights from Logbooks Using AI","date":"2024-05-25","arxiv_id":"2406.12881","n_code_links":0,"syntology":null},{"paper":null,"slug":"an-evaluation-of-estimative-uncertainty-in","title":"An Evaluation of Estimative Uncertainty in Large Language Models","date":"2024-05-24","arxiv_id":"2405.15185","n_code_links":0,"syntology":null},{"paper":null,"slug":"benchmarking-pre-trained-large-language","title":"Benchmarking the Performance of Pre-trained LLMs across Urdu NLP Tasks","date":"2024-05-24","arxiv_id":"2405.15453","n_code_links":0,"syntology":null},{"paper":"/paper/culturepark-boosting-cross-cultural","slug":"culturepark-boosting-cross-cultural","title":"CulturePark: Boosting Cross-cultural Understanding in Large Language Models","date":"2024-05-24","arxiv_id":"2405.15145","n_code_links":1,"syntology":{"ran":0,"of":7,"n_ran_checked":0,"n_instrument":0,"unverified":7,"pointer_only":7,"phrase":"0 ran · 7 unverified","official":{"repos":["scarelette/culturepark"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":7,"ran_from_kinds":[]}}},{"paper":null,"slug":"enhancing-augmentative-and-alternative","title":"Enhancing Augmentative and Alternative Communication with Card Prediction and Colourful Semantics","date":"2024-05-24","arxiv_id":"2405.15896","n_code_links":0,"syntology":null},{"paper":"/paper/evaluating-the-adversarial-robustness-of-1","slug":"evaluating-the-adversarial-robustness-of-1","title":"Evaluating and Safeguarding the Adversarial Robustness of Retrieval-Based In-Context Learning","date":"2024-05-24","arxiv_id":"2405.15984","n_code_links":1,"syntology":{"ran":5,"of":6,"n_ran_checked":5,"n_instrument":0,"unverified":1,"pointer_only":6,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["simonucl/adv-retreival-icl"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/generalizable-and-scalable-multistage","slug":"generalizable-and-scalable-multistage","title":"Generalizable and Scalable Multistage Biomedical Concept Normalization Leveraging Large Language Models","date":"2024-05-24","arxiv_id":"2405.15122","n_code_links":1,"syntology":null},{"paper":null,"slug":"gpt-is-not-an-annotator-the-necessity-of","title":"GPT is Not an Annotator: The Necessity of Human Annotation in Fairness Benchmark Construction","date":"2024-05-24","arxiv_id":"2405.15760","n_code_links":0,"syntology":null},{"paper":"/paper/learning-the-language-of-protein-structure","slug":"learning-the-language-of-protein-structure","title":"Learning the Language of Protein Structure","date":"2024-05-24","arxiv_id":"2405.15840","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["instadeepai/protein-structure-tokenizer"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"textit-comet-a-underline-com-munication","title":"Comet: A Communication-efficient and Performant Approximation for Private Transformer Inference","date":"2024-05-24","arxiv_id":"2405.17485","n_code_links":0,"syntology":null},{"paper":null,"slug":"the-impact-and-opportunities-of-generative-ai","title":"The Impact and Opportunities of Generative AI in Fact-Checking","date":"2024-05-24","arxiv_id":"2405.15985","n_code_links":0,"syntology":null},{"paper":null,"slug":"towards-understanding-how-transformer-perform","title":"The Buffer Mechanism for Multi-Step Information Reasoning in Language Models","date":"2024-05-24","arxiv_id":"2405.15302","n_code_links":0,"syntology":null},{"paper":"/paper/a-structure-aware-framework-for-learning","slug":"a-structure-aware-framework-for-learning","title":"A Structure-Aware Framework for Learning Device Placements on Computation Graphs","date":"2024-05-23","arxiv_id":"2405.14185","n_code_links":1,"syntology":null},{"paper":"/paper/ceebert-cross-domain-inference-in-early-exit","slug":"ceebert-cross-domain-inference-in-early-exit","title":"CEEBERT: Cross-Domain Inference in Early Exit BERT","date":"2024-05-23","arxiv_id":"2405.15039","n_code_links":1,"syntology":{"ran":2,"of":3,"n_ran_checked":2,"n_instrument":0,"unverified":1,"pointer_only":3,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["Div290/CeeBERT"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/editworld-simulating-world-dynamics-for","slug":"editworld-simulating-world-dynamics-for","title":"EditWorld: Simulating World Dynamics for Instruction-Following Image Editing","date":"2024-05-23","arxiv_id":"2405.14785","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 1 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["yangling0818/editworld"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/eliciting-informative-text-evaluations-with","slug":"eliciting-informative-text-evaluations-with","title":"Eliciting Informative Text Evaluations with Large Language Models","date":"2024-05-23","arxiv_id":"2405.15077","n_code_links":1,"syntology":{"ran":10,"of":14,"n_ran_checked":8,"n_instrument":2,"unverified":4,"pointer_only":14,"phrase":"10 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 0 violated, 8 with no contract checked; 2 where Syntology's instrument failed) · 4 unverified","official":{"repos":["yx-lu/eliciting-informative-text-evaluations-with-large-language-models"],"state":"official (archive's flag): 10 ran","n_ran":10,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":"/paper/from-explicit-cot-to-implicit-cot-learning-to","slug":"from-explicit-cot-to-implicit-cot-learning-to","title":"From Explicit CoT to Implicit CoT: Learning to Internalize CoT Step by Step","date":"2024-05-23","arxiv_id":"2405.14838","n_code_links":1,"syntology":{"ran":4,"of":4,"n_ran_checked":4,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["da03/internalize_cot_step_by_step"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/hipporag-neurobiologically-inspired-long-term","slug":"hipporag-neurobiologically-inspired-long-term","title":"HippoRAG: Neurobiologically Inspired Long-Term Memory for Large Language Models","date":"2024-05-23","arxiv_id":"2405.14831","n_code_links":2,"syntology":{"ran":5,"of":5,"n_ran_checked":5,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["osu-nlp-group/hipporag"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":null,"slug":"large-language-models-can-self-correct-with","title":"Large Language Models Can Self-Correct with Key Condition Verification","date":"2024-05-23","arxiv_id":"2405.14092","n_code_links":0,"syntology":null},{"paper":"/paper/not-all-language-model-features-are-linear","slug":"not-all-language-model-features-are-linear","title":"Not All Language Model Features Are Linear","date":"2024-05-23","arxiv_id":"2405.14860","n_code_links":1,"syntology":{"ran":6,"of":10,"n_ran_checked":6,"n_instrument":0,"unverified":4,"pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","official":{"repos":["joshengels/multidimensionalfeatures"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":"/paper/phinets-brain-inspired-non-contrastive","slug":"phinets-brain-inspired-non-contrastive","title":"PhiNets: Brain-inspired Non-contrastive Learning Based on Temporal Prediction Hypothesis","date":"2024-05-23","arxiv_id":"2405.14650","n_code_links":0,"syntology":{"ran":4,"of":5,"n_ran_checked":3,"n_instrument":1,"unverified":1,"pointer_only":5,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":null}},{"paper":null,"slug":"rafe-ranking-feedback-improves-query","title":"RaFe: Ranking Feedback Improves Query Rewriting for RAG","date":"2024-05-23","arxiv_id":"2405.14431","n_code_links":0,"syntology":null},{"paper":"/paper/vihatet5-enhancing-hate-speech-detection-in","slug":"vihatet5-enhancing-hate-speech-detection-in","title":"ViHateT5: Enhancing Hate Speech Detection in Vietnamese With A Unified Text-to-Text Transformer Model","date":"2024-05-23","arxiv_id":"2405.14141","n_code_links":1,"syntology":null},{"paper":"/paper/wise-rethinking-the-knowledge-memory-for","slug":"wise-rethinking-the-knowledge-memory-for","title":"WISE: Rethinking the Knowledge Memory for Lifelong Model Editing of Large Language Models","date":"2024-05-23","arxiv_id":"2405.14768","n_code_links":1,"syntology":{"ran":13,"of":16,"n_ran_checked":4,"n_instrument":9,"unverified":3,"pointer_only":0,"phrase":"13 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 1 honoured, 0 violated, 3 with no contract checked; 9 where Syntology's instrument failed) · 3 unverified","official":{"repos":["zjunlp/easyedit"],"state":"official (archive's flag): 13 ran","n_ran":13,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":"/paper/automated-evaluation-of-retrieval-augmented","slug":"automated-evaluation-of-retrieval-augmented","title":"Automated Evaluation of Retrieval-Augmented Language Models with Task-Specific Exam Generation","date":"2024-05-22","arxiv_id":"2405.13622","n_code_links":1,"syntology":null},{"paper":null,"slug":"automatically-identifying-local-and-global","title":"Automatically Identifying Local and Global Circuits with Linear Computation Graphs","date":"2024-05-22","arxiv_id":"2405.13868","n_code_links":0,"syntology":null},{"paper":"/paper/evaluating-large-language-models-with-human","slug":"evaluating-large-language-models-with-human","title":"Evaluating Large Language Models with Human Feedback: Establishing a Swedish Benchmark","date":"2024-05-22","arxiv_id":"2405.14006","n_code_links":1,"syntology":null},{"paper":"/paper/flashrag-a-modular-toolkit-for-efficient","slug":"flashrag-a-modular-toolkit-for-efficient","title":"FlashRAG: A Modular Toolkit for Efficient Retrieval-Augmented Generation Research","date":"2024-05-22","arxiv_id":"2405.13576","n_code_links":1,"syntology":null},{"paper":null,"slug":"how-to-set-adamw-s-weight-decay-as-you-scale","title":"How to set AdamW's weight decay as you scale model and dataset size","date":"2024-05-22","arxiv_id":"2405.13698","n_code_links":0,"syntology":null},{"paper":null,"slug":"ku-dmis-at-ehrsql-2024-generating-sql-query","title":"KU-DMIS at EHRSQL 2024:Generating SQL query via question templatization in EHR","date":"2024-05-22","arxiv_id":"2406.00014","n_code_links":0,"syntology":null},{"paper":"/paper/topa-extend-large-language-models-for-video","slug":"topa-extend-large-language-models-for-video","title":"TOPA: Extending Large Language Models for Video Understanding via Text-Only Pre-Alignment","date":"2024-05-22","arxiv_id":"2405.13911","n_code_links":1,"syntology":{"ran":6,"of":9,"n_ran_checked":4,"n_instrument":2,"unverified":3,"pointer_only":3,"phrase":"6 ran (of which 3 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 1 violated, 3 with no contract checked; 2 where Syntology's instrument failed) · 3 unverified","official":{"repos":["dhg-wei/topa"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":3,"n_ran_no_instrument_failure":3,"n_unverified":3,"ran_from_kinds":["official","unlocated"]}}},{"paper":"/paper/trojanrag-retrieval-augmented-generation-can","slug":"trojanrag-retrieval-augmented-generation-can","title":"TrojanRAG: Retrieval-Augmented Generation Can Be Backdoor Driver in Large Language Models","date":"2024-05-22","arxiv_id":"2405.13401","n_code_links":1,"syntology":null},{"paper":null,"slug":"unleashing-the-power-of-unlabeled-data-a-self","title":"Unleashing the Power of Unlabeled Data: A Self-supervised Learning Framework for Cyber Attack Detection in Smart Grids","date":"2024-05-22","arxiv_id":"2405.13965","n_code_links":0,"syntology":null},{"paper":"/paper/fadam-adam-is-a-natural-gradient-optimizer","slug":"fadam-adam-is-a-natural-gradient-optimizer","title":"FAdam: Adam is a natural gradient optimizer using diagonal empirical Fisher information","date":"2024-05-21","arxiv_id":"2405.12807","n_code_links":1,"syntology":null},{"paper":null,"slug":"generative-ai-and-large-language-models-for","title":"Generative AI in Cybersecurity: A Comprehensive Review of LLM Applications and Vulnerabilities","date":"2024-05-21","arxiv_id":"2405.12750","n_code_links":0,"syntology":null},{"paper":null,"slug":"how-reliable-ai-chatbots-are-for-disease","title":"How Reliable AI Chatbots are for Disease Prediction from Patient Complaints?","date":"2024-05-21","arxiv_id":"2405.13219","n_code_links":0,"syntology":null},{"paper":null,"slug":"investigating-persuasion-techniques-in-arabic","title":"Investigating Persuasion Techniques in Arabic: An Empirical Study Leveraging Large Language Models","date":"2024-05-21","arxiv_id":"2405.12884","n_code_links":0,"syntology":null},{"paper":"/paper/quantifying-emergence-in-large-language","slug":"quantifying-emergence-in-large-language","title":"Quantifying Semantic Emergence in Language Models","date":"2024-05-21","arxiv_id":"2405.12617","n_code_links":1,"syntology":{"ran":1,"of":6,"n_ran_checked":1,"n_instrument":0,"unverified":5,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 5 unverified","official":{"repos":["zodiark-ch/emergence-of-llms"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":5,"ran_from_kinds":["official"]}}},{"paper":"/paper/the-2nd-futuredial-challenge-dialog-systems","slug":"the-2nd-futuredial-challenge-dialog-systems","title":"The 2nd FutureDial Challenge: Dialog Systems with Retrieval Augmented Generation (FutureDial-RAG)","date":"2024-05-21","arxiv_id":"2405.13084","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-review-on-the-use-of-large-language-models","title":"A review on the use of large language models as virtual tutors","date":"2024-05-20","arxiv_id":"2405.11983","n_code_links":0,"syntology":null},{"paper":null,"slug":"crema-crisis-response-through-computational","title":"CReMa: Crisis Response through Computational Identification and Matching of Cross-Lingual Requests and Offers Shared on Social Media","date":"2024-05-20","arxiv_id":"2405.11897","n_code_links":0,"syntology":null},{"paper":null,"slug":"degree-of-irrationality-sentiment-and-implied","title":"Degree of Irrationality: Sentiment and Implied Volatility Surface","date":"2024-05-20","arxiv_id":"2405.11730","n_code_links":0,"syntology":null},{"paper":"/paper/evaluating-and-modeling-social-intelligence-a","slug":"evaluating-and-modeling-social-intelligence-a","title":"Evaluating and Modeling Social Intelligence: A Comparative Study of Human and AI Capabilities","date":"2024-05-20","arxiv_id":"2405.11841","n_code_links":1,"syntology":null},{"paper":null,"slug":"question-based-retrieval-using-atomic-units","title":"Question-Based Retrieval using Atomic Units for Enterprise RAG","date":"2024-05-20","arxiv_id":"2405.12363","n_code_links":0,"syntology":null},{"paper":null,"slug":"davinci-at-semeval-2024-task-9-few-shot","title":"DaVinci at SemEval-2024 Task 9: Few-shot prompting GPT-3.5 for Unconventional Reasoning","date":"2024-05-19","arxiv_id":"2405.11559","n_code_links":0,"syntology":null},{"paper":"/paper/human-centered-llm-agent-user-interface-a","slug":"human-centered-llm-agent-user-interface-a","title":"Human-Centered LLM-Agent User Interface: A Position Paper","date":"2024-05-19","arxiv_id":"2405.13050","n_code_links":1,"syntology":null},{"paper":"/paper/your-transformer-is-secretly-linear","slug":"your-transformer-is-secretly-linear","title":"Your Transformer is Secretly Linear","date":"2024-05-19","arxiv_id":"2405.12250","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":3,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["AIRI-Institute/LLM-Microscope"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/zero-shot-stance-detection-using-contextual","slug":"zero-shot-stance-detection-using-contextual","title":"Zero-Shot Stance Detection using Contextual Data Generation with LLMs","date":"2024-05-19","arxiv_id":"2405.11637","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":2,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["Babakbehkamkia/GPT-Stance-Detection"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"cross-language-assessment-of-mathematical","title":"Cross-Language Assessment of Mathematical Capability of ChatGPT","date":"2024-05-18","arxiv_id":"2405.11264","n_code_links":0,"syntology":null},{"paper":null,"slug":"exploring-speech-style-spaces-with-language","title":"Exploring speech style spaces with language models: Emotional TTS without emotion labels","date":"2024-05-18","arxiv_id":"2405.11413","n_code_links":0,"syntology":null},{"paper":null,"slug":"meta-reinforcement-learning-for-resource-1","title":"Meta Reinforcement Learning for Resource Allocation in Multi-Antenna UAV Network with Rate Splitting Multiple Access","date":"2024-05-18","arxiv_id":"2405.11306","n_code_links":0,"syntology":null},{"paper":null,"slug":"activellm-large-language-model-based-active","title":"ActiveLLM: Large Language Model-based Active Learning for Textual Few-Shot Scenarios","date":"2024-05-17","arxiv_id":"2405.10808","n_code_links":0,"syntology":null},{"paper":null,"slug":"empowering-prior-to-court-legal-analysis-a","title":"Empowering Prior to Court Legal Analysis: A Transparent and Accessible Dataset for Defensive Statement Classification and Interpretation","date":"2024-05-17","arxiv_id":"2405.10702","n_code_links":0,"syntology":null},{"paper":"/paper/evaluation-of-large-language-model","slug":"evaluation-of-large-language-model","title":"Evaluation of large language model performance on the Biomedical Language Understanding and Reasoning Benchmark","date":"2024-05-17","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/language-models-can-exploit-cross-task-in","slug":"language-models-can-exploit-cross-task-in","title":"Language Models can Exploit Cross-Task In-context Learning for Data-Scarce Novel Tasks","date":"2024-05-17","arxiv_id":"2405.10548","n_code_links":1,"syntology":{"ran":7,"of":9,"n_ran_checked":7,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["c-anwoy/cross-task-icl"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":7,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"neuroassist-enhancing-cognitive-computer","title":"NeuroAssist: Enhancing Cognitive-Computer Synergy with Adaptive AI and Advanced Neural Decoding for Efficient EEG Signal Classification","date":"2024-05-17","arxiv_id":"2406.01600","n_code_links":0,"syntology":null},{"paper":"/paper/tailoring-vaccine-messaging-with-common","slug":"tailoring-vaccine-messaging-with-common","title":"Tailoring Vaccine Messaging with Common-Ground Opinions","date":"2024-05-17","arxiv_id":"2405.10861","n_code_links":1,"syntology":null},{"paper":null,"slug":"fintextqa-a-dataset-for-long-form-financial","title":"FinTextQA: A Dataset for Long-form Financial Question Answering","date":"2024-05-16","arxiv_id":"2405.09980","n_code_links":0,"syntology":null},{"paper":null,"slug":"gpt-store-mining-and-analysis","title":"GPT Store Mining and Analysis","date":"2024-05-16","arxiv_id":"2405.10210","n_code_links":0,"syntology":null},{"paper":"/paper/hw-gpt-bench-hardware-aware-architecture","slug":"hw-gpt-bench-hardware-aware-architecture","title":"HW-GPT-Bench: Hardware-Aware Architecture Benchmark for Language Models","date":"2024-05-16","arxiv_id":"2405.10299","n_code_links":2,"syntology":null},{"paper":null,"slug":"optimization-techniques-for-sentiment","title":"Optimization Techniques for Sentiment Analysis Based on LLM (GPT-3)","date":"2024-05-16","arxiv_id":"2405.09770","n_code_links":0,"syntology":null},{"paper":null,"slug":"bridging-the-gap-in-online-hate-speech","title":"Bridging the gap in online hate speech detection: a comparative analysis of BERT and traditional models for homophobic content identification on X/Twitter","date":"2024-05-15","arxiv_id":"2405.09221","n_code_links":0,"syntology":null},{"paper":null,"slug":"im-rag-multi-round-retrieval-augmented","title":"IM-RAG: Multi-Round Retrieval-Augmented Generation Through Learning Inner Monologues","date":"2024-05-15","arxiv_id":"2405.13021","n_code_links":0,"syntology":null},{"paper":"/paper/lora-learns-less-and-forgets-less","slug":"lora-learns-less-and-forgets-less","title":"LoRA Learns Less and Forgets Less","date":"2024-05-15","arxiv_id":"2405.09673","n_code_links":1,"syntology":null},{"paper":null,"slug":"matching-domain-experts-by-training-from","title":"Matching domain experts by training from scratch on domain knowledge","date":"2024-05-15","arxiv_id":"2405.09395","n_code_links":0,"syntology":null},{"paper":null,"slug":"transfer-learning-in-pre-trained-large","title":"Transfer Learning in Pre-Trained Large Language Models for Malware Detection Based on System Calls","date":"2024-05-15","arxiv_id":"2405.09318","n_code_links":0,"syntology":null},{"paper":null,"slug":"beyond-scaling-laws-understanding-transformer","title":"Beyond Scaling Laws: Understanding Transformer Performance with Associative Memory","date":"2024-05-14","arxiv_id":"2405.08707","n_code_links":0,"syntology":null}],"record_sha256":"c17b3971ef1380e20a6a785c1945421a18efa76c5728a5c4656a876697383629","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}