{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/linear-warmup-with-cosine-annealing/papers/21","list_of":"/method/linear-warmup-with-cosine-annealing","method":"Linear Warmup With Cosine Annealing","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":21,"pages_in_order":38,"rows_per_page":100,"rows":[2001,2100],"of":3797,"counts":{"archive_papers_tagged":3797,"with_a_code_link":1655,"where_syntology_ran_a_sample":602,"not_listed_spam_title":0,"listed":3797,"listed_where_code_ran":602,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":490,"every_run_a_failure_of_syntologys_instrument":112,"listed_with_a_run_with_no_instrument_failure":490,"listed_every_run_a_failure_of_syntologys_instrument":112,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/linear-warmup-with-cosine-annealing","prev":"/method/linear-warmup-with-cosine-annealing/papers/20","next":"/method/linear-warmup-with-cosine-annealing/papers/22","papers":[{"paper":null,"slug":"alltogether-investigating-the-efficacy-of","title":"AllTogether: Investigating the Efficacy of Spliced Prompt for Web Navigation using Large Language Models","date":"2023-10-20","arxiv_id":"2310.18331","n_code_links":0,"syntology":null},{"paper":"/paper/cache-me-if-you-can-an-online-cost-aware","slug":"cache-me-if-you-can-an-online-cost-aware","title":"Cache me if you Can: an Online Cost-aware Teacher-Student framework to Reduce the Calls to Large Language Models","date":"2023-10-20","arxiv_id":"2310.13395","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["stoyian/OCaTS"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"challenges-and-contributing-factors-in-the","title":"Challenges and Contributing Factors in the Utilization of Large Language Models (LLMs)","date":"2023-10-20","arxiv_id":"2310.13343","n_code_links":0,"syntology":null},{"paper":null,"slug":"design-inclusive-language-models-for","title":"She had Cobalt Blue Eyes: Prompt Testing to Create Aligned and Sustainable Language Models","date":"2023-10-20","arxiv_id":"2310.18333","n_code_links":0,"syntology":null},{"paper":null,"slug":"equivariant-transformer-is-all-you-need","title":"Equivariant Transformer is all you need","date":"2023-10-20","arxiv_id":"2310.13222","n_code_links":0,"syntology":null},{"paper":null,"slug":"foundation-model-s-embedded-representations","title":"Foundation Model's Embedded Representations May Detect Distribution Shift","date":"2023-10-20","arxiv_id":"2310.13836","n_code_links":0,"syntology":null},{"paper":null,"slug":"robust-training-for-conversational-question","title":"Robust Training for Conversational Question Answering Models with Reinforced Reformulation Generation","date":"2023-10-20","arxiv_id":"2310.13505","n_code_links":0,"syntology":null},{"paper":null,"slug":"the-perils-promises-of-fact-checking-with","title":"The Perils & Promises of Fact-checking with Large Language Models","date":"2023-10-20","arxiv_id":"2310.13549","n_code_links":0,"syntology":null},{"paper":null,"slug":"wordart-designer-user-driven-artistic","title":"WordArt Designer: User-Driven Artistic Typography Synthesis using Large Language Models","date":"2023-10-20","arxiv_id":"2310.18332","n_code_links":0,"syntology":null},{"paper":"/paper/agenttuning-enabling-generalized-agent","slug":"agenttuning-enabling-generalized-agent","title":"AgentTuning: Enabling Generalized Agent Abilities for LLMs","date":"2023-10-19","arxiv_id":"2310.12823","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["thudm/agenttuning"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"an-exploration-of-in-context-learning-for","title":"Exploring In-Context Learning of Textless Speech Language Model for Speech Classification Tasks","date":"2023-10-19","arxiv_id":"2310.12477","n_code_links":0,"syntology":null},{"paper":null,"slug":"experimental-narratives-a-comparison-of-human","title":"Experimental Narratives: A Comparison of Human Crowdsourced Storytelling and AI Storytelling","date":"2023-10-19","arxiv_id":"2310.12902","n_code_links":0,"syntology":null},{"paper":"/paper/identifying-and-adapting-transformer","slug":"identifying-and-adapting-transformer","title":"Identifying and Adapting Transformer-Components Responsible for Gender Bias in an English Language Model","date":"2023-10-19","arxiv_id":"2310.12611","n_code_links":1,"syntology":{"ran":4,"of":5,"n_ran_checked":4,"n_instrument":0,"unverified":1,"pointer_only":5,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["iabhijith/bias-causal-analysis"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"laser-linear-compression-in-wireless","title":"LASER: Linear Compression in Wireless Distributed Optimization","date":"2023-10-19","arxiv_id":"2310.13033","n_code_links":0,"syntology":null},{"paper":null,"slug":"not-all-countries-celebrate-thanksgiving-on","title":"Not All Countries Celebrate Thanksgiving: On the Cultural Dominance in Large Language Models","date":"2023-10-19","arxiv_id":"2310.12481","n_code_links":0,"syntology":null},{"paper":"/paper/product-attribute-value-extraction-using","slug":"product-attribute-value-extraction-using","title":"ExtractGPT: Exploring the Potential of Large Language Models for Product Attribute Value Extraction","date":"2023-10-19","arxiv_id":"2310.12537","n_code_links":1,"syntology":null},{"paper":"/paper/the-shifted-and-the-overlooked-a-task","slug":"the-shifted-and-the-overlooked-a-task","title":"The Shifted and The Overlooked: A Task-oriented Investigation of User-GPT Interactions","date":"2023-10-19","arxiv_id":"2310.12418","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":0,"n_instrument":2,"unverified":0,"pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":{"repos":["ozyyshr/sharegpt_investigation"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"evaluating-the-symbol-binding-ability-of","title":"Evaluating the Symbol Binding Ability of Large Language Models for Multiple-Choice Questions in Vietnamese General Education","date":"2023-10-18","arxiv_id":"2310.12059","n_code_links":0,"syntology":null},{"paper":null,"slug":"solving-the-multiplication-problem-of-a-large","title":"Solving the multiplication problem of a large language model system using a graph-based method","date":"2023-10-18","arxiv_id":"2310.13016","n_code_links":0,"syntology":null},{"paper":null,"slug":"emergent-ai-assisted-discourse-case-study-of","title":"Emergent AI-Assisted Discourse: Case Study of a Second Language Writer Authoring with ChatGPT","date":"2023-10-17","arxiv_id":"2310.10903","n_code_links":0,"syntology":null},{"paper":"/paper/evaluating-llms-for-privilege-escalation","slug":"evaluating-llms-for-privilege-escalation","title":"LLMs as Hackers: Autonomous Linux Privilege Escalation Attacks","date":"2023-10-17","arxiv_id":"2310.11409","n_code_links":1,"syntology":null},{"paper":"/paper/intent-detection-and-slot-filling-for-home","slug":"intent-detection-and-slot-filling-for-home","title":"Intent Detection and Slot Filling for Home Assistants: Dataset and Analysis for Bangla and Sylheti","date":"2023-10-17","arxiv_id":"2310.10935","n_code_links":1,"syntology":null},{"paper":"/paper/probing-the-creativity-of-large-language","slug":"probing-the-creativity-of-large-language","title":"Probing the Creativity of Large Language Models: Can models produce divergent semantic association?","date":"2023-10-17","arxiv_id":"2310.11158","n_code_links":1,"syntology":{"ran":1,"of":2,"n_ran_checked":0,"n_instrument":1,"unverified":1,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":{"repos":["dingnlab/probing_creativity"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"utilising-a-large-language-model-to-annotate","title":"Utilising a Large Language Model to Annotate Subject Metadata: A Case Study in an Australian National Research Data Catalogue","date":"2023-10-17","arxiv_id":"2310.11318","n_code_links":0,"syntology":null},{"paper":"/paper/approximating-two-layer-feedforward-networks","slug":"approximating-two-layer-feedforward-networks","title":"Approximating Two-Layer Feedforward Networks for Efficient Transformers","date":"2023-10-16","arxiv_id":"2310.10837","n_code_links":2,"syntology":{"ran":3,"of":4,"n_ran_checked":1,"n_instrument":2,"unverified":1,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 1 unverified","official":{"repos":["robertcsordas/moe","robertcsordas/moe_layer"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"battle-of-the-large-language-models-dolly-vs","title":"Battle of the Large Language Models: Dolly vs LLaMA vs Vicuna vs Guanaco vs Bard vs ChatGPT -- A Text-to-SQL Parsing Comparison","date":"2023-10-16","arxiv_id":"2310.10190","n_code_links":0,"syntology":null},{"paper":"/paper/bioplanner-automatic-evaluation-of-llms-on","slug":"bioplanner-automatic-evaluation-of-llms-on","title":"BioPlanner: Automatic Evaluation of LLMs on Protocol Planning in Biology","date":"2023-10-16","arxiv_id":"2310.10632","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":2,"n_instrument":0,"unverified":0,"pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["bioplanner/bioplanner"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/data-contamination-through-the-lens-of-time","slug":"data-contamination-through-the-lens-of-time","title":"Data Contamination Through the Lens of Time","date":"2023-10-16","arxiv_id":"2310.10628","n_code_links":1,"syntology":{"ran":2,"of":4,"n_ran_checked":2,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["abacusai/to-the-cutoff"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"fine-tuning-chatgpt-for-automatic-scoring","title":"Fine-tuning ChatGPT for Automatic Scoring","date":"2023-10-16","arxiv_id":"2310.10072","n_code_links":0,"syntology":null},{"paper":null,"slug":"moconvq-unified-physics-based-motion-control","title":"MoConVQ: Unified Physics-Based Motion Control via Scalable Discrete Representations","date":"2023-10-16","arxiv_id":"2310.10198","n_code_links":0,"syntology":null},{"paper":null,"slug":"prediction-of-arabic-legal-rulings-using","title":"Prediction of Arabic Legal Rulings using Large Language Models","date":"2023-10-16","arxiv_id":"2310.10260","n_code_links":0,"syntology":null},{"paper":"/paper/transom-an-efficient-fault-tolerant-system","slug":"transom-an-efficient-fault-tolerant-system","title":"TRANSOM: An Efficient Fault-Tolerant System for Training LLMs","date":"2023-10-16","arxiv_id":"2310.10046","n_code_links":1,"syntology":{"ran":0,"of":3,"n_ran_checked":0,"n_instrument":0,"unverified":3,"pointer_only":0,"phrase":"0 ran · 3 unverified","official":{"repos":["SenseCore/transom-checkpoint-engine"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":3,"ran_from_kinds":[]}}},{"paper":null,"slug":"configuration-validation-with-large-language","title":"Configuration Validation with Large Language Models","date":"2023-10-15","arxiv_id":"2310.09690","n_code_links":0,"syntology":null},{"paper":null,"slug":"image-augmentation-with-controlled-diffusion","title":"Image Augmentation with Controlled Diffusion for Weakly-Supervised Semantic Segmentation","date":"2023-10-15","arxiv_id":"2310.09760","n_code_links":0,"syntology":null},{"paper":null,"slug":"large-language-model-aware-in-context","title":"Large Language Model-Aware In-Context Learning for Code Generation","date":"2023-10-15","arxiv_id":"2310.09748","n_code_links":0,"syntology":null},{"paper":"/paper/large-language-models-for-in-context-student","slug":"large-language-models-for-in-context-student","title":"Large Language Models for In-Context Student Modeling: Synthesizing Student's Behavior in Visual Programming","date":"2023-10-15","arxiv_id":"2310.10690","n_code_links":1,"syntology":null},{"paper":null,"slug":"efficient-model-agnostic-multi-group","title":"Efficient Model-Agnostic Multi-Group Equivariant Networks","date":"2023-10-14","arxiv_id":"2310.09675","n_code_links":0,"syntology":null},{"paper":"/paper/a-systematic-evaluation-of-large-language-1","slug":"a-systematic-evaluation-of-large-language-1","title":"Assessing and Enhancing the Robustness of Large Language Models with Task Structure Variations for Logical Reasoning","date":"2023-10-13","arxiv_id":"2310.09430","n_code_links":1,"syntology":{"ran":6,"of":8,"n_ran_checked":6,"n_instrument":0,"unverified":2,"pointer_only":1,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["strong-ai-lab/logical-and-abstract-reasoning"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"from-words-and-exercises-to-wellness-farsi","title":"From Words and Exercises to Wellness: Farsi Chatbot for Self-Attachment Technique","date":"2023-10-13","arxiv_id":"2310.09362","n_code_links":0,"syntology":null},{"paper":"/paper/human-in-the-loop-machine-translation-with","slug":"human-in-the-loop-machine-translation-with","title":"Human-in-the-loop Machine Translation with Large Language Model","date":"2023-10-13","arxiv_id":"2310.08908","n_code_links":1,"syntology":null},{"paper":null,"slug":"table-gpt-table-tuned-gpt-for-diverse-table","title":"Table-GPT: Table-tuned GPT for Diverse Table Tasks","date":"2023-10-13","arxiv_id":"2310.09263","n_code_links":0,"syntology":null},{"paper":"/paper/towards-end-to-end-4-bit-inference-on","slug":"towards-end-to-end-4-bit-inference-on","title":"QUIK: Towards End-to-End 4-Bit Inference on Generative Large Language Models","date":"2023-10-13","arxiv_id":"2310.09259","n_code_links":1,"syntology":{"ran":9,"of":13,"n_ran_checked":9,"n_instrument":0,"unverified":4,"pointer_only":0,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 9 with no instrument failure: 0 honoured, 0 violated, 9 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","official":{"repos":["ist-daslab/quik"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":9,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":"/paper/jailbreaking-black-box-large-language-models","slug":"jailbreaking-black-box-large-language-models","title":"Jailbreaking Black Box Large Language Models in Twenty Queries","date":"2023-10-12","arxiv_id":"2310.08419","n_code_links":1,"syntology":{"ran":2,"of":7,"n_ran_checked":1,"n_instrument":1,"unverified":5,"pointer_only":0,"phrase":"2 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 1 where Syntology's instrument failed) · 5 unverified","official":{"repos":["patrickrchao/jailbreakingllms"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":5,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"large-language-models-can-replicate-cross","title":"Large language models can replicate cross-cultural differences in personality","date":"2023-10-12","arxiv_id":"2310.10679","n_code_links":0,"syntology":null},{"paper":null,"slug":"multiclass-classification-of-policy-documents","title":"Multiclass Classification of Policy Documents with Large Language Models","date":"2023-10-12","arxiv_id":"2310.08167","n_code_links":0,"syntology":null},{"paper":null,"slug":"promptor-a-conversational-and-autonomous","title":"Promptor: A Conversational and Autonomous Prompt Generation Agent for Intelligent Text Entry Techniques","date":"2023-10-12","arxiv_id":"2310.08101","n_code_links":0,"syntology":null},{"paper":"/paper/qasina-religious-domain-question-answering","slug":"qasina-religious-domain-question-answering","title":"QASiNa: Religious Domain Question Answering using Sirah Nabawiyah","date":"2023-10-12","arxiv_id":"2310.08102","n_code_links":1,"syntology":null},{"paper":null,"slug":"training-generative-question-answering-on","title":"Training Generative Question-Answering on Synthetic Data Obtained from an Instruct-tuned Model","date":"2023-10-12","arxiv_id":"2310.08072","n_code_links":0,"syntology":null},{"paper":null,"slug":"diversity-of-thought-improves-reasoning","title":"Diversity of Thought Improves Reasoning Abilities of LLMs","date":"2023-10-11","arxiv_id":"2310.07088","n_code_links":0,"syntology":null},{"paper":null,"slug":"exploring-the-landscape-of-large-language","title":"Do Large Language Models have Shared Weaknesses in Medical Question Answering?","date":"2023-10-11","arxiv_id":"2310.07225","n_code_links":0,"syntology":null},{"paper":"/paper/found-in-the-middle-permutation-self","slug":"found-in-the-middle-permutation-self","title":"Found in the Middle: Permutation Self-Consistency Improves Listwise Ranking in Large Language Models","date":"2023-10-11","arxiv_id":"2310.07712","n_code_links":1,"syntology":{"ran":5,"of":5,"n_ran_checked":5,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["castorini/perm-sc"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/instructretro-instruction-tuning-post","slug":"instructretro-instruction-tuning-post","title":"InstructRetro: Instruction Tuning post Retrieval-Augmented Pretraining","date":"2023-10-11","arxiv_id":"2310.07713","n_code_links":1,"syntology":null},{"paper":"/paper/large-language-models-are-zero-shot-time-1","slug":"large-language-models-are-zero-shot-time-1","title":"Large Language Models Are Zero-Shot Time Series Forecasters","date":"2023-10-11","arxiv_id":"2310.07820","n_code_links":2,"syntology":{"ran":7,"of":7,"n_ran_checked":7,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 1 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["ngruver/llmtime"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["listed","official"]}}},{"paper":"/paper/uncovering-hidden-connections-iterative","slug":"uncovering-hidden-connections-iterative","title":"Uncovering Hidden Connections: Iterative Search and Reasoning for Video-grounded Dialog","date":"2023-10-11","arxiv_id":"2310.07259","n_code_links":2,"syntology":{"ran":3,"of":6,"n_ran_checked":1,"n_instrument":2,"unverified":3,"pointer_only":6,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 2 where Syntology's instrument failed) · 3 unverified","official":{"repos":["hyu-zhang/itr","Hyu-Zhang/ISR"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"wigenai-the-symphony-of-wireless-and","title":"Diffusion Models for Wireless Communications","date":"2023-10-11","arxiv_id":"2310.07312","n_code_links":0,"syntology":null},{"paper":null,"slug":"automated-clinical-coding-using-off-the-shelf","title":"Automated clinical coding using off-the-shelf large language models","date":"2023-10-10","arxiv_id":"2310.06552","n_code_links":0,"syntology":null},{"paper":"/paper/geollm-extracting-geospatial-knowledge-from","slug":"geollm-extracting-geospatial-knowledge-from","title":"GeoLLM: Extracting Geospatial Knowledge from Large Language Models","date":"2023-10-10","arxiv_id":"2310.06213","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":2,"n_instrument":1,"unverified":0,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 2 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["rohinmanvi/GeoLLM"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"gpt-4-as-an-agronomist-assistant-answering","title":"GPT-4 as an Agronomist Assistant? Answering Agriculture Exams Using Large Language Models","date":"2023-10-10","arxiv_id":"2310.06225","n_code_links":0,"syntology":null},{"paper":"/paper/humans-and-language-models-diverge-when","slug":"humans-and-language-models-diverge-when","title":"Humans and language models diverge when predicting repeating text","date":"2023-10-10","arxiv_id":"2310.06408","n_code_links":1,"syntology":{"ran":1,"of":2,"n_ran_checked":0,"n_instrument":1,"unverified":1,"pointer_only":2,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":{"repos":["huthlab/lm-repeating-text"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/large-language-models-for-propaganda","slug":"large-language-models-for-propaganda","title":"Large Language Models for Propaganda Detection","date":"2023-10-10","arxiv_id":"2310.06422","n_code_links":2,"syntology":null},{"paper":"/paper/longllmlingua-accelerating-and-enhancing-llms","slug":"longllmlingua-accelerating-and-enhancing-llms","title":"LongLLMLingua: Accelerating and Enhancing LLMs in Long Context Scenarios via Prompt Compression","date":"2023-10-10","arxiv_id":"2310.06839","n_code_links":3,"syntology":null},{"paper":null,"slug":"cabbage-sweeter-than-cake-analysing-the","title":"Cabbage Sweeter than Cake? Analysing the Potential of Large Language Models for Learning Conceptual Spaces","date":"2023-10-09","arxiv_id":"2310.05481","n_code_links":0,"syntology":null},{"paper":null,"slug":"foundation-models-meet-visualizations","title":"Foundation Models Meet Visualizations: Challenges and Opportunities","date":"2023-10-09","arxiv_id":"2310.05771","n_code_links":0,"syntology":null},{"paper":"/paper/mbbc-exploring-the-multilingual-maze","slug":"mbbc-exploring-the-multilingual-maze","title":"Exploring the Maze of Multilingual Modeling","date":"2023-10-09","arxiv_id":"2310.05404","n_code_links":0,"syntology":null},{"paper":null,"slug":"sc-safety-a-multi-round-open-ended-question","title":"SC-Safety: A Multi-round Open-ended Question Adversarial Safety Benchmark for Large Language Models in Chinese","date":"2023-10-09","arxiv_id":"2310.05818","n_code_links":0,"syntology":null},{"paper":null,"slug":"the-program-testing-ability-of-large-language","title":"The Program Testing Ability of Large Language Models for Code","date":"2023-10-09","arxiv_id":"2310.05727","n_code_links":0,"syntology":null},{"paper":null,"slug":"are-emily-and-greg-still-more-employable-than","title":"Are Emily and Greg Still More Employable than Lakisha and Jamal? Investigating Algorithmic Hiring Bias in the Era of ChatGPT","date":"2023-10-08","arxiv_id":"2310.05135","n_code_links":0,"syntology":null},{"paper":"/paper/distantly-supervised-joint-entity-and","slug":"distantly-supervised-joint-entity-and","title":"Distantly-Supervised Joint Extraction with Noise-Robust Learning","date":"2023-10-08","arxiv_id":"2310.04994","n_code_links":1,"syntology":null},{"paper":"/paper/llm4vv-developing-llm-driven-testsuite-for","slug":"llm4vv-developing-llm-driven-testsuite-for","title":"LLM4VV: Developing LLM-Driven Testsuite for Compiler Validation","date":"2023-10-08","arxiv_id":"2310.04963","n_code_links":1,"syntology":null},{"paper":"/paper/zero-shot-detection-of-machine-generated","slug":"zero-shot-detection-of-machine-generated","title":"Zero-Shot Detection of Machine-Generated Codes","date":"2023-10-08","arxiv_id":"2310.05103","n_code_links":1,"syntology":null},{"paper":null,"slug":"do-self-supervised-speech-and-language-models","title":"Do self-supervised speech and language models extract similar representations as human brain?","date":"2023-10-07","arxiv_id":"2310.04645","n_code_links":0,"syntology":null},{"paper":"/paper/large-language-models-only-pass-primary","slug":"large-language-models-only-pass-primary","title":"Large Language Models Only Pass Primary School Exams in Indonesia: A Comprehensive Test on IndoMMLU","date":"2023-10-07","arxiv_id":"2310.04928","n_code_links":1,"syntology":{"ran":2,"of":5,"n_ran_checked":1,"n_instrument":1,"unverified":3,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 1 where Syntology's instrument failed) · 3 unverified","official":{"repos":["fajri91/indommlu"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":"/paper/lauragpt-listen-attend-understand-and","slug":"lauragpt-listen-attend-understand-and","title":"LauraGPT: Listen, Attend, Understand, and Regenerate Audio with GPT","date":"2023-10-07","arxiv_id":"2310.04673","n_code_links":2,"syntology":{"ran":3,"of":3,"n_ran_checked":3,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":null}},{"paper":null,"slug":"question-focused-summarization-by-decomposing","title":"Question-focused Summarization by Decomposing Articles into Facts and Opinions and Retrieving Entities","date":"2023-10-07","arxiv_id":"2310.04880","n_code_links":0,"syntology":null},{"paper":"/paper/copy-suppression-comprehensively","slug":"copy-suppression-comprehensively","title":"Copy Suppression: Comprehensively Understanding an Attention Head","date":"2023-10-06","arxiv_id":"2310.04625","n_code_links":1,"syntology":{"ran":3,"of":5,"n_ran_checked":3,"n_instrument":0,"unverified":2,"pointer_only":3,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["callummcdougall/seri-mats-2023-streamlit-pages"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"keyword-augmented-retrieval-novel-framework","title":"Keyword Augmented Retrieval: Novel framework for Information Retrieval integrated with speech interface","date":"2023-10-06","arxiv_id":"2310.04205","n_code_links":0,"syntology":null},{"paper":"/paper/language-agent-tree-search-unifies-reasoning","slug":"language-agent-tree-search-unifies-reasoning","title":"Language Agent Tree Search Unifies Reasoning Acting and Planning in Language Models","date":"2023-10-06","arxiv_id":"2310.04406","n_code_links":2,"syntology":{"ran":9,"of":9,"n_ran_checked":8,"n_instrument":1,"unverified":0,"pointer_only":2,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 1 honoured, 0 violated, 7 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["lapisrocks/languageagenttreesearch","andyz245/LanguageAgentTreeSearch"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/agent-instructs-large-language-models-to-be","slug":"agent-instructs-large-language-models-to-be","title":"Agent Instructs Large Language Models to be General Zero-Shot Reasoners","date":"2023-10-05","arxiv_id":"2310.03710","n_code_links":1,"syntology":{"ran":9,"of":10,"n_ran_checked":9,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 9 with no instrument failure: 0 honoured, 0 violated, 9 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["wang-research-lab/agentinstruct"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":9,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/automating-human-tutor-style-programming","slug":"automating-human-tutor-style-programming","title":"Automating Human Tutor-Style Programming Feedback: Leveraging GPT-4 Tutor Model for Hint Generation and GPT-3.5 Student Model for Hint Validation","date":"2023-10-05","arxiv_id":"2310.03780","n_code_links":2,"syntology":{"ran":10,"of":10,"n_ran_checked":10,"n_instrument":0,"unverified":0,"pointer_only":10,"phrase":"10 ran (of which 0 constructed an object rather than computing a result; 10 with no instrument failure: 0 honoured, 0 violated, 10 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["machine-teaching-group/lak2024_gpt4-hints-gpt3.5val","machine-teaching-group/lak2024_gpt4hints-gpt3.5val"],"state":"official (archive's flag): 10 ran","n_ran":10,"n_constructed":0,"n_ran_no_instrument_failure":10,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/dspy-compiling-declarative-language-model","slug":"dspy-compiling-declarative-language-model","title":"DSPy: Compiling Declarative Language Model Calls into Self-Improving Pipelines","date":"2023-10-05","arxiv_id":"2310.03714","n_code_links":3,"syntology":{"ran":3,"of":7,"n_ran_checked":3,"n_instrument":0,"unverified":4,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","official":{"repos":["stanfordnlp/dspy"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/fine-tuning-aligned-language-models","slug":"fine-tuning-aligned-language-models","title":"Fine-tuning Aligned Language Models Compromises Safety, Even When Users Do Not Intend To!","date":"2023-10-05","arxiv_id":"2310.03693","n_code_links":1,"syntology":{"ran":0,"of":1,"n_ran_checked":0,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"0 ran · 1 unverified","official":{"repos":["llm-tuning-safety/llms-finetuning-safety"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":[]}}},{"paper":"/paper/smoothllm-defending-large-language-models","slug":"smoothllm-defending-large-language-models","title":"SmoothLLM: Defending Large Language Models Against Jailbreaking Attacks","date":"2023-10-05","arxiv_id":"2310.03684","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-survey-of-gpt-3-family-large-language","title":"A Survey of GPT-3 Family Large Language Models Including ChatGPT and GPT-4","date":"2023-10-04","arxiv_id":"2310.12321","n_code_links":0,"syntology":null},{"paper":"/paper/large-language-model-cascades-with-mixture-of","slug":"large-language-model-cascades-with-mixture-of","title":"Large Language Model Cascades with Mixture of Thoughts Representations for Cost-efficient Reasoning","date":"2023-10-04","arxiv_id":"2310.03094","n_code_links":1,"syntology":null},{"paper":"/paper/memoria-hebbian-memory-architecture-for-human","slug":"memoria-hebbian-memory-architecture-for-human","title":"Memoria: Resolving Fateful Forgetting Problem through Human-Inspired Memory Architecture","date":"2023-10-04","arxiv_id":"2310.03052","n_code_links":1,"syntology":{"ran":1,"of":10,"n_ran_checked":0,"n_instrument":1,"unverified":9,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 9 unverified","official":{"repos":["cosmoquester/memoria"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":9,"ran_from_kinds":["official"]}}},{"paper":"/paper/nola-networks-as-linear-combination-of-low","slug":"nola-networks-as-linear-combination-of-low","title":"NOLA: Compressing LoRA using Linear Combination of Random Basis","date":"2023-10-04","arxiv_id":"2310.02556","n_code_links":1,"syntology":{"ran":5,"of":5,"n_ran_checked":2,"n_instrument":3,"unverified":0,"pointer_only":0,"phrase":"5 ran (of which 2 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 3 where Syntology's instrument failed) · 0 unverified","official":{"repos":["UCDvision/NOLA"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":2,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"retrieval-meets-long-context-large-language","title":"Retrieval meets Long Context Large Language Models","date":"2023-10-04","arxiv_id":"2310.03025","n_code_links":0,"syntology":null},{"paper":"/paper/instance-needs-more-care-rewriting-prompts","slug":"instance-needs-more-care-rewriting-prompts","title":"Instances Need More Care: Rewriting Prompts for Instances with LLMs in the Loop Yields Better Zero-Shot Performance","date":"2023-10-03","arxiv_id":"2310.02107","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":2,"n_instrument":0,"unverified":0,"pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["salokr/propmted"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/gpt-driver-learning-to-drive-with-gpt","slug":"gpt-driver-learning-to-drive-with-gpt","title":"GPT-Driver: Learning to Drive with GPT","date":"2023-10-02","arxiv_id":"2310.01415","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":1,"n_instrument":2,"unverified":0,"pointer_only":3,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 1 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":{"repos":["pointscoder/gpt-driver"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/llm-lies-hallucinations-are-not-bugs-but","slug":"llm-lies-hallucinations-are-not-bugs-but","title":"LLM Lies: Hallucinations are not Bugs, but Features as Adversarial Examples","date":"2023-10-02","arxiv_id":"2310.01469","n_code_links":1,"syntology":{"ran":1,"of":3,"n_ran_checked":1,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["pku-yuangroup/hallucination-attack"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"polysketchformer-fast-transformers-via","title":"PolySketchFormer: Fast Transformers via Sketching Polynomial Kernels","date":"2023-10-02","arxiv_id":"2310.01655","n_code_links":0,"syntology":null},{"paper":"/paper/analyzing-and-mitigating-object-hallucination","slug":"analyzing-and-mitigating-object-hallucination","title":"Analyzing and Mitigating Object Hallucination in Large Vision-Language Models","date":"2023-10-01","arxiv_id":"2310.00754","n_code_links":1,"syntology":{"ran":7,"of":8,"n_ran_checked":2,"n_instrument":5,"unverified":1,"pointer_only":8,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 1 violated, 1 with no contract checked; 5 where Syntology's instrument failed) · 1 unverified","official":{"repos":["yiyangzhou/lure"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/booookscore-a-systematic-exploration-of-book","slug":"booookscore-a-systematic-exploration-of-book","title":"BooookScore: A systematic exploration of book-length summarization in the era of LLMs","date":"2023-10-01","arxiv_id":"2310.00785","n_code_links":2,"syntology":{"ran":3,"of":3,"n_ran_checked":3,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["lilakk/booookscore"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/rolellm-benchmarking-eliciting-and-enhancing","slug":"rolellm-benchmarking-eliciting-and-enhancing","title":"RoleLLM: Benchmarking, Eliciting, and Enhancing Role-Playing Abilities of Large Language Models","date":"2023-10-01","arxiv_id":"2310.00746","n_code_links":2,"syntology":null},{"paper":null,"slug":"gaze-driven-sentence-simplification-for","title":"Gaze-Driven Sentence Simplification for Language Learners: Enhancing Comprehension and Readability","date":"2023-09-30","arxiv_id":"2310.00355","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-large-language-model-approach-to","title":"A Large Language Model Approach to Educational Survey Feedback Analysis","date":"2023-09-29","arxiv_id":"2309.17447","n_code_links":0,"syntology":null},{"paper":null,"slug":"an-evaluation-of-gpt-models-for-phenotype","title":"An evaluation of GPT models for phenotype concept recognition","date":"2023-09-29","arxiv_id":"2309.17169","n_code_links":0,"syntology":null},{"paper":"/paper/benchmarking-the-abilities-of-large-language","slug":"benchmarking-the-abilities-of-large-language","title":"Benchmarking the Abilities of Large Language Models for RDF Knowledge Graph Creation and Comprehension: How Well Do LLMs Speak Turtle?","date":"2023-09-29","arxiv_id":"2309.17122","n_code_links":3,"syntology":null},{"paper":"/paper/dyval-graph-informed-dynamic-evaluation-of","slug":"dyval-graph-informed-dynamic-evaluation-of","title":"DyVal: Dynamic Evaluation of Large Language Models for Reasoning Tasks","date":"2023-09-29","arxiv_id":"2309.17167","n_code_links":1,"syntology":null},{"paper":"/paper/llm-deliberation-evaluating-llms-with","slug":"llm-deliberation-evaluating-llms-with","title":"Cooperation, Competition, and Maliciousness: LLM-Stakeholders Interactive Negotiation","date":"2023-09-29","arxiv_id":"2309.17234","n_code_links":2,"syntology":{"ran":3,"of":3,"n_ran_checked":1,"n_instrument":2,"unverified":0,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":{"repos":["s-abdelnabi/llm-deliberation"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}}],"record_sha256":"f48ea89212d940a65c17dba82017f062270d5c2533e0142fcae207755392cafb","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}