{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/linear-warmup-with-cosine-annealing/papers/19","list_of":"/method/linear-warmup-with-cosine-annealing","method":"Linear Warmup With Cosine Annealing","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":19,"pages_in_order":38,"rows_per_page":100,"rows":[1801,1900],"of":3797,"counts":{"archive_papers_tagged":3797,"with_a_code_link":1655,"where_syntology_ran_a_sample":602,"not_listed_spam_title":0,"listed":3797,"listed_where_code_ran":602,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":490,"every_run_a_failure_of_syntologys_instrument":112,"listed_with_a_run_with_no_instrument_failure":490,"listed_every_run_a_failure_of_syntologys_instrument":112,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/linear-warmup-with-cosine-annealing","prev":"/method/linear-warmup-with-cosine-annealing/papers/18","next":"/method/linear-warmup-with-cosine-annealing/papers/20","papers":[{"paper":null,"slug":"prospective-role-of-foundation-models-in","title":"Prospective Role of Foundation Models in Advancing Autonomous Vehicles","date":"2023-12-08","arxiv_id":"2405.02288","n_code_links":0,"syntology":null},{"paper":null,"slug":"user-aware-prefix-tuning-is-a-good-learner","title":"User-Aware Prefix-Tuning is a Good Learner for Personalized Image Captioning","date":"2023-12-08","arxiv_id":"2312.04793","n_code_links":0,"syntology":null},{"paper":null,"slug":"on-sarcasm-detection-with-openai-gpt-based","title":"On Sarcasm Detection with OpenAI GPT-based Models","date":"2023-12-07","arxiv_id":"2312.04642","n_code_links":0,"syntology":null},{"paper":null,"slug":"purple-llama-cyberseceval-a-secure-coding","title":"Purple Llama CyberSecEval: A Secure Coding Benchmark for Language Models","date":"2023-12-07","arxiv_id":"2312.04724","n_code_links":0,"syntology":null},{"paper":null,"slug":"holmes-towards-distributed-training-across","title":"Holmes: Towards Distributed Training Across Clusters with Heterogeneous NIC Environment","date":"2023-12-06","arxiv_id":"2312.03549","n_code_links":0,"syntology":null},{"paper":"/paper/not-all-large-language-models-llms-succumb-to","slug":"not-all-large-language-models-llms-succumb-to","title":"Exploring the Reversal Curse and Other Deductive Logical Reasoning in BERT and GPT-Based Large Language Models","date":"2023-12-06","arxiv_id":"2312.03633","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-hardware-evaluation-framework-for-large","title":"A Hardware Evaluation Framework for Large Language Model Inference","date":"2023-12-05","arxiv_id":"2312.03134","n_code_links":0,"syntology":null},{"paper":"/paper/draft-dense-retrieval-augmented-few-shot","slug":"draft-dense-retrieval-augmented-few-shot","title":"DRAFT: Dense Retrieval Augmented Few-shot Topic classifier Framework","date":"2023-12-05","arxiv_id":"2312.02532","n_code_links":1,"syntology":null},{"paper":null,"slug":"gpt-vs-human-for-scientific-reviews-a-dual","title":"GPT vs Human for Scientific Reviews: A Dual Source Review on Applications of ChatGPT in Science","date":"2023-12-05","arxiv_id":"2312.03769","n_code_links":0,"syntology":null},{"paper":null,"slug":"rank-without-gpt-building-gpt-independent","title":"Rank-without-GPT: Building GPT-Independent Listwise Rerankers on Open-Source Large Language Models","date":"2023-12-05","arxiv_id":"2312.02969","n_code_links":0,"syntology":null},{"paper":null,"slug":"towards-more-unified-in-context-visual","title":"Towards More Unified In-context Visual Understanding","date":"2023-12-05","arxiv_id":"2312.02520","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-survey-on-large-language-model-llm-security","title":"A Survey on Large Language Model (LLM) Security and Privacy: The Good, the Bad, and the Ugly","date":"2023-12-04","arxiv_id":"2312.02003","n_code_links":0,"syntology":null},{"paper":null,"slug":"jellyfish-a-large-language-model-for-data","title":"Jellyfish: A Large Language Model for Data Preprocessing","date":"2023-12-04","arxiv_id":"2312.01678","n_code_links":0,"syntology":null},{"paper":"/paper/tree-of-attacks-jailbreaking-black-box-llms","slug":"tree-of-attacks-jailbreaking-black-box-llms","title":"Tree of Attacks: Jailbreaking Black-Box LLMs Automatically","date":"2023-12-04","arxiv_id":"2312.02119","n_code_links":2,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["ricommunity/tap"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/nlebench-norglm-a-comprehensive-empirical","slug":"nlebench-norglm-a-comprehensive-empirical","title":"NLEBench+NorGLM: A Comprehensive Empirical Analysis and Benchmark Dataset for Generative Language Models in Norwegian","date":"2023-12-03","arxiv_id":"2312.01314","n_code_links":1,"syntology":{"ran":0,"of":4,"n_ran_checked":0,"n_instrument":0,"unverified":4,"pointer_only":4,"phrase":"0 ran · 4 unverified","official":{"repos":["smartmedia-ai/norglm"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":4,"ran_from_kinds":[]}}},{"paper":"/paper/a-ripple-in-time-a-discontinuity-in-american","slug":"a-ripple-in-time-a-discontinuity-in-american","title":"A ripple in time: a discontinuity in American history","date":"2023-12-02","arxiv_id":"2312.01185","n_code_links":1,"syntology":null},{"paper":"/paper/harnessing-the-power-of-prompt-based","slug":"harnessing-the-power-of-prompt-based","title":"Harnessing the Power of Prompt-based Techniques for Generating School-Level Questions using Large Language Models","date":"2023-12-02","arxiv_id":"2312.01032","n_code_links":1,"syntology":null},{"paper":"/paper/large-language-models-are-zero-shot-text","slug":"large-language-models-are-zero-shot-text","title":"Large Language Models Are Zero-Shot Text Classifiers","date":"2023-12-02","arxiv_id":"2312.01044","n_code_links":1,"syntology":null},{"paper":"/paper/gift-generative-interpretable-fine-tuning","slug":"gift-generative-interpretable-fine-tuning","title":"Generative Parameter-Efficient Fine-Tuning","date":"2023-12-01","arxiv_id":"2312.00700","n_code_links":1,"syntology":{"ran":12,"of":16,"n_ran_checked":12,"n_instrument":0,"unverified":4,"pointer_only":10,"phrase":"12 ran (of which 0 constructed an object rather than computing a result; 12 with no instrument failure: 0 honoured, 0 violated, 12 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","official":{"repos":["savadikarc/gift"],"state":"official (archive's flag): 12 ran","n_ran":12,"n_constructed":0,"n_ran_no_instrument_failure":12,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"applying-large-language-models-and-chain-of","title":"Applying Large Language Models and Chain-of-Thought for Automatic Scoring","date":"2023-11-30","arxiv_id":"2312.03748","n_code_links":0,"syntology":null},{"paper":null,"slug":"iag-induction-augmented-generation-framework","title":"IAG: Induction-Augmented Generation Framework for Answering Reasoning Questions","date":"2023-11-30","arxiv_id":"2311.18397","n_code_links":0,"syntology":null},{"paper":"/paper/robust-concept-erasure-via-kernelized-rate-1","slug":"robust-concept-erasure-via-kernelized-rate-1","title":"Robust Concept Erasure via Kernelized Rate-Distortion Maximization","date":"2023-11-30","arxiv_id":"2312.00194","n_code_links":1,"syntology":{"ran":20,"of":21,"n_ran_checked":18,"n_instrument":2,"unverified":1,"pointer_only":2,"phrase":"20 ran (of which 0 constructed an object rather than computing a result; 18 with no instrument failure: 0 honoured, 0 violated, 18 with no contract checked; 2 where Syntology's instrument failed) · 1 unverified","official":{"repos":["brcsomnath/kram"],"state":"official (archive's flag): 20 ran","n_ran":20,"n_constructed":0,"n_ran_no_instrument_failure":18,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/biomedical-knowledge-graph-enhanced-prompt","slug":"biomedical-knowledge-graph-enhanced-prompt","title":"Biomedical knowledge graph-optimized prompt generation for large language models","date":"2023-11-29","arxiv_id":"2311.17330","n_code_links":1,"syntology":{"ran":0,"of":3,"n_ran_checked":0,"n_instrument":0,"unverified":3,"pointer_only":0,"phrase":"0 ran · 3 unverified","official":{"repos":["BaranziniLab/KG_RAG"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":3,"ran_from_kinds":[]}}},{"paper":null,"slug":"improving-the-robustness-of-transformer-based","title":"Improving the Robustness of Transformer-based Large Language Models with Dynamic Attention","date":"2023-11-29","arxiv_id":"2311.17400","n_code_links":0,"syntology":null},{"paper":null,"slug":"timelygpt-recurrent-convolutional-transformer","title":"TimelyGPT: Extrapolatable Transformer Pre-training for Long-term Time-Series Forecasting in Healthcare","date":"2023-11-29","arxiv_id":"2312.00817","n_code_links":0,"syntology":null},{"paper":"/paper/characterglm-customizing-chinese","slug":"characterglm-customizing-chinese","title":"CharacterGLM: Customizing Chinese Conversational AI Characters with Large Language Models","date":"2023-11-28","arxiv_id":"2311.16832","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["thu-coai/characterglm-6b"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/chatgpt-s-one-year-anniversary-are-open","slug":"chatgpt-s-one-year-anniversary-are-open","title":"ChatGPT's One-year Anniversary: Are Open-Source Large Language Models Catching up?","date":"2023-11-28","arxiv_id":"2311.16989","n_code_links":1,"syntology":null},{"paper":null,"slug":"cole-a-hierarchical-generation-framework-for","title":"COLE: A Hierarchical Generation Framework for Multi-Layered and Editable Graphic Design","date":"2023-11-28","arxiv_id":"2311.16974","n_code_links":0,"syntology":null},{"paper":null,"slug":"comparing-generative-chatbots-based-on","title":"Comparing Generative Chatbots Based on Process Requirements","date":"2023-11-28","arxiv_id":"2312.03741","n_code_links":0,"syntology":null},{"paper":"/paper/seed-bench-2-benchmarking-multimodal-large","slug":"seed-bench-2-benchmarking-multimodal-large","title":"SEED-Bench-2: Benchmarking Multimodal Large Language Models","date":"2023-11-28","arxiv_id":"2311.17092","n_code_links":2,"syntology":null},{"paper":"/paper/bert-goes-off-topic-investigating-the-domain","slug":"bert-goes-off-topic-investigating-the-domain","title":"BERT Goes Off-Topic: Investigating the Domain Transfer Challenge using Genre Classification","date":"2023-11-27","arxiv_id":"2311.16083","n_code_links":1,"syntology":null},{"paper":null,"slug":"decoding-logic-errors-a-comparative-study-on","title":"Decoding Logic Errors: A Comparative Study on Bug Detection by Students and Large Language Models","date":"2023-11-27","arxiv_id":"2311.16017","n_code_links":0,"syntology":null},{"paper":null,"slug":"real-customization-or-just-marketing-are","title":"Real Customization or Just Marketing: Are Customized Versions of Chat GPT Useful?","date":"2023-11-27","arxiv_id":"2312.03728","n_code_links":0,"syntology":null},{"paper":null,"slug":"leveraging-ai-derived-data-for-carbon","title":"Leveraging AI-derived Data for Carbon Accounting: Information Extraction from Alternative Sources","date":"2023-11-26","arxiv_id":"2312.03722","n_code_links":0,"syntology":null},{"paper":"/paper/machine-generated-text-detection-using-deep","slug":"machine-generated-text-detection-using-deep","title":"Machine-Generated Text Detection using Deep Learning","date":"2023-11-26","arxiv_id":"2311.15425","n_code_links":1,"syntology":null},{"paper":"/paper/uhgeval-benchmarking-the-hallucination-of","slug":"uhgeval-benchmarking-the-hallucination-of","title":"UHGEval: Benchmarking the Hallucination of Chinese Large Language Models via Unconstrained Generation","date":"2023-11-26","arxiv_id":"2311.15296","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":3,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["IAAR-Shanghai/UHGEval"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"cmed-gpt-prompt-tuning-for-entity-aware","title":"CMed-GPT: Prompt Tuning for Entity-Aware Chinese Medical Dialogue Generation","date":"2023-11-24","arxiv_id":"2311.14539","n_code_links":0,"syntology":null},{"paper":"/paper/data-to-text-bilingual-generation","slug":"data-to-text-bilingual-generation","title":"Data-to-Text Bilingual Generation","date":"2023-11-24","arxiv_id":"2311.14808","n_code_links":2,"syntology":null},{"paper":"/paper/gpt-struct-me-probing-gpt-models-on-narrative","slug":"gpt-struct-me-probing-gpt-models-on-narrative","title":"GPT Struct Me: Probing GPT Models on Narrative Entity Extraction","date":"2023-11-24","arxiv_id":"2311.14583","n_code_links":1,"syntology":null},{"paper":null,"slug":"large-language-models-as-automated-aligners","title":"Large Language Models as Automated Aligners for benchmarking Vision-Language Models","date":"2023-11-24","arxiv_id":"2311.14580","n_code_links":0,"syntology":null},{"paper":null,"slug":"machine-translation-for-ge-ez-language","title":"Machine Translation for Ge'ez Language","date":"2023-11-24","arxiv_id":"2311.14530","n_code_links":0,"syntology":null},{"paper":"/paper/a-cross-attention-approach-to-diagnostic","slug":"a-cross-attention-approach-to-diagnostic","title":"A Cross Attention Approach to Diagnostic Explainability using Clinical Practice Guidelines for Depression","date":"2023-11-23","arxiv_id":"2311.13852","n_code_links":1,"syntology":null},{"paper":"/paper/hardware-resilience-properties-of-text-guided-1","slug":"hardware-resilience-properties-of-text-guided-1","title":"Hardware Resilience Properties of Text-Guided Image Classifiers","date":"2023-11-23","arxiv_id":"2311.14062","n_code_links":1,"syntology":{"ran":3,"of":4,"n_ran_checked":1,"n_instrument":2,"unverified":1,"pointer_only":4,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 2 where Syntology's instrument failed) · 1 unverified","official":{"repos":["talalwasim/textguidedresilience"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"minimizing-factual-inconsistency-and","title":"Minimizing Factual Inconsistency and Hallucination in Large Language Models","date":"2023-11-23","arxiv_id":"2311.13878","n_code_links":0,"syntology":null},{"paper":null,"slug":"towards-auditing-large-language-models","title":"Towards Auditing Large Language Models: Improving Text-based Stereotype Detection","date":"2023-11-23","arxiv_id":"2311.14126","n_code_links":0,"syntology":null},{"paper":"/paper/comparison-of-pipeline-sequence-to-sequence","slug":"comparison-of-pipeline-sequence-to-sequence","title":"Comparison of pipeline, sequence-to-sequence, and GPT models for end-to-end relation extraction: experiments with the rare disease use-case","date":"2023-11-22","arxiv_id":"2311.13729","n_code_links":1,"syntology":null},{"paper":null,"slug":"drilling-down-into-the-discourse-structure","title":"Drilling Down into the Discourse Structure with LLMs for Long Document Question Answering","date":"2023-11-22","arxiv_id":"2311.13565","n_code_links":0,"syntology":null},{"paper":null,"slug":"generation-of-explanations-for-logic","title":"Generation of Explanations for Logic Reasoning","date":"2023-11-22","arxiv_id":"2311.13455","n_code_links":0,"syntology":null},{"paper":null,"slug":"nova-generative-language-models-for-binaries","title":"Nova: Generative Language Models for Assembly Code with Hierarchical Attention and Contrastive Learning","date":"2023-11-22","arxiv_id":"2311.13721","n_code_links":0,"syntology":null},{"paper":"/paper/pg-video-llava-pixel-grounding-large-video","slug":"pg-video-llava-pixel-grounding-large-video","title":"PG-Video-LLaVA: Pixel Grounding Large Video-Language Models","date":"2023-11-22","arxiv_id":"2311.13435","n_code_links":1,"syntology":{"ran":3,"of":5,"n_ran_checked":2,"n_instrument":1,"unverified":2,"pointer_only":5,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","official":{"repos":["mbzuai-oryx/video-llava"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/speak-like-a-native-prompting-large-language","slug":"speak-like-a-native-prompting-large-language","title":"AlignedCoT: Prompting Large Language Models via Native-Speaking Demonstrations","date":"2023-11-22","arxiv_id":"2311.13538","n_code_links":1,"syntology":null},{"paper":null,"slug":"ve-a-chatbot-for-latin","title":"@ve: A Chatbot for Latin","date":"2023-11-22","arxiv_id":"2311.14741","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-survey-on-large-language-models-for-1","title":"A Survey on Large Language Models for Personalized and Explainable Recommendations","date":"2023-11-21","arxiv_id":"2311.12338","n_code_links":0,"syntology":null},{"paper":null,"slug":"academicgpt-empowering-academic-research","title":"AcademicGPT: Empowering Academic Research","date":"2023-11-21","arxiv_id":"2311.12315","n_code_links":0,"syntology":null},{"paper":"/paper/alpha-anomalous-physiological-health","slug":"alpha-anomalous-physiological-health","title":"ALPHA: AnomaLous Physiological Health Assessment Using Large Language Models","date":"2023-11-21","arxiv_id":"2311.12524","n_code_links":1,"syntology":null},{"paper":"/paper/descriptor-and-word-soups-overcoming-the","slug":"descriptor-and-word-soups-overcoming-the","title":"Descriptor and Word Soups: Overcoming the Parameter Efficiency Accuracy Tradeoff for Out-of-Distribution Few-shot Learning","date":"2023-11-21","arxiv_id":"2311.13612","n_code_links":1,"syntology":null},{"paper":"/paper/extracting-definienda-in-mathematical","slug":"extracting-definienda-in-mathematical","title":"Extracting Definienda in Mathematical Scholarly Articles with Transformers","date":"2023-11-21","arxiv_id":"2311.12448","n_code_links":2,"syntology":null},{"paper":null,"slug":"gpt4motion-scripting-physical-motions-in-text","title":"GPT4Motion: Scripting Physical Motions in Text-to-Video Generation via Blender-Oriented GPT Planning","date":"2023-11-21","arxiv_id":"2311.12631","n_code_links":0,"syntology":null},{"paper":null,"slug":"interprompt-interpretable-prompting-for","title":"InterPrompt: Interpretable Prompting for Interrelated Interpersonal Risk Factors in Reddit Posts","date":"2023-11-21","arxiv_id":"2311.12404","n_code_links":0,"syntology":null},{"paper":"/paper/assessing-prompt-injection-risks-in-200","slug":"assessing-prompt-injection-risks-in-200","title":"Assessing Prompt Injection Risks in 200+ Custom GPTs","date":"2023-11-20","arxiv_id":"2311.11538","n_code_links":1,"syntology":null},{"paper":"/paper/evil-geniuses-delving-into-the-safety-of-llm","slug":"evil-geniuses-delving-into-the-safety-of-llm","title":"Evil Geniuses: Delving into the Safety of LLM-based Agents","date":"2023-11-20","arxiv_id":"2311.11855","n_code_links":1,"syntology":{"ran":5,"of":5,"n_ran_checked":5,"n_instrument":0,"unverified":0,"pointer_only":5,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["t1ans1r/evil-geniuses"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/how-to-use-large-language-models-for-text","slug":"how-to-use-large-language-models-for-text","title":"Towards Human-Level Text Coding with LLMs: The Case of Fatherhood Roles in Public Policy Documents","date":"2023-11-20","arxiv_id":"2311.11844","n_code_links":1,"syntology":null},{"paper":null,"slug":"memorycompanion-a-smart-healthcare-solution","title":"MemoryCompanion: A Smart Healthcare Solution to Empower Efficient Alzheimer's Care Via Unleashing Generative AI","date":"2023-11-20","arxiv_id":"2311.14730","n_code_links":0,"syntology":null},{"paper":null,"slug":"refactoring-programs-using-large-language","title":"Refactoring Programs Using Large Language Models with Few-Shot Examples","date":"2023-11-20","arxiv_id":"2311.11690","n_code_links":0,"syntology":null},{"paper":null,"slug":"spot-the-bot-distinguishing-human-written-and","title":"Spot the Bot: Distinguishing Human-Written and Bot-Generated Texts Using Clustering and Information Theory Techniques","date":"2023-11-19","arxiv_id":"2311.11441","n_code_links":0,"syntology":null},{"paper":null,"slug":"behavior-optimized-image-generation","title":"Behavior Optimized Image Generation","date":"2023-11-18","arxiv_id":"2311.10995","n_code_links":0,"syntology":null},{"paper":null,"slug":"advancements-in-generative-ai-a-comprehensive","title":"Advancements in Generative AI: A Comprehensive Review of GANs, GPT, Autoencoders, Diffusion Model, and Transformers","date":"2023-11-17","arxiv_id":"2311.10242","n_code_links":0,"syntology":null},{"paper":null,"slug":"bias-a-head-analyzing-bias-in-transformer","title":"Bias A-head? Analyzing Bias in Transformer-Based Language Model Attention Heads","date":"2023-11-17","arxiv_id":"2311.10395","n_code_links":0,"syntology":null},{"paper":"/paper/camels-in-a-changing-climate-enhancing-lm","slug":"camels-in-a-changing-climate-enhancing-lm","title":"Camels in a Changing Climate: Enhancing LM Adaptation with Tulu 2","date":"2023-11-17","arxiv_id":"2311.10702","n_code_links":3,"syntology":null},{"paper":"/paper/dynapipe-optimizing-multi-task-training","slug":"dynapipe-optimizing-multi-task-training","title":"DynaPipe: Optimizing Multi-task Training through Dynamic Pipelines","date":"2023-11-17","arxiv_id":"2311.10418","n_code_links":2,"syntology":null},{"paper":"/paper/event-causality-is-key-to-computational-story","slug":"event-causality-is-key-to-computational-story","title":"Event Causality Is Key to Computational Story Understanding","date":"2023-11-16","arxiv_id":"2311.09648","n_code_links":1,"syntology":{"ran":6,"of":9,"n_ran_checked":6,"n_instrument":0,"unverified":3,"pointer_only":9,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":{"repos":["insundaycathy/event-causality-extraction"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"fumbling-in-babel-an-investigation-into","title":"Fumbling in Babel: An Investigation into ChatGPT's Language Identification Ability","date":"2023-11-16","arxiv_id":"2311.09696","n_code_links":0,"syntology":null},{"paper":null,"slug":"generative-ai-for-hate-speech-detection","title":"Generative AI for Hate Speech Detection: Evaluation and Findings","date":"2023-11-16","arxiv_id":"2311.09993","n_code_links":0,"syntology":null},{"paper":null,"slug":"human-still-wins-over-llm-an-empirical-study","title":"Human Still Wins over LLM: An Empirical Study of Active Learning on Domain-Specific Annotation Tasks","date":"2023-11-16","arxiv_id":"2311.09825","n_code_links":0,"syntology":null},{"paper":"/paper/intervenor-prompt-the-coding-ability-of-large","slug":"intervenor-prompt-the-coding-ability-of-large","title":"INTERVENOR: Prompting the Coding Ability of Large Language Models with the Interactive Chain of Repair","date":"2023-11-16","arxiv_id":"2311.09868","n_code_links":1,"syntology":null},{"paper":"/paper/knowledgemath-knowledge-intensive-math-word","slug":"knowledgemath-knowledge-intensive-math-word","title":"FinanceMath: Knowledge-Intensive Math Reasoning in Finance Domains","date":"2023-11-16","arxiv_id":"2311.09797","n_code_links":1,"syntology":{"ran":12,"of":15,"n_ran_checked":12,"n_instrument":0,"unverified":3,"pointer_only":15,"phrase":"12 ran (of which 0 constructed an object rather than computing a result; 12 with no instrument failure: 0 honoured, 0 violated, 12 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":{"repos":["yale-nlp/knowledgemath"],"state":"official (archive's flag): 12 ran","n_ran":12,"n_constructed":0,"n_ran_no_instrument_failure":12,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"on-retrieval-augmentation-and-the-limitations","title":"On Retrieval Augmentation and the Limitations of Language Model Training","date":"2023-11-16","arxiv_id":"2311.09615","n_code_links":0,"syntology":null},{"paper":null,"slug":"predictive-minds-llms-as-atypical-active","title":"Predictive Minds: LLMs As Atypical Active Inference Agents","date":"2023-11-16","arxiv_id":"2311.10215","n_code_links":0,"syntology":null},{"paper":null,"slug":"reducing-privacy-risks-in-online-self","title":"Reducing Privacy Risks in Online Self-Disclosures with Language Models","date":"2023-11-16","arxiv_id":"2311.09538","n_code_links":0,"syntology":null},{"paper":"/paper/can-large-language-models-follow-concept","slug":"can-large-language-models-follow-concept","title":"Can Large Language Models Follow Concept Annotation Guidelines? A Case Study on Scientific and Financial Domains","date":"2023-11-15","arxiv_id":"2311.08704","n_code_links":1,"syntology":null},{"paper":null,"slug":"evaluating-gender-bias-in-the-translation-of","title":"Evaluating Gender Bias in the Translation of Gender-Neutral Languages into English","date":"2023-11-15","arxiv_id":"2311.08836","n_code_links":0,"syntology":null},{"paper":null,"slug":"loke-linked-open-knowledge-extraction-for","title":"LOKE: Linked Open Knowledge Extraction for Automated Knowledge Graph Construction","date":"2023-11-15","arxiv_id":"2311.09366","n_code_links":0,"syntology":null},{"paper":"/paper/tooltalk-evaluating-tool-usage-in-a","slug":"tooltalk-evaluating-tool-usage-in-a","title":"ToolTalk: Evaluating Tool-Usage in a Conversational Setting","date":"2023-11-15","arxiv_id":"2311.10775","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":null}},{"paper":"/paper/we-demand-justice-towards-grounding-political","slug":"we-demand-justice-towards-grounding-political","title":"\"We Demand Justice!\": Towards Social Context Grounding of Political Texts","date":"2023-11-15","arxiv_id":"2311.09106","n_code_links":1,"syntology":null},{"paper":"/paper/a-survey-on-language-models-for-code","slug":"a-survey-on-language-models-for-code","title":"Unifying the Perspectives of NLP and Software Engineering: A Survey on Language Models for Code","date":"2023-11-14","arxiv_id":"2311.07989","n_code_links":1,"syntology":null},{"paper":null,"slug":"cpopqa-ranking-cultural-concept-popularity-by","title":"CPopQA: Ranking Cultural Concept Popularity by LLMs","date":"2023-11-14","arxiv_id":"2311.07897","n_code_links":0,"syntology":null},{"paper":null,"slug":"evaluating-llms-on-document-based-qa-exact","title":"Evaluating LLMs on Document-Based QA: Exact Answer Selection and Numerical Extraction using Cogtale dataset","date":"2023-11-14","arxiv_id":"2311.07878","n_code_links":0,"syntology":null},{"paper":"/paper/fair-abstractive-summarization-of-diverse","slug":"fair-abstractive-summarization-of-diverse","title":"Fair Abstractive Summarization of Diverse Perspectives","date":"2023-11-14","arxiv_id":"2311.07884","n_code_links":1,"syntology":null},{"paper":"/paper/language-models-are-better-bug-detector","slug":"language-models-are-better-bug-detector","title":"Language Models are Better Bug Detector Through Code-Pair Classification","date":"2023-11-14","arxiv_id":"2311.07957","n_code_links":1,"syntology":null},{"paper":null,"slug":"large-language-model-driven-classroom","title":"Large Language Model-Driven Classroom Flipping: Empowering Student-Centric Peer Questioning with Flipped Interaction","date":"2023-11-14","arxiv_id":"2311.14708","n_code_links":0,"syntology":null},{"paper":"/paper/memory-efficient-stochastic-methods-for","slug":"memory-efficient-stochastic-methods-for","title":"Memory-efficient Stochastic methods for Memory-based Transformers","date":"2023-11-14","arxiv_id":"2311.08123","n_code_links":1,"syntology":null},{"paper":"/paper/do-large-language-models-and-humans-have","slug":"do-large-language-models-and-humans-have","title":"Do large language models and humans have similar behaviors in causal inference with script knowledge?","date":"2023-11-13","arxiv_id":"2311.07311","n_code_links":1,"syntology":null},{"paper":"/paper/in-context-learning-generalizes-but-not","slug":"in-context-learning-generalizes-but-not","title":"In-context Learning Generalizes, But Not Always Robustly: The Case of Syntax","date":"2023-11-13","arxiv_id":"2311.07811","n_code_links":1,"syntology":null},{"paper":"/paper/it-s-not-easy-being-wrong-evaluating-process","slug":"it-s-not-easy-being-wrong-evaluating-process","title":"It's Not Easy Being Wrong: Large Language Models Struggle with Process of Elimination Reasoning","date":"2023-11-13","arxiv_id":"2311.07532","n_code_links":1,"syntology":null},{"paper":null,"slug":"language-model-in-the-loop-data-optimal","title":"Language Model-In-The-Loop: Data Optimal Approach to Learn-To-Recommend Actions in Text Games","date":"2023-11-13","arxiv_id":"2311.07687","n_code_links":0,"syntology":null},{"paper":null,"slug":"megaverse-benchmarking-large-language-models","title":"MEGAVERSE: Benchmarking Large Language Models Across Languages, Modalities, Models and Tasks","date":"2023-11-13","arxiv_id":"2311.07463","n_code_links":0,"syntology":null},{"paper":null,"slug":"on-the-truthfulness-of-surprisingly-likely","title":"On The Truthfulness of 'Surprisingly Likely' Responses of Large Language Models","date":"2023-11-13","arxiv_id":"2311.07692","n_code_links":0,"syntology":null},{"paper":null,"slug":"speech-based-slot-filling-using-large","title":"Speech-based Slot Filling using Large Language Models","date":"2023-11-13","arxiv_id":"2311.07418","n_code_links":0,"syntology":null},{"paper":"/paper/steer-unified-style-transfer-with-expert","slug":"steer-unified-style-transfer-with-expert","title":"STEER: Unified Style Transfer with Expert Reinforcement","date":"2023-11-13","arxiv_id":"2311.07167","n_code_links":1,"syntology":null},{"paper":null,"slug":"from-complex-to-simple-unraveling-the","title":"From Complex to Simple: Unraveling the Cognitive Tree for Reasoning with Small Language Models","date":"2023-11-12","arxiv_id":"2311.06754","n_code_links":0,"syntology":null}],"record_sha256":"de52cec8ef0dd5cdb03161662d3b43ab73048bc29b70a3231d9cc31f73eda7af","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}