{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/linear-warmup-with-cosine-annealing/papers/5","list_of":"/method/linear-warmup-with-cosine-annealing","method":"Linear Warmup With Cosine Annealing","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":5,"pages_in_order":38,"rows_per_page":100,"rows":[401,500],"of":3797,"counts":{"archive_papers_tagged":3797,"with_a_code_link":1655,"where_syntology_ran_a_sample":602,"not_listed_spam_title":0,"listed":3797,"listed_where_code_ran":602,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":490,"every_run_a_failure_of_syntologys_instrument":112,"listed_with_a_run_with_no_instrument_failure":490,"listed_every_run_a_failure_of_syntologys_instrument":112,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/linear-warmup-with-cosine-annealing","prev":"/method/linear-warmup-with-cosine-annealing/papers/4","next":"/method/linear-warmup-with-cosine-annealing/papers/6","papers":[{"paper":null,"slug":"auto-generating-earnings-report-analysis-via","title":"Auto-Generating Earnings Report Analysis via a Financial-Augmented LLM","date":"2024-12-11","arxiv_id":"2412.08179","n_code_links":0,"syntology":null},{"paper":"/paper/exploiting-the-index-gradients-for","slug":"exploiting-the-index-gradients-for","title":"Exploiting the Index Gradients for Optimization-Based Jailbreaking on Large Language Models","date":"2024-12-11","arxiv_id":"2412.08615","n_code_links":1,"syntology":null},{"paper":null,"slug":"graphtool-instruction-revolutionizing-graph","title":"GraphTool-Instruction: Revolutionizing Graph Reasoning in LLMs through Decomposed Subtask Instruction","date":"2024-12-11","arxiv_id":"2412.12152","n_code_links":0,"syntology":null},{"paper":null,"slug":"imitate-before-detect-aligning-machine","title":"Imitate Before Detect: Aligning Machine Stylistic Preference for Machine-Revised Text Detection","date":"2024-12-11","arxiv_id":"2412.10432","n_code_links":0,"syntology":null},{"paper":null,"slug":"large-language-models-still-face-challenges","title":"Large Language Models Still Face Challenges in Multi-Hop Reasoning with External Knowledge","date":"2024-12-11","arxiv_id":"2412.08317","n_code_links":0,"syntology":null},{"paper":"/paper/causal-world-representation-in-the-gpt-model","slug":"causal-world-representation-in-the-gpt-model","title":"A Causal World Model Underlying Next Token Prediction: Exploring GPT in a Controlled Environment","date":"2024-12-10","arxiv_id":"2412.07446","n_code_links":1,"syntology":null},{"paper":null,"slug":"gpt-2-through-the-lens-of-vector-symbolic","title":"GPT-2 Through the Lens of Vector Symbolic Architectures","date":"2024-12-10","arxiv_id":"2412.07947","n_code_links":0,"syntology":null},{"paper":"/paper/intellectseeker-a-personalized-literature","slug":"intellectseeker-a-personalized-literature","title":"IntellectSeeker: A Personalized Literature Management System with the Probabilistic Model and Large Language Model","date":"2024-12-10","arxiv_id":"2412.07213","n_code_links":1,"syntology":null},{"paper":"/paper/rag-based-question-answering-over","slug":"rag-based-question-answering-over","title":"RAG-based Question Answering over Heterogeneous Data and Text","date":"2024-12-10","arxiv_id":"2412.07420","n_code_links":0,"syntology":null},{"paper":"/paper/superficial-consciousness-hypothesis-for","slug":"superficial-consciousness-hypothesis-for","title":"Superficial Consciousness Hypothesis for Autoregressive Transformers","date":"2024-12-10","arxiv_id":"2412.07278","n_code_links":1,"syntology":null},{"paper":null,"slug":"towards-predictive-communication-with-brain","title":"Towards Predictive Communication with Brain-Computer Interfaces integrating Large Language Models","date":"2024-12-10","arxiv_id":"2412.07355","n_code_links":0,"syntology":null},{"paper":"/paper/batchtopk-sparse-autoencoders","slug":"batchtopk-sparse-autoencoders","title":"BatchTopK Sparse Autoencoders","date":"2024-12-09","arxiv_id":"2412.06410","n_code_links":2,"syntology":{"ran":2,"of":3,"n_ran_checked":2,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["bartbussmann/batchtopk"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"exploring-memorization-and-copyright","title":"Exploring Memorization and Copyright Violation in Frontier LLMs: A Study of the New York Times v. OpenAI 2023 Lawsuit","date":"2024-12-09","arxiv_id":"2412.06370","n_code_links":0,"syntology":null},{"paper":null,"slug":"optimizing-multi-task-learning-for-enhanced","title":"Optimizing Multi-Task Learning for Enhanced Performance in Large Language Models","date":"2024-12-09","arxiv_id":"2412.06249","n_code_links":0,"syntology":null},{"paper":null,"slug":"the-rosetta-paradox-domain-specific","title":"The Rosetta Paradox: Domain-Specific Performance Inversions in Large Language Models","date":"2024-12-09","arxiv_id":"2412.17821","n_code_links":0,"syntology":null},{"paper":"/paper/m-3-20m-a-large-scale-multi-modal-molecule","slug":"m-3-20m-a-large-scale-multi-modal-molecule","title":"M$^{3}$-20M: A Large-Scale Multi-Modal Molecule Dataset for AI-driven Drug Design and Discovery","date":"2024-12-08","arxiv_id":"2412.06847","n_code_links":1,"syntology":null},{"paper":"/paper/characterbox-evaluating-the-role-playing","slug":"characterbox-evaluating-the-role-playing","title":"CharacterBox: Evaluating the Role-Playing Capabilities of LLMs in Text-Based Virtual Worlds","date":"2024-12-07","arxiv_id":"2412.05631","n_code_links":1,"syntology":{"ran":3,"of":5,"n_ran_checked":3,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["paitesanshi/characterbox"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"exploring-the-use-of-llms-for-sql-equivalence","title":"Can the Rookies Cut the Tough Cookie? Exploring the Use of LLMs for SQL Equivalence Checking","date":"2024-12-07","arxiv_id":"2412.05561","n_code_links":0,"syntology":null},{"paper":"/paper/privagent-agentic-based-red-teaming-for-llm","slug":"privagent-agentic-based-red-teaming-for-llm","title":"PrivAgent: Agentic-based Red-teaming for LLM Privacy Leakage","date":"2024-12-07","arxiv_id":"2412.05734","n_code_links":1,"syntology":null},{"paper":null,"slug":"100-hallucination-elimination-using-acurai","title":"100% Elimination of Hallucinations on RAGTruth for GPT-4 and GPT-3.5 Turbo","date":"2024-12-06","arxiv_id":"2412.05223","n_code_links":0,"syntology":null},{"paper":null,"slug":"are-frontier-large-language-models-suitable","title":"Are Frontier Large Language Models Suitable for Q&A in Science Centres?","date":"2024-12-06","arxiv_id":"2412.05200","n_code_links":0,"syntology":null},{"paper":null,"slug":"queen-a-large-language-model-for-quechua","title":"QueEn: A Large Language Model for Quechua-English Translation","date":"2024-12-06","arxiv_id":"2412.05184","n_code_links":0,"syntology":null},{"paper":null,"slug":"how-good-is-chatgpt-in-giving-adaptive","title":"How Good is ChatGPT in Giving Adaptive Guidance Using Knowledge Graphs in E-Learning Environments?","date":"2024-12-05","arxiv_id":"2412.03856","n_code_links":0,"syntology":null},{"paper":null,"slug":"controlling-the-mutation-in-large-language","title":"Controlling the Mutation in Large Language Models for the Efficient Evolution of Algorithms","date":"2024-12-04","arxiv_id":"2412.03250","n_code_links":0,"syntology":null},{"paper":null,"slug":"compressing-kv-cache-for-long-context-llm","title":"Compressing KV Cache for Long-Context LLM Inference with Inter-Layer Attention Similarity","date":"2024-12-03","arxiv_id":"2412.02252","n_code_links":0,"syntology":null},{"paper":"/paper/dp-2stage-adapting-language-models-as","slug":"dp-2stage-adapting-language-models-as","title":"DP-2Stage: Adapting Language Models as Differentially Private Tabular Data Generators","date":"2024-12-03","arxiv_id":"2412.02467","n_code_links":1,"syntology":null},{"paper":null,"slug":"flattering-to-deceive-the-impact-of","title":"Flattering to Deceive: The Impact of Sycophantic Behavior on User Trust in Large Language Model","date":"2024-12-03","arxiv_id":"2412.02802","n_code_links":0,"syntology":null},{"paper":null,"slug":"impact-of-data-snooping-on-deep-learning","title":"Impact of Data Snooping on Deep Learning Models for Locating Vulnerabilities in Lifted Code","date":"2024-12-03","arxiv_id":"2412.02048","n_code_links":0,"syntology":null},{"paper":null,"slug":"the-asymptotic-behavior-of-attention-in","title":"The Asymptotic Behavior of Attention in Transformers","date":"2024-12-03","arxiv_id":"2412.02682","n_code_links":0,"syntology":null},{"paper":null,"slug":"su-roberta-a-semi-supervised-approach-to","title":"Su-RoBERTa: A Semi-supervised Approach to Predicting Suicide Risk through Social Media using Base Language Models","date":"2024-12-02","arxiv_id":"2412.01353","n_code_links":0,"syntology":null},{"paper":null,"slug":"the-promise-and-peril-of-generative-ai","title":"The Promise and Peril of Generative AI: Evidence from GPT-4 as Sell-Side Analysts","date":"2024-12-02","arxiv_id":"2412.01069","n_code_links":0,"syntology":null},{"paper":null,"slug":"tokenizing-3d-molecule-structure-with","title":"Tokenizing 3D Molecule Structure with Quantized Spherical Coordinates","date":"2024-12-02","arxiv_id":"2412.01564","n_code_links":0,"syntology":null},{"paper":"/paper/a-comprehensive-guide-to-explainable-ai-from","slug":"a-comprehensive-guide-to-explainable-ai-from","title":"A Comprehensive Guide to Explainable AI: From Classical Models to LLMs","date":"2024-12-01","arxiv_id":"2412.00800","n_code_links":1,"syntology":null},{"paper":null,"slug":"eventgpt-event-stream-understanding-with","title":"EventGPT: Event Stream Understanding with Multimodal Large Language Models","date":"2024-12-01","arxiv_id":"2412.00832","n_code_links":0,"syntology":null},{"paper":null,"slug":"cdemapper-enhancing-nih-common-data-element","title":"CDEMapper: Enhancing NIH Common Data Element Normalization using Large Language Models","date":"2024-11-30","arxiv_id":"2412.00491","n_code_links":0,"syntology":null},{"paper":null,"slug":"cognitive-biases-in-large-language-models-a","title":"Cognitive Biases in Large Language Models: A Survey and Mitigation Experiments","date":"2024-11-30","arxiv_id":"2412.00323","n_code_links":0,"syntology":null},{"paper":null,"slug":"empowering-the-deaf-and-hard-of-hearing","title":"Empowering the Deaf and Hard of Hearing Community: Enhancing Video Captions Using Large Language Models","date":"2024-11-30","arxiv_id":"2412.00342","n_code_links":0,"syntology":null},{"paper":null,"slug":"forma-mentis-networks-predict-creativity","title":"Forma mentis networks predict creativity ratings of short texts via interpretable artificial intelligence in human and GPT-simulated raters","date":"2024-11-30","arxiv_id":"2412.00530","n_code_links":0,"syntology":null},{"paper":null,"slug":"automatic-prompt-generation-and-grounding","title":"Automatic Prompt Generation and Grounding Object Detection for Zero-Shot Image Anomaly Detection","date":"2024-11-28","arxiv_id":"2411.19220","n_code_links":0,"syntology":null},{"paper":null,"slug":"beautimeter-harnessing-gpt-for-assessing","title":"Beautimeter: Harnessing GPT for Assessing Architectural and Urban Beauty based on the 15 Properties of Living Structure","date":"2024-11-28","arxiv_id":"2411.19094","n_code_links":0,"syntology":null},{"paper":"/paper/deniahl-in-context-features-influence-llm","slug":"deniahl-in-context-features-influence-llm","title":"DENIAHL: In-Context Features Influence LLM Needle-In-A-Haystack Abilities","date":"2024-11-28","arxiv_id":"2411.19360","n_code_links":1,"syntology":null},{"paper":null,"slug":"habit-coach-customising-rag-based-chatbots-to","title":"Habit Coach: Customising RAG-based chatbots to support behavior change","date":"2024-11-28","arxiv_id":"2411.19229","n_code_links":0,"syntology":null},{"paper":null,"slug":"smartllmsentry-a-comprehensive-llm-based","title":"SmartLLMSentry: A Comprehensive LLM Based Smart Contract Vulnerability Detection Framework","date":"2024-11-28","arxiv_id":"2411.19234","n_code_links":0,"syntology":null},{"paper":null,"slug":"the-impact-of-example-selection-in-few-shot","title":"The Impact of Example Selection in Few-Shot Prompting on Automated Essay Scoring Using GPT Models","date":"2024-11-28","arxiv_id":"2411.18924","n_code_links":0,"syntology":null},{"paper":null,"slug":"automated-literature-review-using-nlp","title":"Automated Literature Review Using NLP Techniques and LLM-Based Retrieval-Augmented Generation","date":"2024-11-27","arxiv_id":"2411.18583","n_code_links":0,"syntology":null},{"paper":null,"slug":"can-bidirectional-encoder-become-the-ultimate","title":"Can bidirectional encoder become the ultimate winner for downstream applications of foundation models?","date":"2024-11-27","arxiv_id":"2411.18021","n_code_links":0,"syntology":null},{"paper":null,"slug":"chatgpt-as-speechwriter-for-the-french","title":"ChatGPT as speechwriter for the French presidents","date":"2024-11-27","arxiv_id":"2411.18382","n_code_links":0,"syntology":null},{"paper":"/paper/drs-deep-question-reformulation-with","slug":"drs-deep-question-reformulation-with","title":"DRS: Deep Question Reformulation With Structured Output","date":"2024-11-27","arxiv_id":"2411.17993","n_code_links":1,"syntology":null},{"paper":"/paper/streamlining-prediction-in-bayesian-deep","slug":"streamlining-prediction-in-bayesian-deep","title":"Streamlining Prediction in Bayesian Deep Learning","date":"2024-11-27","arxiv_id":"2411.18425","n_code_links":1,"syntology":{"ran":8,"of":10,"n_ran_checked":8,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 0 violated, 8 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["aaltoml/suq"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/training-and-evaluating-language-models-with","slug":"training-and-evaluating-language-models-with","title":"Training and Evaluating Language Models with Template-based Data Generation","date":"2024-11-27","arxiv_id":"2411.18104","n_code_links":1,"syntology":{"ran":3,"of":4,"n_ran_checked":1,"n_instrument":2,"unverified":1,"pointer_only":4,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 2 where Syntology's instrument failed) · 1 unverified","official":{"repos":["iiis-ai/templatemath"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"advancing-content-moderation-evaluating-large","title":"Advancing Content Moderation: Evaluating Large Language Models for Detecting Sensitive Content Across Text, Images, and Videos","date":"2024-11-26","arxiv_id":"2411.17123","n_code_links":0,"syntology":null},{"paper":null,"slug":"can-artificial-intelligence-predict-clinical","title":"Can artificial intelligence predict clinical trial outcomes?","date":"2024-11-26","arxiv_id":"2411.17595","n_code_links":0,"syntology":null},{"paper":"/paper/clover-constrained-learning-with-orthonormal","slug":"clover-constrained-learning-with-orthonormal","title":"CLOVER: Cross-Layer Orthogonal Vectors Pruning and Fine-Tuning","date":"2024-11-26","arxiv_id":"2411.17426","n_code_links":1,"syntology":null},{"paper":"/paper/distributed-sign-momentum-with-local-steps","slug":"distributed-sign-momentum-with-local-steps","title":"Distributed Sign Momentum with Local Steps for Training Transformers","date":"2024-11-26","arxiv_id":"2411.17866","n_code_links":1,"syntology":null},{"paper":null,"slug":"give-me-the-code-log-analysis-of-first-year","title":"\"Give me the code\" -- Log Analysis of First-Year CS Students' Interactions With GPT","date":"2024-11-26","arxiv_id":"2411.17855","n_code_links":0,"syntology":null},{"paper":null,"slug":"on-limitations-of-llm-as-annotator-for-low","title":"On Limitations of LLM as Annotator for Low Resource Languages","date":"2024-11-26","arxiv_id":"2411.17637","n_code_links":0,"syntology":null},{"paper":"/paper/pretrained-llm-adapted-with-lora-as-a","slug":"pretrained-llm-adapted-with-lora-as-a","title":"Pretrained LLM Adapted with LoRA as a Decision Transformer for Offline RL in Quantitative Trading","date":"2024-11-26","arxiv_id":"2411.17900","n_code_links":1,"syntology":null},{"paper":null,"slug":"adaptive-circuit-behavior-and-generalization","title":"Adaptive Circuit Behavior and Generalization in Mechanistic Interpretability","date":"2024-11-25","arxiv_id":"2411.16105","n_code_links":0,"syntology":null},{"paper":null,"slug":"are-transformers-truly-foundational-for","title":"Are Transformers Truly Foundational for Robotics?","date":"2024-11-25","arxiv_id":"2411.16917","n_code_links":0,"syntology":null},{"paper":null,"slug":"can-ai-grade-your-essays-a-comparative","title":"Can AI grade your essays? A comparative analysis of large language models and teacher ratings in multidimensional essay scoring","date":"2024-11-25","arxiv_id":"2411.16337","n_code_links":0,"syntology":null},{"paper":null,"slug":"fine-tuning-llms-with-noisy-data-for","title":"Fine-Tuning LLMs with Noisy Data for Political Argument Generation and Post Guidance","date":"2024-11-25","arxiv_id":"2411.16813","n_code_links":0,"syntology":null},{"paper":"/paper/marketgpt-developing-a-pre-trained","slug":"marketgpt-developing-a-pre-trained","title":"MarketGPT: Developing a Pre-trained transformer (GPT) for Modeling Financial Time Series","date":"2024-11-25","arxiv_id":"2411.16585","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":1,"n_instrument":2,"unverified":0,"pointer_only":3,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":{"repos":["aaron-wheeler/marketgpt"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"predictive-power-of-llms-in-financial-markets","title":"Predictive Power of LLMs in Financial Markets","date":"2024-11-25","arxiv_id":"2411.16569","n_code_links":0,"syntology":null},{"paper":"/paper/development-of-pre-trained-transformer-based","slug":"development-of-pre-trained-transformer-based","title":"Development of Pre-Trained Transformer-based Models for the Nepali Language","date":"2024-11-24","arxiv_id":"2411.15734","n_code_links":0,"syntology":null},{"paper":"/paper/all-that-glitters-approaches-to-evaluations","slug":"all-that-glitters-approaches-to-evaluations","title":"\"All that Glitters\": Approaches to Evaluations with Unreliable Model and Human Annotations","date":"2024-11-23","arxiv_id":"2411.15634","n_code_links":1,"syntology":null},{"paper":null,"slug":"chatbci-a-p300-speller-bci-leveraging-large","title":"ChatBCI: A P300 Speller BCI Leveraging Large Language Models for Improved Sentence Composition in Realistic Scenarios","date":"2024-11-23","arxiv_id":"2411.15395","n_code_links":0,"syntology":null},{"paper":null,"slug":"improving-next-tokens-via-second-last","title":"Improving Next Tokens via Second-Last Predictions with Generate and Refine","date":"2024-11-23","arxiv_id":"2411.15661","n_code_links":0,"syntology":null},{"paper":null,"slug":"comparative-analysis-of-pooling-mechanisms-in","title":"Comparative Analysis of Pooling Mechanisms in LLMs: A Sentiment Analysis Perspective","date":"2024-11-22","arxiv_id":"2411.14654","n_code_links":0,"syntology":null},{"paper":null,"slug":"assessment-of-llm-responses-to-end-user","title":"Assessment of LLM Responses to End-user Security Questions","date":"2024-11-21","arxiv_id":"2411.14571","n_code_links":0,"syntology":null},{"paper":"/paper/evaluating-the-robustness-of-analogical","slug":"evaluating-the-robustness-of-analogical","title":"Evaluating the Robustness of Analogical Reasoning in Large Language Models","date":"2024-11-21","arxiv_id":"2411.14215","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 1 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["marthaflinderslewis/robust-analogy"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"ai-driven-agents-with-prompts-designed-for","title":"AI-Driven Agents with Prompts Designed for High Agreeableness Increase the Likelihood of Being Mistaken for a Human in the Turing Test","date":"2024-11-20","arxiv_id":"2411.13749","n_code_links":0,"syntology":null},{"paper":"/paper/combining-autoregressive-and-autoencoder","slug":"combining-autoregressive-and-autoencoder","title":"Combining Autoregressive and Autoencoder Language Models for Text Classification","date":"2024-11-20","arxiv_id":"2411.13282","n_code_links":1,"syntology":null},{"paper":null,"slug":"exploring-large-language-models-for-climate","title":"Exploring Large Language Models for Climate Forecasting","date":"2024-11-20","arxiv_id":"2411.13724","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-combined-encoder-and-transformer-approach","title":"A Combined Encoder and Transformer Approach for Coherent and High-Quality Text Generation","date":"2024-11-19","arxiv_id":"2411.12157","n_code_links":0,"syntology":null},{"paper":null,"slug":"leveraging-virtual-reality-and-ai-tutoring","title":"Leveraging Virtual Reality and AI Tutoring for Language Learning: A Case Study of a Virtual Campus Environment with OpenAI GPT Integration with Unity 3D","date":"2024-11-19","arxiv_id":"2411.12619","n_code_links":0,"syntology":null},{"paper":null,"slug":"can-open-source-llms-enhance-data","title":"Can Open-source LLMs Enhance Data Synthesis for Toxic Detection?: An Experimental Study","date":"2024-11-18","arxiv_id":"2411.15175","n_code_links":0,"syntology":null},{"paper":null,"slug":"chapter-7-review-of-data-driven-generative-ai","title":"Chapter 7 Review of Data-Driven Generative AI Models for Knowledge Extraction from Scientific Literature in Healthcare","date":"2024-11-18","arxiv_id":"2411.11635","n_code_links":0,"syntology":null},{"paper":"/paper/cnmbert-a-model-for-hanyu-pinyin-abbreviation","slug":"cnmbert-a-model-for-hanyu-pinyin-abbreviation","title":"CNMBERT: A Model for Converting Hanyu Pinyin Abbreviations to Chinese Characters","date":"2024-11-18","arxiv_id":"2411.11770","n_code_links":1,"syntology":null},{"paper":"/paper/perfcodegen-improving-performance-of-llm","slug":"perfcodegen-improving-performance-of-llm","title":"PerfCodeGen: Improving Performance of LLM Generated Code with Execution Feedback","date":"2024-11-18","arxiv_id":"2412.03578","n_code_links":1,"syntology":{"ran":1,"of":3,"n_ran_checked":1,"n_instrument":0,"unverified":2,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["SalesforceAIResearch/perfcodegen"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"understanding-student-sentiment-on-mental","title":"Understanding Student Sentiment on Mental Health Support in Colleges Using Large Language Models","date":"2024-11-18","arxiv_id":"2412.04326","n_code_links":0,"syntology":null},{"paper":"/paper/versatune-fine-tuning-multi-ability-llms","slug":"versatune-fine-tuning-multi-ability-llms","title":"VersaTune: An Efficient Data Composition Framework for Training Multi-Capability LLMs","date":"2024-11-18","arxiv_id":"2411.11266","n_code_links":1,"syntology":null},{"paper":null,"slug":"does-prompt-formatting-have-any-impact-on-llm","title":"Does Prompt Formatting Have Any Impact on LLM Performance?","date":"2024-11-15","arxiv_id":"2411.10541","n_code_links":0,"syntology":null},{"paper":"/paper/mars-unleashing-the-power-of-variance","slug":"mars-unleashing-the-power-of-variance","title":"MARS: Unleashing the Power of Variance Reduction for Training Large Models","date":"2024-11-15","arxiv_id":"2411.10438","n_code_links":2,"syntology":{"ran":3,"of":5,"n_ran_checked":2,"n_instrument":1,"unverified":2,"pointer_only":1,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 1 honoured, 1 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","official":{"repos":["AGI-Arena/MARS"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":2,"ran_from_kinds":["official","unlocated"]}}},{"paper":null,"slug":"prompting-and-fine-tuning-large-language","title":"Prompting and Fine-tuning Large Language Models for Automated Code Review Comment Generation","date":"2024-11-15","arxiv_id":"2411.10129","n_code_links":0,"syntology":null},{"paper":null,"slug":"take-package-as-language-anomaly-detection","title":"Take Package as Language: Anomaly Detection Using Transformer","date":"2024-11-15","arxiv_id":"2412.04473","n_code_links":0,"syntology":null},{"paper":null,"slug":"babylm-challenge-exploring-the-effect-of","title":"BabyLM Challenge: Exploring the Effect of Variation Sets on Language Model Training Efficiency","date":"2024-11-14","arxiv_id":"2411.09587","n_code_links":0,"syntology":null},{"paper":null,"slug":"beyond-static-tools-evaluating-large-language","title":"Beyond Static Tools: Evaluating Large Language Models for Cryptographic Misuse Detection","date":"2024-11-14","arxiv_id":"2411.09772","n_code_links":0,"syntology":null},{"paper":null,"slug":"hategpt-unleashing-gpt-3-5-turbo-to-combat","title":"HateGPT: Unleashing GPT-3.5 Turbo to Combat Hate Speech on X","date":"2024-11-14","arxiv_id":"2411.09214","n_code_links":0,"syntology":null},{"paper":null,"slug":"llmstinger-jailbreaking-llms-using-rl-fine","title":"LLMStinger: Jailbreaking LLMs using RL fine-tuned LLMs","date":"2024-11-13","arxiv_id":"2411.08862","n_code_links":0,"syntology":null},{"paper":null,"slug":"responsible-ai-in-construction-safety","title":"Responsible AI in Construction Safety: Systematic Evaluation of Large Language Models and Prompt Engineering","date":"2024-11-13","arxiv_id":"2411.08320","n_code_links":0,"syntology":null},{"paper":null,"slug":"valtest-automated-validation-of-language","title":"VALTEST: Automated Validation of Language Model Generated Test Cases","date":"2024-11-13","arxiv_id":"2411.08254","n_code_links":0,"syntology":null},{"paper":"/paper/evaluating-chatgpt-3-5-efficiency-in-solving","slug":"evaluating-chatgpt-3-5-efficiency-in-solving","title":"Evaluating ChatGPT-3.5 Efficiency in Solving Coding Problems of Different Complexity Levels: An Empirical Analysis","date":"2024-11-12","arxiv_id":"2411.07529","n_code_links":1,"syntology":null},{"paper":"/paper/fair-summarization-bridging-quality-and","slug":"fair-summarization-bridging-quality-and","title":"Fair Summarization: Bridging Quality and Diversity in Extractive Summaries","date":"2024-11-12","arxiv_id":"2411.07521","n_code_links":1,"syntology":{"ran":4,"of":5,"n_ran_checked":1,"n_instrument":3,"unverified":1,"pointer_only":5,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 1 violated, 0 with no contract checked; 3 where Syntology's instrument failed) · 1 unverified","official":{"repos":["PortNLP/FairEXTSummarizer"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":["official","unlocated"]}}},{"paper":null,"slug":"llm-app-squatting-and-cloning","title":"LLM App Squatting and Cloning","date":"2024-11-12","arxiv_id":"2411.07518","n_code_links":0,"syntology":null},{"paper":null,"slug":"ambient-ai-scribing-support-comparing-the","title":"Ambient AI Scribing Support: Comparing the Performance of Specialized AI Agentic Architecture to Leading Foundational Models","date":"2024-11-11","arxiv_id":"2411.06713","n_code_links":0,"syntology":null},{"paper":"/paper/autonomous-droplet-microfluidic-design","slug":"autonomous-droplet-microfluidic-design","title":"Autonomous Droplet Microfluidic Design Framework with Large Language Models","date":"2024-11-11","arxiv_id":"2411.06691","n_code_links":1,"syntology":null},{"paper":null,"slug":"cancer-answer-empowering-cancer-care-with","title":"Cancer-Answer: Empowering Cancer Care with Advanced Large Language Models","date":"2024-11-11","arxiv_id":"2411.06946","n_code_links":0,"syntology":null},{"paper":null,"slug":"explore-the-reasoning-capability-of-llms-in","title":"Explore the Reasoning Capability of LLMs in the Chess Testbed","date":"2024-11-11","arxiv_id":"2411.06655","n_code_links":0,"syntology":null},{"paper":null,"slug":"on-active-privacy-auditing-in-supervised-fine","title":"On Active Privacy Auditing in Supervised Fine-tuning for White-Box Language Models","date":"2024-11-11","arxiv_id":"2411.07070","n_code_links":0,"syntology":null},{"paper":null,"slug":"spartan-a-sparse-transformer-learning-local","title":"SPARTAN: A Sparse Transformer Learning Local Causation","date":"2024-11-11","arxiv_id":"2411.06890","n_code_links":0,"syntology":null}],"record_sha256":"f5c70302712da2791d3f70701106d6ac186c4d09a4d1ddaf936955b9a4bf991b","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}