{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/weight-decay/papers/44","list_of":"/method/weight-decay","method":"Weight Decay","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":44,"pages_in_order":108,"rows_per_page":100,"rows":[4301,4400],"of":10713,"counts":{"archive_papers_tagged":10713,"with_a_code_link":4533,"where_syntology_ran_a_sample":1291,"not_listed_spam_title":0,"listed":10713,"listed_where_code_ran":1291,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":1064,"every_run_a_failure_of_syntologys_instrument":227,"listed_with_a_run_with_no_instrument_failure":1064,"listed_every_run_a_failure_of_syntologys_instrument":227,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/weight-decay","prev":"/method/weight-decay/papers/43","next":"/method/weight-decay/papers/45","papers":[{"paper":null,"slug":"an-evaluation-of-gpt-models-for-phenotype","title":"An evaluation of GPT models for phenotype concept recognition","date":"2023-09-29","arxiv_id":"2309.17169","n_code_links":0,"syntology":null},{"paper":"/paper/benchmarking-the-abilities-of-large-language","slug":"benchmarking-the-abilities-of-large-language","title":"Benchmarking the Abilities of Large Language Models for RDF Knowledge Graph Creation and Comprehension: How Well Do LLMs Speak Turtle?","date":"2023-09-29","arxiv_id":"2309.17122","n_code_links":3,"syntology":null},{"paper":"/paper/dyval-graph-informed-dynamic-evaluation-of","slug":"dyval-graph-informed-dynamic-evaluation-of","title":"DyVal: Dynamic Evaluation of Large Language Models for Reasoning Tasks","date":"2023-09-29","arxiv_id":"2309.17167","n_code_links":1,"syntology":null},{"paper":null,"slug":"intuitive-or-dependent-investigating-llms","title":"Intuitive or Dependent? Investigating LLMs' Behavior Style to Conflicting Prompts","date":"2023-09-29","arxiv_id":"2309.17415","n_code_links":0,"syntology":null},{"paper":"/paper/llm-deliberation-evaluating-llms-with","slug":"llm-deliberation-evaluating-llms-with","title":"Cooperation, Competition, and Maliciousness: LLM-Stakeholders Interactive Negotiation","date":"2023-09-29","arxiv_id":"2309.17234","n_code_links":2,"syntology":{"ran":3,"of":3,"n_ran_checked":1,"n_instrument":2,"unverified":0,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":{"repos":["s-abdelnabi/llm-deliberation"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"revolutionizing-mobile-interaction-enabling-a","title":"Revolutionizing Mobile Interaction: Enabling a 3 Billion Parameter GPT LLM on Mobile","date":"2023-09-29","arxiv_id":"2310.01434","n_code_links":0,"syntology":null},{"paper":null,"slug":"split-and-merge-aligning-position-biases-in","title":"Split and Merge: Aligning Position Biases in LLM-based Evaluators","date":"2023-09-29","arxiv_id":"2310.01432","n_code_links":0,"syntology":null},{"paper":null,"slug":"symmetry-leads-to-structured-constraint-of","title":"Symmetry Induces Structure and Constraint of Learning","date":"2023-09-29","arxiv_id":"2309.16932","n_code_links":0,"syntology":null},{"paper":null,"slug":"training-and-inference-of-large-language","title":"Training and inference of large language models using 8-bit floating point","date":"2023-09-29","arxiv_id":"2309.17224","n_code_links":0,"syntology":null},{"paper":null,"slug":"ae-gpt-using-large-language-models-to-extract","title":"AE-GPT: Using Large Language Models to Extract Adverse Events from Surveillance Reports-A Use Case with Influenza Vaccine Adverse Events","date":"2023-09-28","arxiv_id":"2309.16150","n_code_links":0,"syntology":null},{"paper":"/paper/gpt-fathom-benchmarking-large-language-models","slug":"gpt-fathom-benchmarking-large-language-models","title":"GPT-Fathom: Benchmarking Large Language Models to Decipher the Evolutionary Path towards GPT-4 and Beyond","date":"2023-09-28","arxiv_id":"2309.16583","n_code_links":1,"syntology":{"ran":5,"of":9,"n_ran_checked":5,"n_instrument":0,"unverified":4,"pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","official":{"repos":["gpt-fathom/gpt-fathom"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":"/paper/hallucination-reduction-in-long-input-text","slug":"hallucination-reduction-in-long-input-text","title":"Hallucination Reduction in Long Input Text Summarization","date":"2023-09-28","arxiv_id":"2309.16781","n_code_links":1,"syntology":null},{"paper":null,"slug":"large-language-model-soft-ideologization-via","title":"Large Language Model Soft Ideologization via AI-Self-Consciousness","date":"2023-09-28","arxiv_id":"2309.16167","n_code_links":0,"syntology":null},{"paper":null,"slug":"stress-testing-chain-of-thought-prompting-for","title":"Stress Testing Chain-of-Thought Prompting for Large Language Models","date":"2023-09-28","arxiv_id":"2309.16621","n_code_links":0,"syntology":null},{"paper":"/paper/maximum-weight-entropy","slug":"maximum-weight-entropy","title":"Deep Out-of-Distribution Uncertainty Quantification via Weight Entropy Maximization","date":"2023-09-27","arxiv_id":"2309.15704","n_code_links":1,"syntology":null},{"paper":null,"slug":"mededit-model-editing-for-medical-question","title":"MKRAG: Medical Knowledge Retrieval Augmented Generation for Medical Question Answering","date":"2023-09-27","arxiv_id":"2309.16035","n_code_links":0,"syntology":null},{"paper":"/paper/mindgpt-interpreting-what-you-see-with-non","slug":"mindgpt-interpreting-what-you-see-with-non","title":"MindGPT: Interpreting What You See with Non-invasive Brain Recordings","date":"2023-09-27","arxiv_id":"2309.15729","n_code_links":1,"syntology":{"ran":3,"of":7,"n_ran_checked":3,"n_instrument":0,"unverified":4,"pointer_only":7,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","official":{"repos":["jxuanc/mindgpt"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":"/paper/nlpbench-evaluating-large-language-models-on","slug":"nlpbench-evaluating-large-language-models-on","title":"NLPBench: Evaluating Large Language Models on Solving NLP Problems","date":"2023-09-27","arxiv_id":"2309.15630","n_code_links":1,"syntology":null},{"paper":null,"slug":"an-nlp-benchmark-dataset-for-assessing","title":"An NLP Benchmark Dataset for Assessing Corporate Climate Policy Engagement","date":"2023-09-26","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/capp-130-a-corpus-of-chinese-application","slug":"capp-130-a-corpus-of-chinese-application","title":"CAPP-130: A Corpus of Chinese Application Privacy Policy Summarization and Interpretation","date":"2023-09-26","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"comparative-analysis-of-artificial","title":"Legal Question-Answering in the Indian Context: Efficacy, Challenges, and Potential of Modern AI Models","date":"2023-09-26","arxiv_id":"2309.14735","n_code_links":0,"syntology":null},{"paper":null,"slug":"exploring-small-language-models-with-prompt","title":"Exploring Small Language Models with Prompt-Learning Paradigm for Efficient Domain-Specific Text Classification","date":"2023-09-26","arxiv_id":"2309.14779","n_code_links":0,"syntology":null},{"paper":"/paper/how-to-catch-an-ai-liar-lie-detection-in","slug":"how-to-catch-an-ai-liar-lie-detection-in","title":"How to Catch an AI Liar: Lie Detection in Black-Box LLMs by Asking Unrelated Questions","date":"2023-09-26","arxiv_id":"2309.15840","n_code_links":1,"syntology":{"ran":3,"of":4,"n_ran_checked":3,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["lorypack/llm-liedetector"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"low-rank-adaptation-of-large-language-model","title":"Low-rank Adaptation of Large Language Model Rescoring for Parameter-Efficient Speech Recognition","date":"2023-09-26","arxiv_id":"2309.15223","n_code_links":0,"syntology":null},{"paper":"/paper/ragas-automated-evaluation-of-retrieval","slug":"ragas-automated-evaluation-of-retrieval","title":"RAGAS: Automated Evaluation of Retrieval Augmented Generation","date":"2023-09-26","arxiv_id":"2309.15217","n_code_links":3,"syntology":{"ran":4,"of":6,"n_ran_checked":4,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":null}},{"paper":"/paper/rankvicuna-zero-shot-listwise-document","slug":"rankvicuna-zero-shot-listwise-document","title":"RankVicuna: Zero-Shot Listwise Document Reranking with Open-Source Large Language Models","date":"2023-09-26","arxiv_id":"2309.15088","n_code_links":3,"syntology":null},{"paper":"/paper/supersonic-learning-to-generate-source-code","slug":"supersonic-learning-to-generate-source-code","title":"Supersonic: Learning to Generate Source Code Optimizations in C/C++","date":"2023-09-26","arxiv_id":"2309.14846","n_code_links":1,"syntology":null},{"paper":null,"slug":"comprehensive-overview-of-named-entity","title":"Comprehensive Overview of Named Entity Recognition: Models, Domain-Specific Applications and Challenges","date":"2023-09-25","arxiv_id":"2309.14084","n_code_links":0,"syntology":null},{"paper":"/paper/enhancing-data-efficiency-in-reinforcement","slug":"enhancing-data-efficiency-in-reinforcement","title":"Enhancing data efficiency in reinforcement learning: a novel imagination mechanism based on mesh information propagation","date":"2023-09-25","arxiv_id":"2309.14243","n_code_links":2,"syntology":null},{"paper":null,"slug":"evaluating-cognitive-maps-and-planning-in","title":"Evaluating Cognitive Maps and Planning in Large Language Models with CogEval","date":"2023-09-25","arxiv_id":"2309.15129","n_code_links":0,"syntology":null},{"paper":"/paper/loggpt-log-anomaly-detection-via-gpt","slug":"loggpt-log-anomaly-detection-via-gpt","title":"LogGPT: Log Anomaly Detection via GPT","date":"2023-09-25","arxiv_id":"2309.14482","n_code_links":1,"syntology":null},{"paper":null,"slug":"watch-your-language-large-language-models-and","title":"Watch Your Language: Investigating Content Moderation with Large Language Models","date":"2023-09-25","arxiv_id":"2309.14517","n_code_links":0,"syntology":null},{"paper":null,"slug":"accelerating-large-batch-training-via","title":"Accelerating Large Batch Training via Gradient Signal to Noise Ratio (GSNR)","date":"2023-09-24","arxiv_id":"2309.13681","n_code_links":0,"syntology":null},{"paper":null,"slug":"does-the-most-sinfully-decadent-cake-ever","title":"Does the \"most sinfully decadent cake ever\" taste good? Answering Yes/No Questions from Figurative Contexts","date":"2023-09-24","arxiv_id":"2309.13748","n_code_links":0,"syntology":null},{"paper":"/paper/seeing-is-not-always-believing-invisible","slug":"seeing-is-not-always-believing-invisible","title":"Seeing Is Not Always Believing: Invisible Collision Attack and Defence on Pre-Trained Models","date":"2023-09-24","arxiv_id":"2309.13579","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-chat-about-boring-problems-studying-gpt","title":"A Chat About Boring Problems: Studying GPT-based text normalization","date":"2023-09-23","arxiv_id":"2309.13426","n_code_links":0,"syntology":null},{"paper":null,"slug":"exploring-large-language-models-cognitive","title":"Probing the Moral Development of Large Language Models through Defining Issues Test","date":"2023-09-23","arxiv_id":"2309.13356","n_code_links":0,"syntology":null},{"paper":"/paper/lexical-squad-multimodal-hate-speech-event","slug":"lexical-squad-multimodal-hate-speech-event","title":"Lexical Squad@Multimodal Hate Speech Event Detection 2023: Multimodal Hate Speech Detection using Fused Ensemble Approach","date":"2023-09-23","arxiv_id":"2309.13354","n_code_links":1,"syntology":null},{"paper":"/paper/amplify-attention-based-mixup-for-performance","slug":"amplify-attention-based-mixup-for-performance","title":"AMPLIFY:Attention-based Mixup for Performance Improvement and Label Smoothing in Transformer","date":"2023-09-22","arxiv_id":"2309.12689","n_code_links":1,"syntology":null},{"paper":null,"slug":"benllmeval-a-comprehensive-evaluation-into","title":"BenLLMEval: A Comprehensive Evaluation into the Potentials and Pitfalls of Large Language Models on Bengali NLP","date":"2023-09-22","arxiv_id":"2309.13173","n_code_links":0,"syntology":null},{"paper":null,"slug":"contextual-emotion-estimation-from-image","title":"Contextual Emotion Estimation from Image Captions","date":"2023-09-22","arxiv_id":"2309.13136","n_code_links":0,"syntology":null},{"paper":"/paper/large-language-models-and-control-mechanisms","slug":"large-language-models-and-control-mechanisms","title":"Investigating Large Language Models and Control Mechanisms to Improve Text Readability of Biomedical Abstracts","date":"2023-09-22","arxiv_id":"2309.13202","n_code_links":1,"syntology":null},{"paper":null,"slug":"large-language-models-are-also-good","title":"Large Language Models Are Also Good Prototypical Commonsense Reasoners","date":"2023-09-22","arxiv_id":"2309.13165","n_code_links":0,"syntology":null},{"paper":null,"slug":"spion-layer-wise-sparse-training-of","title":"SPION: Layer-Wise Sparse Training of Transformer via Convolutional Flood Filling","date":"2023-09-22","arxiv_id":"2309.12578","n_code_links":0,"syntology":null},{"paper":"/paper/toproberta-topology-aware-authorship","slug":"toproberta-topology-aware-authorship","title":"TOPFORMER: Topology-Aware Authorship Attribution of Deepfake Texts with Diverse Writing Styles","date":"2023-09-22","arxiv_id":"2309.12934","n_code_links":1,"syntology":null},{"paper":"/paper/a-chinese-prompt-attack-dataset-for-llms-with","slug":"a-chinese-prompt-attack-dataset-for-llms-with","title":"Goal-Oriented Prompt Attack and Safety Evaluation for LLMs","date":"2023-09-21","arxiv_id":"2309.11830","n_code_links":2,"syntology":null},{"paper":null,"slug":"a-long-n-step-surrogate-stage-reward-for-deep","title":"A Long $N$-step Surrogate Stage Reward for Deep Reinforcement Learning","date":"2023-09-21","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/bad-actor-good-advisor-exploring-the-role-of","slug":"bad-actor-good-advisor-exploring-the-role-of","title":"Bad Actor, Good Advisor: Exploring the Role of Large Language Models in Fake News Detection","date":"2023-09-21","arxiv_id":"2309.12247","n_code_links":1,"syntology":null},{"paper":"/paper/bayestune-bayesian-sparse-deep-model-fine","slug":"bayestune-bayesian-sparse-deep-model-fine","title":"BayesTune: Bayesian Sparse Deep Model Fine-tuning","date":"2023-09-21","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"constraints-first-a-new-mdd-based-model-to","title":"Constraints First: A New MDD-based Model to Generate Sentences Under Constraints","date":"2023-09-21","arxiv_id":"2309.12415","n_code_links":0,"syntology":null},{"paper":"/paper/deconstructing-data-reconstruction-multiclass","slug":"deconstructing-data-reconstruction-multiclass","title":"Deconstructing Data Reconstruction: Multiclass, Weight Decay and General Losses","date":"2023-09-21","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/double-gumbel-q-learning","slug":"double-gumbel-q-learning","title":"Double Gumbel Q-Learning","date":"2023-09-21","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/fednar-federated-optimization-with-normalized","slug":"fednar-federated-optimization-with-normalized","title":"FedNAR: Federated Optimization with Normalized Annealing Regularization","date":"2023-09-21","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/implicit-differentiable-outlier-detection","slug":"implicit-differentiable-outlier-detection","title":"Implicit Differentiable Outlier Detection Enable Robust Deep Multimodal Analysis","date":"2023-09-21","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/making-scalable-meta-learning-practical","slug":"making-scalable-meta-learning-practical","title":"Making Scalable Meta Learning Practical","date":"2023-09-21","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/marich-a-query-efficient-distributionally-1","slug":"marich-a-query-efficient-distributionally-1","title":"Marich: A Query-efficient Distributionally Equivalent Model Extraction Attack","date":"2023-09-21","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/metamath-bootstrap-your-own-mathematical","slug":"metamath-bootstrap-your-own-mathematical","title":"MetaMath: Bootstrap Your Own Mathematical Questions for Large Language Models","date":"2023-09-21","arxiv_id":"2309.12284","n_code_links":1,"syntology":{"ran":15,"of":22,"n_ran_checked":1,"n_instrument":14,"unverified":7,"pointer_only":0,"phrase":"15 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 1 violated, 0 with no contract checked; 14 where Syntology's instrument failed) · 7 unverified","official":{"repos":["meta-math/MetaMath"],"state":"official (archive's flag): 15 ran","n_ran":15,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":7,"ran_from_kinds":["official"]}}},{"paper":"/paper/on-the-relationship-between-skill-neurons-and","slug":"on-the-relationship-between-skill-neurons-and","title":"On the Relationship between Skill Neurons and Robustness in Prompt Tuning","date":"2023-09-21","arxiv_id":"2309.12263","n_code_links":1,"syntology":null},{"paper":null,"slug":"re-exploring-the-role-of-grammar-and-word","title":"[Re] Exploring the Role of Grammar and Word Choice in Bias Toward African American English (AAE) in Hate Speech Classification","date":"2023-09-21","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"safe-hierarchical-reinforcement-learning-for","title":"Safe Hierarchical Reinforcement Learning for CubeSat Task Scheduling Based on Energy Consumption","date":"2023-09-21","arxiv_id":"2309.12004","n_code_links":0,"syntology":null},{"paper":null,"slug":"slhcat-mapping-wikipedia-categories-and-lists","title":"SLHCat: Mapping Wikipedia Categories and Lists to DBpedia by Leveraging Semantic, Lexical, and Hierarchical Features","date":"2023-09-21","arxiv_id":"2309.11791","n_code_links":0,"syntology":null},{"paper":null,"slug":"spiced-news-similarity-detection-dataset-with","title":"SPICED: News Similarity Detection Dataset with Multiple Topics and Complexity Levels","date":"2023-09-21","arxiv_id":"2309.13080","n_code_links":0,"syntology":null},{"paper":null,"slug":"stock-market-sentiment-classification-and","title":"Stock Market Sentiment Classification and Backtesting via Fine-tuned BERT","date":"2023-09-21","arxiv_id":"2309.11979","n_code_links":0,"syntology":null},{"paper":"/paper/tart-a-plug-and-play-transformer-module-for","slug":"tart-a-plug-and-play-transformer-module-for","title":"TART: A plug-and-play Transformer module for task-agnostic reasoning","date":"2023-09-21","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/the-cambridge-law-corpus-a-corpus-for-legal-1","slug":"the-cambridge-law-corpus-a-corpus-for-legal-1","title":"The Cambridge Law Corpus: A Dataset for Legal AI Research","date":"2023-09-21","arxiv_id":"2309.12269","n_code_links":0,"syntology":null},{"paper":"/paper/the-reversal-curse-llms-trained-on-a-is-b","slug":"the-reversal-curse-llms-trained-on-a-is-b","title":"The Reversal Curse: LLMs trained on \"A is B\" fail to learn \"B is A\"","date":"2023-09-21","arxiv_id":"2309.12288","n_code_links":2,"syntology":{"ran":9,"of":11,"n_ran_checked":9,"n_instrument":0,"unverified":2,"pointer_only":11,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 9 with no instrument failure: 0 honoured, 0 violated, 9 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["lukasberglund/reversal_curse"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":9,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"toa-task-oriented-active-vqa","title":"TOA: Task-oriented Active VQA","date":"2023-09-21","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"towards-efficient-pre-trained-language-model","title":"Towards Efficient Pre-Trained Language Model via Feature Correlation Distillation","date":"2023-09-21","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/a-paradigm-shift-in-machine-translation","slug":"a-paradigm-shift-in-machine-translation","title":"A Paradigm Shift in Machine Translation: Boosting Translation Performance of Large Language Models","date":"2023-09-20","arxiv_id":"2309.11674","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":0,"n_instrument":3,"unverified":0,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 3 where Syntology's instrument failed) · 0 unverified","official":{"repos":["fe1ixxu/alma"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"ai-driven-patient-monitoring-with-multi-agent","title":"Adaptive Multi-Agent Deep Reinforcement Learning for Timely Healthcare Interventions","date":"2023-09-20","arxiv_id":"2309.10980","n_code_links":0,"syntology":null},{"paper":null,"slug":"attentionmix-data-augmentation-method-that","title":"AttentionMix: Data augmentation method that relies on BERT attention mechanism","date":"2023-09-20","arxiv_id":"2309.11104","n_code_links":0,"syntology":null},{"paper":"/paper/controlled-generation-with-prompt-insertion","slug":"controlled-generation-with-prompt-insertion","title":"Controlled Generation with Prompt Insertion for Natural Language Explanations in Grammatical Error Correction","date":"2023-09-20","arxiv_id":"2309.11439","n_code_links":1,"syntology":null},{"paper":"/paper/cot-bert-enhancing-unsupervised-sentence","slug":"cot-bert-enhancing-unsupervised-sentence","title":"CoT-BERT: Enhancing Unsupervised Sentence Representation through Chain-of-Thought","date":"2023-09-20","arxiv_id":"2309.11143","n_code_links":2,"syntology":null},{"paper":"/paper/design-of-chain-of-thought-in-math-problem","slug":"design-of-chain-of-thought-in-math-problem","title":"Design of Chain-of-Thought in Math Problem Solving","date":"2023-09-20","arxiv_id":"2309.11054","n_code_links":1,"syntology":null},{"paper":null,"slug":"fictional-worlds-real-connections-developing","title":"Fictional Worlds, Real Connections: Developing Community Storytelling Social Chatbots through LLMs","date":"2023-09-20","arxiv_id":"2309.11478","n_code_links":0,"syntology":null},{"paper":null,"slug":"generative-ai-in-mafia-like-game-simulation","title":"Generative AI in Mafia-like Game Simulation","date":"2023-09-20","arxiv_id":"2309.11672","n_code_links":0,"syntology":null},{"paper":null,"slug":"gpt-molberta-gpt-molecular-features-language","title":"GPT-MolBERTa: GPT Molecular Features Language Model for molecular property prediction","date":"2023-09-20","arxiv_id":"2310.03030","n_code_links":0,"syntology":null},{"paper":"/paper/safurai-001-new-qualitative-approach-for-code","slug":"safurai-001-new-qualitative-approach-for-code","title":"Safurai 001: New Qualitative Approach for Code LLM Evaluation","date":"2023-09-20","arxiv_id":"2309.11385","n_code_links":1,"syntology":null},{"paper":"/paper/sequence-to-sequence-spanish-pre-trained","slug":"sequence-to-sequence-spanish-pre-trained","title":"Sequence-to-Sequence Spanish Pre-trained Language Models","date":"2023-09-20","arxiv_id":"2309.11259","n_code_links":1,"syntology":null},{"paper":"/paper/the-languini-kitchen-enabling-language","slug":"the-languini-kitchen-enabling-language","title":"The Languini Kitchen: Enabling Language Modelling Research at Different Scales of Compute","date":"2023-09-20","arxiv_id":"2309.11197","n_code_links":1,"syntology":null},{"paper":null,"slug":"language-as-the-medium-multimodal-video","title":"Language as the Medium: Multimodal Video Classification through text only","date":"2023-09-19","arxiv_id":"2309.10783","n_code_links":0,"syntology":null},{"paper":null,"slug":"mixed-distil-bert-code-mixed-language","title":"Mixed-Distil-BERT: Code-mixed Language Modeling for Bangla, English, and Hindi","date":"2023-09-19","arxiv_id":"2309.10272","n_code_links":0,"syntology":null},{"paper":null,"slug":"rigorously-assessing-natural-language","title":"Rigorously Assessing Natural Language Explanations of Neurons","date":"2023-09-19","arxiv_id":"2309.10312","n_code_links":0,"syntology":null},{"paper":null,"slug":"writer-defined-ai-personas-for-on-demand","title":"Writer-Defined AI Personas for On-Demand Feedback Generation","date":"2023-09-19","arxiv_id":"2309.10433","n_code_links":0,"syntology":null},{"paper":null,"slug":"evaluation-of-gpt-3-for-anti-cancer-drug","title":"Evaluation of GPT-3 for Anti-Cancer Drug Sensitivity Prediction","date":"2023-09-18","arxiv_id":"2309.10016","n_code_links":0,"syntology":null},{"paper":"/paper/facilitating-nsfw-text-detection-in-open","slug":"facilitating-nsfw-text-detection-in-open","title":"Facilitating NSFW Text Detection in Open-Domain Dialogue Systems via Knowledge Distillation","date":"2023-09-18","arxiv_id":"2309.09749","n_code_links":1,"syntology":null},{"paper":null,"slug":"proposition-from-the-perspective-of-chinese","title":"Proposition from the Perspective of Chinese Language: A Chinese Proposition Classification Evaluation Benchmark","date":"2023-09-18","arxiv_id":"2309.09602","n_code_links":0,"syntology":null},{"paper":"/paper/recap-retrieval-augmented-audio-captioning","slug":"recap-retrieval-augmented-audio-captioning","title":"RECAP: Retrieval-Augmented Audio Captioning","date":"2023-09-18","arxiv_id":"2309.09836","n_code_links":1,"syntology":{"ran":5,"of":8,"n_ran_checked":5,"n_instrument":0,"unverified":3,"pointer_only":8,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":{"repos":["sreyan88/recap"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"towards-ontology-construction-with-language","title":"Towards Ontology Construction with Language Models","date":"2023-09-18","arxiv_id":"2309.09898","n_code_links":0,"syntology":null},{"paper":null,"slug":"contrastive-decoding-improves-reasoning-in","title":"Contrastive Decoding Improves Reasoning in Large Language Models","date":"2023-09-17","arxiv_id":"2309.09117","n_code_links":0,"syntology":null},{"paper":"/paper/detecting-covariate-drift-in-text-data-using","slug":"detecting-covariate-drift-in-text-data-using","title":"Detecting covariate drift in text data using document embeddings and dimensionality reduction","date":"2023-09-17","arxiv_id":"2309.10000","n_code_links":1,"syntology":null},{"paper":null,"slug":"do-large-gpt-models-discover-moral-dimensions","title":"Do Large GPT Models Discover Moral Dimensions in Language Representations? A Topological Study Of Sentence Embeddings","date":"2023-09-17","arxiv_id":"2309.09397","n_code_links":0,"syntology":null},{"paper":null,"slug":"from-cooking-recipes-to-robot-task-trees","title":"From Cooking Recipes to Robot Task Trees -- Improving Planning Correctness and Task Efficiency by Leveraging LLMs with a Knowledge Network","date":"2023-09-17","arxiv_id":"2309.09181","n_code_links":0,"syntology":null},{"paper":null,"slug":"decoder-only-architecture-for-speech","title":"Decoder-only Architecture for Speech Recognition with CTC Prompts and Text Data Augmentation","date":"2023-09-16","arxiv_id":"2309.08876","n_code_links":0,"syntology":null},{"paper":"/paper/has-sentiment-returned-to-the-pre-pandemic","slug":"has-sentiment-returned-to-the-pre-pandemic","title":"Has Sentiment Returned to the Pre-pandemic Level? A Sentiment Analysis Using U.S. College Subreddit Data from 2019 to 2022","date":"2023-09-16","arxiv_id":"2309.08845","n_code_links":1,"syntology":null},{"paper":"/paper/struc-bench-are-large-language-models-really","slug":"struc-bench-are-large-language-models-really","title":"Struc-Bench: Are Large Language Models Really Good at Generating Complex Structured Data?","date":"2023-09-16","arxiv_id":"2309.08963","n_code_links":1,"syntology":null},{"paper":"/paper/a-modern-turkish-poet-fine-tuned-gpt-2","slug":"a-modern-turkish-poet-fine-tuned-gpt-2","title":"A Modern Turkish Poet: Fine-Tuned GPT-2","date":"2023-09-15","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/advancing-the-evaluation-of-traditional","slug":"advancing-the-evaluation-of-traditional","title":"Advancing the Evaluation of Traditional Chinese Language Models: Towards a Comprehensive Benchmark Suite","date":"2023-09-15","arxiv_id":"2309.08448","n_code_links":1,"syntology":null},{"paper":"/paper/albner-a-corpus-for-named-entity-recognition","slug":"albner-a-corpus-for-named-entity-recognition","title":"AlbNER: A Corpus for Named Entity Recognition in Albanian","date":"2023-09-15","arxiv_id":"2309.08741","n_code_links":0,"syntology":null},{"paper":"/paper/casteist-but-not-racist-quantifying","slug":"casteist-but-not-racist-quantifying","title":"Indian-BhED: A Dataset for Measuring India-Centric Biases in Large Language Models","date":"2023-09-15","arxiv_id":"2309.08573","n_code_links":1,"syntology":null}],"record_sha256":"c81c287b487e26375797f1e7f83d945fb9d600df76b13fd7f1f46bed09d93ab8","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}