{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/weight-decay/papers/38","list_of":"/method/weight-decay","method":"Weight Decay","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":38,"pages_in_order":108,"rows_per_page":100,"rows":[3701,3800],"of":10713,"counts":{"archive_papers_tagged":10713,"with_a_code_link":4533,"where_syntology_ran_a_sample":1291,"not_listed_spam_title":0,"listed":10713,"listed_where_code_ran":1291,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":1064,"every_run_a_failure_of_syntologys_instrument":227,"listed_with_a_run_with_no_instrument_failure":1064,"listed_every_run_a_failure_of_syntologys_instrument":227,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/weight-decay","prev":"/method/weight-decay/papers/37","next":"/method/weight-decay/papers/39","papers":[{"paper":"/paper/text2mdt-extracting-medical-decision-trees","slug":"text2mdt-extracting-medical-decision-trees","title":"Text2MDT: Extracting Medical Decision Trees from Medical Texts","date":"2024-01-04","arxiv_id":"2401.02034","n_code_links":1,"syntology":null},{"paper":"/paper/a-first-look-at-information-highlighting-in","slug":"a-first-look-at-information-highlighting-in","title":"Studying and Recommending Information Highlighting in Stack Overflow Answers","date":"2024-01-03","arxiv_id":"2401.01472","n_code_links":1,"syntology":null},{"paper":null,"slug":"enhancing-multilingual-information-retrieval","title":"Enhancing Multilingual Information Retrieval in Mixed Human Resources Environments: A RAG Model Implementation for Multicultural Enterprise","date":"2024-01-03","arxiv_id":"2401.01511","n_code_links":0,"syntology":null},{"paper":null,"slug":"iot-in-the-era-of-generative-ai-vision-and","title":"The Internet of Things in the Era of Generative AI: Vision and Challenges","date":"2024-01-03","arxiv_id":"2401.01923","n_code_links":0,"syntology":null},{"paper":null,"slug":"iterative-mask-filling-an-effective-text","title":"Iterative Mask Filling: An Effective Text Augmentation Method Using Masked Language Modeling","date":"2024-01-03","arxiv_id":"2401.01830","n_code_links":0,"syntology":null},{"paper":null,"slug":"mlps-compass-what-is-learned-when-mlps-are","title":"MLPs Compass: What is learned when MLPs are combined with PLMs?","date":"2024-01-03","arxiv_id":"2401.01667","n_code_links":0,"syntology":null},{"paper":null,"slug":"natural-language-processing-and-multimodal","title":"Natural Language Processing and Multimodal Stock Price Prediction","date":"2024-01-03","arxiv_id":"2401.01487","n_code_links":0,"syntology":null},{"paper":"/paper/revisiting-zero-shot-abstractive","slug":"revisiting-zero-shot-abstractive","title":"Revisiting Zero-Shot Abstractive Summarization in the Era of Large Language Models from the Perspective of Position Bias","date":"2024-01-03","arxiv_id":"2401.01989","n_code_links":1,"syntology":null},{"paper":"/paper/vietnamese-poem-generation-the-prospect-of","slug":"vietnamese-poem-generation-the-prospect-of","title":"Vietnamese Poem Generation & The Prospect Of Cross-Language Poem-To-Poem Translation","date":"2024-01-02","arxiv_id":"2401.01078","n_code_links":1,"syntology":null},{"paper":"/paper/a-b-b-a-triggering-logical-reasoning-failures","slug":"a-b-b-a-triggering-logical-reasoning-failures","title":"LogicAsker: Evaluating and Improving the Logical Reasoning Ability of Large Language Models","date":"2024-01-01","arxiv_id":"2401.00757","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":3,"n_instrument":0,"unverified":0,"pointer_only":3,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["yxwan123/logicasker"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/a-computational-framework-for-behavioral","slug":"a-computational-framework-for-behavioral","title":"A Computational Framework for Behavioral Assessment of LLM Therapists","date":"2024-01-01","arxiv_id":"2401.00820","n_code_links":1,"syntology":null},{"paper":"/paper/adapt-or-perish-adaptive-sparse-transformer","slug":"adapt-or-perish-adaptive-sparse-transformer","title":"Adapt or Perish: Adaptive Sparse Transformer with Attentive Feature Refinement for Image Restoration","date":"2024-01-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"large-language-models-aren-t-all-that-you","title":"Large Language Models aren't all that you need","date":"2024-01-01","arxiv_id":"2401.00698","n_code_links":0,"syntology":null},{"paper":"/paper/revisiting-counterfactual-problems-in","slug":"revisiting-counterfactual-problems-in","title":"Revisiting Counterfactual Problems in Referring Expression Comprehension","date":"2024-01-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/seed-bench-benchmarking-multimodal-large","slug":"seed-bench-benchmarking-multimodal-large","title":"SEED-Bench: Benchmarking Multimodal Large Language Models","date":"2024-01-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"an-analysis-of-embedding-layers-and","title":"An Analysis of Embedding Layers and Similarity Scores using Siamese Neural Networks","date":"2023-12-31","arxiv_id":"2401.00582","n_code_links":0,"syntology":null},{"paper":"/paper/generative-model-driven-synthetic-training","slug":"generative-model-driven-synthetic-training","title":"Generative Model-Driven Synthetic Training Image Generation: An Approach to Cognition in Rail Defect Detection","date":"2023-12-31","arxiv_id":"2401.00393","n_code_links":1,"syntology":null},{"paper":"/paper/ragtruth-a-hallucination-corpus-for","slug":"ragtruth-a-hallucination-corpus-for","title":"RAGTruth: A Hallucination Corpus for Developing Trustworthy Retrieval-Augmented Language Models","date":"2023-12-31","arxiv_id":"2401.00396","n_code_links":3,"syntology":{"ran":3,"of":5,"n_ran_checked":3,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["particlemedia/ragtruth"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/advancing-ttp-analysis-harnessing-the-power","slug":"advancing-ttp-analysis-harnessing-the-power","title":"Advancing TTP Analysis: Harnessing the Power of Large Language Models with Retrieval Augmented Generation","date":"2023-12-30","arxiv_id":"2401.00280","n_code_links":1,"syntology":null},{"paper":null,"slug":"trace-and-edit-relation-associations-in-gpt","title":"Trace and Edit Relation Associations in GPT","date":"2023-12-30","arxiv_id":"2401.02976","n_code_links":0,"syntology":null},{"paper":"/paper/why-is-the-user-interface-a-dark-pattern","slug":"why-is-the-user-interface-a-dark-pattern","title":"Why is the User Interface a Dark Pattern? : Explainable Auto-Detection and its Analysis","date":"2023-12-30","arxiv_id":"2401.04119","n_code_links":1,"syntology":null},{"paper":"/paper/gemini-in-reasoning-unveiling-commonsense-in","slug":"gemini-in-reasoning-unveiling-commonsense-in","title":"Gemini in Reasoning: Unveiling Commonsense in Multimodal Large Language Models","date":"2023-12-29","arxiv_id":"2312.17661","n_code_links":1,"syntology":null},{"paper":"/paper/jatmo-prompt-injection-defense-by-task","slug":"jatmo-prompt-injection-defense-by-task","title":"Jatmo: Prompt Injection Defense by Task-Specific Finetuning","date":"2023-12-29","arxiv_id":"2312.17673","n_code_links":1,"syntology":{"ran":10,"of":18,"n_ran_checked":10,"n_instrument":0,"unverified":8,"pointer_only":18,"phrase":"10 ran (of which 0 constructed an object rather than computing a result; 10 with no instrument failure: 0 honoured, 0 violated, 10 with no contract checked; 0 where Syntology's instrument failed) · 8 unverified","official":{"repos":["wagner-group/prompt-injection-defense"],"state":"official (archive's flag): 10 ran","n_ran":10,"n_constructed":0,"n_ran_no_instrument_failure":10,"n_unverified":8,"ran_from_kinds":["official"]}}},{"paper":"/paper/mosaicbert-a-bidirectional-encoder-optimized-1","slug":"mosaicbert-a-bidirectional-encoder-optimized-1","title":"MosaicBERT: A Bidirectional Encoder Optimized for Fast Pretraining","date":"2023-12-29","arxiv_id":"2312.17482","n_code_links":1,"syntology":null},{"paper":"/paper/tupy-e-detecting-hate-speech-in-brazilian","slug":"tupy-e-detecting-hate-speech-in-brazilian","title":"TuPy-E: detecting hate speech in Brazilian Portuguese social media with a novel dataset and comprehensive analysis of models","date":"2023-12-29","arxiv_id":"2312.17704","n_code_links":1,"syntology":null},{"paper":null,"slug":"evaluating-the-performance-of-large-language-1","title":"Evaluating the Performance of Large Language Models for Spanish Language in Undergraduate Admissions Exams","date":"2023-12-28","arxiv_id":"2312.16845","n_code_links":0,"syntology":null},{"paper":null,"slug":"language-model-as-an-annotator-unsupervised","title":"Language Model as an Annotator: Unsupervised Context-aware Quality Phrase Generation","date":"2023-12-28","arxiv_id":"2312.17349","n_code_links":0,"syntology":null},{"paper":null,"slug":"large-language-model-for-causal-decision","title":"LLM4Causal: Democratized Causal Tools for Everyone via Large Language Model","date":"2023-12-28","arxiv_id":"2312.17122","n_code_links":0,"syntology":null},{"paper":"/paper/sentinellms-encrypted-input-adaptation-and","slug":"sentinellms-encrypted-input-adaptation-and","title":"SentinelLMs: Encrypted Input Adaptation and Fine-tuning of Language Models for Private and Secure Inference","date":"2023-12-28","arxiv_id":"2312.17342","n_code_links":1,"syntology":null},{"paper":null,"slug":"pangu-p-enhancing-language-model","title":"PanGu-$π$: Enhancing Language Model Architectures via Nonlinearity Compensation","date":"2023-12-27","arxiv_id":"2312.17276","n_code_links":0,"syntology":null},{"paper":null,"slug":"relationship-between-auditory-and-semantic","title":"Relationship between auditory and semantic entrainment using Deep Neural Networks (DNN)","date":"2023-12-27","arxiv_id":"2312.16599","n_code_links":0,"syntology":null},{"paper":null,"slug":"chartbench-a-benchmark-for-complex-visual","title":"ChartBench: A Benchmark for Complex Visual Reasoning in Charts","date":"2023-12-26","arxiv_id":"2312.15915","n_code_links":0,"syntology":null},{"paper":"/paper/principled-instructions-are-all-you-need-for","slug":"principled-instructions-are-all-you-need-for","title":"Principled Instructions Are All You Need for Questioning LLaMA-1/2, GPT-3.5/4","date":"2023-12-26","arxiv_id":"2312.16171","n_code_links":2,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["vila-lab/atlas"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/secqa-a-concise-question-answering-dataset","slug":"secqa-a-concise-question-answering-dataset","title":"SecQA: A Concise Question-Answering Dataset for Evaluating Large Language Models in Computer Security","date":"2023-12-26","arxiv_id":"2312.15838","n_code_links":1,"syntology":{"ran":6,"of":7,"n_ran_checked":5,"n_instrument":1,"unverified":1,"pointer_only":1,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":{"repos":["zefang-liu/lm-evaluation-harness"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"task-contamination-language-models-may-not-be","title":"Task Contamination: Language Models May Not Be Few-Shot Anymore","date":"2023-12-26","arxiv_id":"2312.16337","n_code_links":0,"syntology":null},{"paper":null,"slug":"compositional-generalization-in-spoken","title":"Compositional Generalization in Spoken Language Understanding","date":"2023-12-25","arxiv_id":"2312.15815","n_code_links":0,"syntology":null},{"paper":"/paper/fairness-aware-structured-pruning-in","slug":"fairness-aware-structured-pruning-in","title":"Fairness-Aware Structured Pruning in Transformers","date":"2023-12-24","arxiv_id":"2312.15398","n_code_links":1,"syntology":{"ran":0,"of":2,"n_ran_checked":0,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"0 ran · 2 unverified","official":{"repos":["chandar-lab/fasp"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":2,"ran_from_kinds":[]}}},{"paper":null,"slug":"multi-level-biomedical-ner-through-multi","title":"Multi-level biomedical NER through multi-granularity embeddings and enhanced labeling","date":"2023-12-24","arxiv_id":"2312.15550","n_code_links":0,"syntology":null},{"paper":null,"slug":"do-llm-agents-exhibit-social-behavior","title":"Do LLM Agents Exhibit Social Behavior?","date":"2023-12-23","arxiv_id":"2312.15198","n_code_links":0,"syntology":null},{"paper":"/paper/understanding-the-potential-of-fpga-based","slug":"understanding-the-potential-of-fpga-based","title":"Understanding the Potential of FPGA-Based Spatial Acceleration for Large Language Model Inference","date":"2023-12-23","arxiv_id":"2312.15159","n_code_links":1,"syntology":null},{"paper":null,"slug":"efficacy-of-machine-generated-instructions","title":"Efficacy of Machine-Generated Instructions","date":"2023-12-22","arxiv_id":"2312.14423","n_code_links":0,"syntology":null},{"paper":null,"slug":"fm-ov3d-foundation-model-based-cross-modal","title":"FM-OV3D: Foundation Model-based Cross-modal Knowledge Blending for Open-Vocabulary 3D Detection","date":"2023-12-22","arxiv_id":"2312.14465","n_code_links":0,"syntology":null},{"paper":null,"slug":"multiagent-copilot-approach-for-shared","title":"Multiagent Copilot Approach for Shared Autonomy between Human EEG and TD3 Deep Reinforcement Learning","date":"2023-12-22","arxiv_id":"2312.14458","n_code_links":0,"syntology":null},{"paper":"/paper/refining-gpt-3-embeddings-with-a-siamese","slug":"refining-gpt-3-embeddings-with-a-siamese","title":"Refining GPT-3 Embeddings with a Siamese Structure for Technical Post Duplicate Detection","date":"2023-12-22","arxiv_id":"2312.15068","n_code_links":1,"syntology":null},{"paper":null,"slug":"towards-detecting-cascades-of-biased-medical","title":"Towards Detecting Cascades of Biased Medical Claims on Twitter","date":"2023-12-22","arxiv_id":"2312.15040","n_code_links":0,"syntology":null},{"paper":null,"slug":"understanding-the-regularity-of-self","title":"How Smooth Is Attention?","date":"2023-12-22","arxiv_id":"2312.14820","n_code_links":0,"syntology":null},{"paper":null,"slug":"unsupervised-auditory-and-semantic","title":"Unsupervised Auditory and Semantic Entrainment Models with Deep Neural Networks","date":"2023-12-22","arxiv_id":"2312.15098","n_code_links":0,"syntology":null},{"paper":"/paper/argue-with-me-tersely-towards-sentence-level","slug":"argue-with-me-tersely-towards-sentence-level","title":"Argue with Me Tersely: Towards Sentence-Level Counter-Argument Generation","date":"2023-12-21","arxiv_id":"2312.13608","n_code_links":1,"syntology":null},{"paper":"/paper/chatgpt-as-a-commenter-to-the-news-can-llms","slug":"chatgpt-as-a-commenter-to-the-news-can-llms","title":"ChatGPT as a commenter to the news: can LLMs generate human-like opinions?","date":"2023-12-21","arxiv_id":"2312.13961","n_code_links":1,"syntology":null},{"paper":"/paper/de-novo-drug-design-using-reinforcement-1","slug":"de-novo-drug-design-using-reinforcement-1","title":"De novo Drug Design using Reinforcement Learning with Multiple GPT Agents","date":"2023-12-21","arxiv_id":"2401.06155","n_code_links":2,"syntology":{"ran":3,"of":4,"n_ran_checked":3,"n_instrument":0,"unverified":1,"pointer_only":4,"phrase":"3 ran (of which 2 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["hxyfighter/molrl-mgpt"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":2,"n_ran_no_instrument_failure":3,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"how-to-prune-your-language-model-recovering","title":"How to Prune Your Language Model: Recovering Accuracy on the \"Sparsity May Cry'' Benchmark","date":"2023-12-21","arxiv_id":"2312.13547","n_code_links":0,"syntology":null},{"paper":null,"slug":"infovisdial-an-informative-visual-dialogue","title":"InfoVisDial: An Informative Visual Dialogue Dataset by Bridging Large Multimodal and Language Models","date":"2023-12-21","arxiv_id":"2312.13503","n_code_links":0,"syntology":null},{"paper":"/paper/provfl-client-driven-interpretability-of","slug":"provfl-client-driven-interpretability-of","title":"TraceFL: Interpretability-Driven Debugging in Federated Learning via Neuron Provenance","date":"2023-12-21","arxiv_id":"2312.13632","n_code_links":2,"syntology":null},{"paper":null,"slug":"team-irisapu-project-description-for-drc2023","title":"Team Irisapu Project Description for DRC2023","date":"2023-12-21","arxiv_id":"2312.13765","n_code_links":0,"syntology":null},{"paper":null,"slug":"typhoon-thai-large-language-models","title":"Typhoon: Thai Large Language Models","date":"2023-12-21","arxiv_id":"2312.13951","n_code_links":0,"syntology":null},{"paper":"/paper/agentcoder-multi-agent-based-code-generation","slug":"agentcoder-multi-agent-based-code-generation","title":"AgentCoder: Multi-Agent-based Code Generation with Iterative Testing and Optimisation","date":"2023-12-20","arxiv_id":"2312.13010","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 1 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["huangd1999/AgentCoder"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"benchmarking-and-analyzing-in-context","title":"Benchmarking and Analyzing In-context Learning, Fine-tuning and Supervised Learning for Biomedical Knowledge Curation: a focused study on chemical entities of biological interest","date":"2023-12-20","arxiv_id":"2312.12989","n_code_links":0,"syntology":null},{"paper":"/paper/domain-specific-code-language-models","slug":"domain-specific-code-language-models","title":"MonoCoder: Domain-Specific Code Language Model for HPC Codes and Tasks","date":"2023-12-20","arxiv_id":"2312.13322","n_code_links":3,"syntology":null},{"paper":null,"slug":"dynamic-fairness-aware-spectrum-auction-for","title":"Dynamic Fairness-Aware Spectrum Auction for Enhanced Licensed Shared Access in 6G Networks","date":"2023-12-20","arxiv_id":"2312.12867","n_code_links":0,"syntology":null},{"paper":"/paper/lookahead-an-inference-acceleration-framework","slug":"lookahead-an-inference-acceleration-framework","title":"Lookahead: An Inference Acceleration Framework for Large Language Model with Lossless Generation Accuracy","date":"2023-12-20","arxiv_id":"2312.12728","n_code_links":1,"syntology":{"ran":2,"of":3,"n_ran_checked":2,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["alipay/PainlessInferenceAcceleration"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"can-chatgpt-be-your-personal-medical","title":"Can ChatGPT be Your Personal Medical Assistant?","date":"2023-12-19","arxiv_id":"2312.12006","n_code_links":0,"syntology":null},{"paper":null,"slug":"can-transformers-learn-sequential-function","title":"Can Transformers Learn Sequential Function Classes In Context?","date":"2023-12-19","arxiv_id":"2312.12655","n_code_links":0,"syntology":null},{"paper":"/paper/detecting-technical-debt-using-natural","slug":"detecting-technical-debt-using-natural","title":"Self-Admitted Technical Debt Detection Approaches: A Decade Systematic Review","date":"2023-12-19","arxiv_id":"2312.15020","n_code_links":1,"syntology":null},{"paper":null,"slug":"efficient-title-reranker-for-fast-and","title":"Efficient Title Reranker for Fast and Improved Knowledge-Intense NLP","date":"2023-12-19","arxiv_id":"2312.12430","n_code_links":0,"syntology":null},{"paper":null,"slug":"large-language-models-in-medical-term","title":"Large Language Models in Medical Term Classification and Unexpected Misalignment Between Response and Reasoning","date":"2023-12-19","arxiv_id":"2312.14184","n_code_links":0,"syntology":null},{"paper":"/paper/an-in-depth-look-at-gemini-s-language","slug":"an-in-depth-look-at-gemini-s-language","title":"An In-depth Look at Gemini's Language Abilities","date":"2023-12-18","arxiv_id":"2312.11444","n_code_links":1,"syntology":{"ran":11,"of":12,"n_ran_checked":11,"n_instrument":0,"unverified":1,"pointer_only":12,"phrase":"11 ran (of which 0 constructed an object rather than computing a result; 11 with no instrument failure: 0 honoured, 0 violated, 11 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["neulab/gemini-benchmark"],"state":"official (archive's flag): 11 ran","n_ran":11,"n_constructed":0,"n_ran_no_instrument_failure":11,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"contextual-reinforcement-learning-for","title":"Contextual Reinforcement Learning for Offshore Wind Farm Bidding","date":"2023-12-18","arxiv_id":"2312.10884","n_code_links":0,"syntology":null},{"paper":null,"slug":"generative-linguistic-representation-for","title":"Generative linguistic representation for spoken language identification","date":"2023-12-18","arxiv_id":"2312.10964","n_code_links":0,"syntology":null},{"paper":"/paper/nomiracl-knowing-when-you-don-t-know-for","slug":"nomiracl-knowing-when-you-don-t-know-for","title":"\"Knowing When You Don't Know\": A Multilingual Relevance Assessment Dataset for Robust Retrieval-Augmented Generation","date":"2023-12-18","arxiv_id":"2312.11361","n_code_links":1,"syntology":{"ran":5,"of":6,"n_ran_checked":5,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["project-miracl/nomiracl"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/retrieval-augmented-generation-for-large","slug":"retrieval-augmented-generation-for-large","title":"Retrieval-Augmented Generation for Large Language Models: A Survey","date":"2023-12-18","arxiv_id":"2312.10997","n_code_links":4,"syntology":null},{"paper":null,"slug":"artificial-intelligence-optical-hardware","title":"Artificial intelligence optical hardware empowers high-resolution hyperspectral video understanding at 1.2 Tb/s","date":"2023-12-17","arxiv_id":"2312.10639","n_code_links":0,"syntology":null},{"paper":"/paper/bengali-intent-classification-with-generative","slug":"bengali-intent-classification-with-generative","title":"Bengali Intent Classification with Generative Adversarial BERT","date":"2023-12-17","arxiv_id":"2312.10679","n_code_links":1,"syntology":null},{"paper":null,"slug":"can-persistent-homology-whiten-transformer","title":"Can persistent homology whiten Transformer-based black-box models? A case study on BERT compression","date":"2023-12-17","arxiv_id":"2312.10702","n_code_links":0,"syntology":null},{"paper":"/paper/decoding-concerns-multi-label-classification","slug":"decoding-concerns-multi-label-classification","title":"Decoding Concerns: Multi-label Classification of Vaccine Sentiments in Social Media","date":"2023-12-17","arxiv_id":"2312.10626","n_code_links":1,"syntology":null},{"paper":null,"slug":"evaluating-ai-vocational-skills-through","title":"Evaluating AI Vocational Skills Through Professional Testing","date":"2023-12-17","arxiv_id":"2312.10603","n_code_links":0,"syntology":null},{"paper":"/paper/hyperpie-hyperparameter-information","slug":"hyperpie-hyperparameter-information","title":"HyperPIE: Hyperparameter Information Extraction from Scientific Publications","date":"2023-12-17","arxiv_id":"2312.10638","n_code_links":1,"syntology":null},{"paper":null,"slug":"investigating-salient-representations-and","title":"Investigating salient representations and label Variance in Dimensional Speech Emotion Analysis","date":"2023-12-17","arxiv_id":"2312.16180","n_code_links":0,"syntology":null},{"paper":null,"slug":"mixed-distillation-helps-smaller-language","title":"Mixed Distillation Helps Smaller Language Model Better Reasoning","date":"2023-12-17","arxiv_id":"2312.10730","n_code_links":0,"syntology":null},{"paper":"/paper/multi-label-classification-of-covid-tweets","slug":"multi-label-classification-of-covid-tweets","title":"Multi-Label Classification of COVID-Tweets Using Large Language Models","date":"2023-12-17","arxiv_id":"2312.10748","n_code_links":1,"syntology":null},{"paper":null,"slug":"t2m-hifigpt-generating-high-quality-human","title":"T2M-HiFiGPT: Generating High Quality Human Motion from Textual Descriptions with Residual Discrete Representations","date":"2023-12-17","arxiv_id":"2312.10628","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-comparative-analysis-of-large-language","title":"A Comparative Analysis of Large Language Models for Code Documentation Generation","date":"2023-12-16","arxiv_id":"2312.10349","n_code_links":0,"syntology":null},{"paper":null,"slug":"cross-linguistic-offensive-language-detection","title":"Cross-Linguistic Offensive Language Detection: BERT-Based Analysis of Bengali, Assamese, & Bodo Conversational Hateful Content from Social Media","date":"2023-12-16","arxiv_id":"2312.10528","n_code_links":0,"syntology":null},{"paper":"/paper/investigating-shallow-and-deep-learning","slug":"investigating-shallow-and-deep-learning","title":"Investigating Shallow and Deep Learning Techniques for Emotion Classification in Short Persian Texts","date":"2023-12-16","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/spt-fine-tuning-transformer-based-language","slug":"spt-fine-tuning-transformer-based-language","title":"SPT: Fine-Tuning Transformer-based Language Models Efficiently with Sparsification","date":"2023-12-16","arxiv_id":"2312.10365","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-novel-dataset-for-financial-education-text","title":"A Novel Dataset for Financial Education Text Simplification in Spanish","date":"2023-12-15","arxiv_id":"2312.09897","n_code_links":0,"syntology":null},{"paper":null,"slug":"algorithms-for-automatic-intents-extraction","title":"Algorithms for automatic intents extraction and utterances classification for goal-oriented dialogue systems","date":"2023-12-15","arxiv_id":"2312.09658","n_code_links":0,"syntology":null},{"paper":null,"slug":"distilling-large-language-models-for-matching","title":"Distilling Large Language Models for Matching Patients to Clinical Trials","date":"2023-12-15","arxiv_id":"2312.09958","n_code_links":0,"syntology":null},{"paper":"/paper/exploring-automatic-text-simplification-of","slug":"exploring-automatic-text-simplification-of","title":"Exploring Automatic Text Simplification of German Narrative Documents","date":"2023-12-15","arxiv_id":"2312.09907","n_code_links":1,"syntology":null},{"paper":"/paper/exploring-multi-level-threats-in-telegram","slug":"exploring-multi-level-threats-in-telegram","title":"Exploring Multi-Level Threats in Telegram Data with AI-Human Annotation: A Preliminary Study","date":"2023-12-15","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"no-skim-towards-efficiency-robustness","title":"No-Skim: Towards Efficiency Robustness Evaluation on Skimming-based Language Models","date":"2023-12-15","arxiv_id":"2312.09494","n_code_links":0,"syntology":null},{"paper":null,"slug":"red-ai-inconsistent-responses-from-gpt3-5","title":"Red AI? Inconsistent Responses from GPT3.5 Models on Political Issues in the US and China","date":"2023-12-15","arxiv_id":"2312.09917","n_code_links":0,"syntology":null},{"paper":"/paper/boosting-llm-reasoning-push-the-limits-of-few","slug":"boosting-llm-reasoning-push-the-limits-of-few","title":"Fewer is More: Boosting LLM Reasoning with Reinforced Context Pruning","date":"2023-12-14","arxiv_id":"2312.08901","n_code_links":0,"syntology":null},{"paper":"/paper/dissecting-vocabulary-biases-datasets-through","slug":"dissecting-vocabulary-biases-datasets-through","title":"Dissecting vocabulary biases datasets through statistical testing and automated data augmentation for artifact mitigation in Natural Language Inference","date":"2023-12-14","arxiv_id":"2312.08747","n_code_links":1,"syntology":null},{"paper":null,"slug":"entity-augmented-code-generation","title":"Dynamic Retrieval-Augmented Generation","date":"2023-12-14","arxiv_id":"2312.08976","n_code_links":0,"syntology":null},{"paper":null,"slug":"inter-layer-scheduling-space-exploration-for","title":"Inter-Layer Scheduling Space Exploration for Multi-model Inference on Heterogeneous Chiplets","date":"2023-12-14","arxiv_id":"2312.09401","n_code_links":0,"syntology":null},{"paper":null,"slug":"motion-flow-matching-for-human-motion","title":"Motion Flow Matching for Human Motion Synthesis and Editing","date":"2023-12-14","arxiv_id":"2312.08895","n_code_links":0,"syntology":null},{"paper":null,"slug":"self-evaluation-improves-selective-generation","title":"Self-Evaluation Improves Selective Generation in Large Language Models","date":"2023-12-14","arxiv_id":"2312.09300","n_code_links":0,"syntology":null},{"paper":null,"slug":"successor-heads-recurring-interpretable","title":"Successor Heads: Recurring, Interpretable Attention Heads In The Wild","date":"2023-12-14","arxiv_id":"2312.09230","n_code_links":0,"syntology":null},{"paper":"/paper/tinygsm-achieving-80-on-gsm8k-with-small","slug":"tinygsm-achieving-80-on-gsm8k-with-small","title":"TinyGSM: achieving >80% on GSM8k with small language models","date":"2023-12-14","arxiv_id":"2312.09241","n_code_links":0,"syntology":null},{"paper":null,"slug":"weak-to-strong-generalization-eliciting","title":"Weak-to-Strong Generalization: Eliciting Strong Capabilities With Weak Supervision","date":"2023-12-14","arxiv_id":"2312.09390","n_code_links":0,"syntology":null}],"record_sha256":"ac558f1d823d92e1fd0948ef2aa7541b19e299d95ad9d677ba9deadc1a6a8a18","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}