{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/attention-dropout/papers/38","list_of":"/method/attention-dropout","method":"Attention Dropout","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":38,"pages_in_order":109,"rows_per_page":100,"rows":[3701,3800],"of":10892,"counts":{"archive_papers_tagged":10892,"with_a_code_link":4634,"where_syntology_ran_a_sample":1270,"not_listed_spam_title":0,"listed":10892,"listed_where_code_ran":1270,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":1043,"every_run_a_failure_of_syntologys_instrument":227,"listed_with_a_run_with_no_instrument_failure":1043,"listed_every_run_a_failure_of_syntologys_instrument":227,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/attention-dropout","prev":"/method/attention-dropout/papers/37","next":"/method/attention-dropout/papers/39","papers":[{"paper":null,"slug":"what-will-my-model-forget-forecasting","title":"What Will My Model Forget? Forecasting Forgotten Examples in Language Model Refinement","date":"2024-02-02","arxiv_id":"2402.01865","n_code_links":0,"syntology":null},{"paper":null,"slug":"generation-distillation-and-evaluation-of","title":"Generation, Distillation and Evaluation of Motivational Interviewing-Style Reflections with a Foundational Language Model","date":"2024-02-01","arxiv_id":"2402.01051","n_code_links":0,"syntology":null},{"paper":null,"slug":"hiqa-a-hierarchical-contextual-augmentation","title":"HiQA: A Hierarchical Contextual Augmentation RAG for Multi-Documents QA","date":"2024-02-01","arxiv_id":"2402.01767","n_code_links":0,"syntology":null},{"paper":"/paper/improving-semantic-control-in-discrete-latent","slug":"improving-semantic-control-in-discrete-latent","title":"Improving Semantic Control in Discrete Latent Spaces with Transformer Quantized Variational Autoencoders","date":"2024-02-01","arxiv_id":"2402.00723","n_code_links":1,"syntology":null},{"paper":null,"slug":"learning-planning-based-reasoning-by","title":"Learning Planning-based Reasoning by Trajectories Collection and Process Reward Synthesizing","date":"2024-02-01","arxiv_id":"2402.00658","n_code_links":0,"syntology":null},{"paper":"/paper/reagent-towards-a-model-agnostic-feature","slug":"reagent-towards-a-model-agnostic-feature","title":"ReAGent: A Model-agnostic Feature Attribution Method for Generative Language Models","date":"2024-02-01","arxiv_id":"2402.00794","n_code_links":1,"syntology":null},{"paper":"/paper/recurrent-transformers-with-dynamic-halt","slug":"recurrent-transformers-with-dynamic-halt","title":"Investigating Recurrent Transformers with Dynamic Halt","date":"2024-02-01","arxiv_id":"2402.00976","n_code_links":1,"syntology":null},{"paper":null,"slug":"self-supervised-contrastive-pre-training-for-1","title":"Self-Supervised Contrastive Pre-Training for Multivariate Point Processes","date":"2024-02-01","arxiv_id":"2402.00987","n_code_links":0,"syntology":null},{"paper":"/paper/sparql-generation-with-entity-pre-trained-gpt","slug":"sparql-generation-with-entity-pre-trained-gpt","title":"SPARQL Generation with Entity Pre-trained GPT for KG Question Answering","date":"2024-02-01","arxiv_id":"2402.00969","n_code_links":1,"syntology":null},{"paper":null,"slug":"tiny-titans-can-smaller-large-language-models","title":"Tiny Titans: Can Smaller Large Language Models Punch Above Their Weight in the Real World for Meeting Summarization?","date":"2024-02-01","arxiv_id":"2402.00841","n_code_links":0,"syntology":null},{"paper":null,"slug":"towards-scalable-robotic-intervention-of","title":"Human-mediated Large Language Models for Robotic Intervention in Children with Autism Spectrum Disorders","date":"2024-02-01","arxiv_id":"2402.00260","n_code_links":0,"syntology":null},{"paper":"/paper/consmax-hardware-friendly-alternative-softmax","slug":"consmax-hardware-friendly-alternative-softmax","title":"ConSmax: Hardware-Friendly Alternative Softmax with Learnable Parameters","date":"2024-01-31","arxiv_id":"2402.10930","n_code_links":1,"syntology":null},{"paper":null,"slug":"global-liar-factuality-of-llms-over-time-and","title":"Global-Liar: Factuality of LLMs over Time and Geographic Regions","date":"2024-01-31","arxiv_id":"2401.17839","n_code_links":0,"syntology":null},{"paper":null,"slug":"making-a-long-story-short-in-conversation","title":"Making a Long Story Short in Conversation Modeling","date":"2024-01-31","arxiv_id":"2402.00143","n_code_links":0,"syntology":null},{"paper":null,"slug":"mitigating-the-problem-of-strong-priors-in","title":"Mitigating the Influence of Distractor Tasks in LMs with Prior-Aware Decoding","date":"2024-01-31","arxiv_id":"2401.17692","n_code_links":0,"syntology":null},{"paper":null,"slug":"paramanu-a-family-of-novel-efficient-indic","title":"Paramanu: A Family of Novel Efficient Generative Foundation Language Models for Indian Languages","date":"2024-01-31","arxiv_id":"2401.18034","n_code_links":0,"syntology":null},{"paper":null,"slug":"rag-fusion-a-new-take-on-retrieval-augmented","title":"RAG-Fusion: a New Take on Retrieval-Augmented Generation","date":"2024-01-31","arxiv_id":"2402.03367","n_code_links":0,"syntology":null},{"paper":null,"slug":"real-sparks-of-artificial-intelligence-and","title":"Real Sparks of Artificial Intelligence and the Importance of Inner Interpretability","date":"2024-01-31","arxiv_id":"2402.00901","n_code_links":0,"syntology":null},{"paper":null,"slug":"uncertainty-aware-explainable-recommendation","title":"Uncertainty-Aware Explainable Recommendation with Large Language Models","date":"2024-01-31","arxiv_id":"2402.03366","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-preliminary-study-on-using-large-language","title":"A Preliminary Study on Using Large Language Models in Software Pentesting","date":"2024-01-30","arxiv_id":"2401.17459","n_code_links":0,"syntology":null},{"paper":null,"slug":"arabic-tweet-act-a-weighted-ensemble-pre","title":"Arabic Tweet Act: A Weighted Ensemble Pre-Trained Transformer Model for Classifying Arabic Speech Acts on Twitter","date":"2024-01-30","arxiv_id":"2401.17373","n_code_links":0,"syntology":null},{"paper":"/paper/breaking-free-transformer-models-task","slug":"breaking-free-transformer-models-task","title":"Breaking Free Transformer Models: Task-specific Context Attribution Promises Improved Generalizability Without Fine-tuning Pre-trained LLMs","date":"2024-01-30","arxiv_id":"2401.16638","n_code_links":1,"syntology":null},{"paper":"/paper/crud-rag-a-comprehensive-chinese-benchmark","slug":"crud-rag-a-comprehensive-chinese-benchmark","title":"CRUD-RAG: A Comprehensive Chinese Benchmark for Retrieval-Augmented Generation of Large Language Models","date":"2024-01-30","arxiv_id":"2401.17043","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":3,"n_instrument":0,"unverified":0,"pointer_only":3,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["iaar-shanghai/crud_rag"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/detecting-mental-disorder-on-social-media-a","slug":"detecting-mental-disorder-on-social-media-a","title":"Detecting mental disorder on social media: a ChatGPT-augmented explainable approach","date":"2024-01-30","arxiv_id":"2401.17477","n_code_links":1,"syntology":null},{"paper":null,"slug":"detecting-racist-text-in-bengali-an-ensemble","title":"Detecting Racist Text in Bengali: An Ensemble Deep Learning Framework","date":"2024-01-30","arxiv_id":"2401.16748","n_code_links":0,"syntology":null},{"paper":null,"slug":"fine-tuning-transformer-based-encoder-for","title":"Fine-tuning Transformer-based Encoder for Turkish Language Understanding Tasks","date":"2024-01-30","arxiv_id":"2401.17396","n_code_links":0,"syntology":null},{"paper":null,"slug":"large-multi-modal-models-lmms-as-universal","title":"Large Multi-Modal Models (LMMs) as Universal Foundation Models for AI-Native Wireless Systems","date":"2024-01-30","arxiv_id":"2402.01748","n_code_links":0,"syntology":null},{"paper":"/paper/llamp-large-language-model-made-powerful-for","slug":"llamp-large-language-model-made-powerful-for","title":"LLaMP: Large Language Model Made Powerful for High-fidelity Materials Knowledge Retrieval and Distillation","date":"2024-01-30","arxiv_id":"2401.17244","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["chiang-yuan/llamp"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/mt-eval-a-multi-turn-capabilities-evaluation","slug":"mt-eval-a-multi-turn-capabilities-evaluation","title":"MT-Eval: A Multi-Turn Capabilities Evaluation Benchmark for Large Language Models","date":"2024-01-30","arxiv_id":"2401.16745","n_code_links":1,"syntology":{"ran":9,"of":11,"n_ran_checked":9,"n_instrument":0,"unverified":2,"pointer_only":2,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 9 with no instrument failure: 0 honoured, 0 violated, 9 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["kwanwaichung/mt-eval"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":9,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"single-word-change-is-all-you-need-designing","title":"Single Word Change is All You Need: Designing Attacks and Defenses for Text Classifiers","date":"2024-01-30","arxiv_id":"2401.17196","n_code_links":0,"syntology":null},{"paper":null,"slug":"towards-generating-informative-textual","title":"Towards Generating Informative Textual Description for Neurons in Language Models","date":"2024-01-30","arxiv_id":"2401.16731","n_code_links":0,"syntology":null},{"paper":null,"slug":"credit-risk-meets-large-language-models","title":"Credit Risk Meets Large Language Models: Building a Risk Indicator from Loan Descriptions in P2P Lending","date":"2024-01-29","arxiv_id":"2401.16458","n_code_links":0,"syntology":null},{"paper":null,"slug":"development-and-testing-of-a-novel-large","title":"Development and Testing of a Novel Large Language Model-Based Clinical Decision Support Systems for Medication Safety in 12 Clinical Specialties","date":"2024-01-29","arxiv_id":"2402.01741","n_code_links":0,"syntology":null},{"paper":null,"slug":"development-and-testing-of-retrieval","title":"Development and Testing of Retrieval Augmented Generation in Large Language Models -- A Case Study Report","date":"2024-01-29","arxiv_id":"2402.01733","n_code_links":0,"syntology":null},{"paper":null,"slug":"diverse-but-divisive-llms-can-exaggerate","title":"Diverse, but Divisive: LLMs Can Exaggerate Gender Differences in Opinion Related to Harms of Misinformation","date":"2024-01-29","arxiv_id":"2401.16558","n_code_links":0,"syntology":null},{"paper":null,"slug":"drbert-unveiling-the-potential-of-masked","title":"BPDec: Unveiling the Potential of Masked Language Modeling Decoder in BERT pretraining","date":"2024-01-29","arxiv_id":"2401.15861","n_code_links":0,"syntology":null},{"paper":"/paper/e-eval-a-comprehensive-chinese-k-12-education","slug":"e-eval-a-comprehensive-chinese-k-12-education","title":"E-EVAL: A Comprehensive Chinese K-12 Education Evaluation Benchmark for Large Language Models","date":"2024-01-29","arxiv_id":"2401.15927","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["ai-edu-lab/e-eval"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"leveraging-professional-radiologists","title":"Leveraging Professional Radiologists' Expertise to Enhance LLMs' Evaluation for Radiology Reports","date":"2024-01-29","arxiv_id":"2401.16578","n_code_links":0,"syntology":null},{"paper":null,"slug":"llm4vuln-a-unified-evaluation-framework-for","title":"LLM4Vuln: A Unified Evaluation Framework for Decoupling and Enhancing LLMs' Vulnerability Reasoning","date":"2024-01-29","arxiv_id":"2401.16185","n_code_links":0,"syntology":null},{"paper":null,"slug":"multi-class-regret-detection-in-hindi","title":"Multi-class Regret Detection in Hindi Devanagari Script","date":"2024-01-29","arxiv_id":"2401.16561","n_code_links":0,"syntology":null},{"paper":"/paper/regal-refactoring-programs-to-discover","slug":"regal-refactoring-programs-to-discover","title":"ReGAL: Refactoring Programs to Discover Generalizable Abstractions","date":"2024-01-29","arxiv_id":"2401.16467","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":2,"n_instrument":0,"unverified":0,"pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["esteng/regal_program_learning"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"security-code-review-by-llms-a-deep-dive-into","title":"An Insight into Security Code Review with LLMs: Capabilities, Obstacles, and Influential Factors","date":"2024-01-29","arxiv_id":"2401.16310","n_code_links":0,"syntology":null},{"paper":"/paper/topro-token-level-prompt-decomposition-for","slug":"topro-token-level-prompt-decomposition-for","title":"ToPro: Token-Level Prompt Decomposition for Cross-Lingual Sequence Labeling Tasks","date":"2024-01-29","arxiv_id":"2401.16589","n_code_links":1,"syntology":null},{"paper":null,"slug":"trackgpt-a-generative-pre-trained-transformer","title":"TrackGPT -- A generative pre-trained transformer for cross-domain entity trajectory forecasting","date":"2024-01-29","arxiv_id":"2402.00066","n_code_links":0,"syntology":null},{"paper":"/paper/contrastive-learning-and-mixture-of-experts","slug":"contrastive-learning-and-mixture-of-experts","title":"Contrastive Learning and Mixture of Experts Enables Precise Vector Embeddings","date":"2024-01-28","arxiv_id":"2401.15713","n_code_links":1,"syntology":null},{"paper":null,"slug":"unmasked-quantifying-gender-biases-in-masked","title":"UnMASKed: Quantifying Gender Biases in Masked Language Models through Linguistically Informed Job Market Prompts","date":"2024-01-28","arxiv_id":"2401.15798","n_code_links":0,"syntology":null},{"paper":"/paper/convosense-overcoming-monotonous-commonsense","slug":"convosense-overcoming-monotonous-commonsense","title":"ConvoSense: Overcoming Monotonous Commonsense Inferences for Conversational AI","date":"2024-01-27","arxiv_id":"2401.15471","n_code_links":1,"syntology":null},{"paper":null,"slug":"enhancing-large-language-model-performance-to","title":"Enhancing Large Language Model Performance To Answer Questions and Extract Information More Accurately","date":"2024-01-27","arxiv_id":"2402.01722","n_code_links":0,"syntology":null},{"paper":null,"slug":"equipping-language-models-with-tool-use","title":"Equipping Language Models with Tool Use Capability for Tabular Data Analysis in Finance","date":"2024-01-27","arxiv_id":"2401.15328","n_code_links":0,"syntology":null},{"paper":null,"slug":"fortifying-ethical-boundaries-in-ai-advanced","title":"Fortifying Ethical Boundaries in AI: Advanced Strategies for Enhancing Security in Large Language Models","date":"2024-01-27","arxiv_id":"2402.01725","n_code_links":0,"syntology":null},{"paper":"/paper/multihop-rag-benchmarking-retrieval-augmented","slug":"multihop-rag-benchmarking-retrieval-augmented","title":"MultiHop-RAG: Benchmarking Retrieval-Augmented Generation for Multi-Hop Queries","date":"2024-01-27","arxiv_id":"2401.15391","n_code_links":2,"syntology":{"ran":4,"of":4,"n_ran_checked":1,"n_instrument":3,"unverified":0,"pointer_only":4,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 1 violated, 0 with no contract checked; 3 where Syntology's instrument failed) · 0 unverified","official":{"repos":["yixuantt/MultiHop-RAG"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/from-rag-to-qa-rag-integrating-generative-ai","slug":"from-rag-to-qa-rag-integrating-generative-ai","title":"From RAG to QA-RAG: Integrating Generative AI for Pharmaceutical Regulatory Compliance Process","date":"2024-01-26","arxiv_id":"2402.01717","n_code_links":1,"syntology":null},{"paper":null,"slug":"geodecoder-empowering-multimodal-map","title":"GeoDecoder: Empowering Multimodal Map Understanding","date":"2024-01-26","arxiv_id":"2401.15118","n_code_links":0,"syntology":null},{"paper":null,"slug":"mptq-vit-mixed-precisionpost","title":"MPTQ-ViT: Mixed-Precision Post-Training Quantization for Vision Transformer","date":"2024-01-26","arxiv_id":"2401.14895","n_code_links":0,"syntology":null},{"paper":null,"slug":"scalable-qualitative-coding-with-llms-chain","title":"Scalable Qualitative Coding with LLMs: Chain-of-Thought Reasoning Matches Human Performance in Some Hermeneutic Tasks","date":"2024-01-26","arxiv_id":"2401.15170","n_code_links":0,"syntology":null},{"paper":"/paper/the-power-of-noise-redefining-retrieval-for","slug":"the-power-of-noise-redefining-retrieval-for","title":"The Power of Noise: Redefining Retrieval for RAG Systems","date":"2024-01-26","arxiv_id":"2401.14887","n_code_links":3,"syntology":{"ran":5,"of":5,"n_ran_checked":5,"n_instrument":0,"unverified":0,"pointer_only":5,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["florin-git/The-Power-of-Noise"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"a-comparative-study-of-zero-shot-inference","title":"A comparative study of zero-shot inference with large language models and supervised modeling in breast cancer pathology classification","date":"2024-01-25","arxiv_id":"2401.13887","n_code_links":0,"syntology":null},{"paper":"/paper/chat-gpt-v-bert-dawn-of-justice-for-semantic","slug":"chat-gpt-v-bert-dawn-of-justice-for-semantic","title":"(Chat)GPT v BERT: Dawn of Justice for Semantic Change Detection","date":"2024-01-25","arxiv_id":"2401.14040","n_code_links":1,"syntology":null},{"paper":"/paper/deepseek-coder-when-the-large-language-model","slug":"deepseek-coder-when-the-large-language-model","title":"DeepSeek-Coder: When the Large Language Model Meets Programming -- The Rise of Code Intelligence","date":"2024-01-25","arxiv_id":"2401.14196","n_code_links":1,"syntology":{"ran":9,"of":10,"n_ran_checked":9,"n_instrument":0,"unverified":1,"pointer_only":1,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 9 with no instrument failure: 0 honoured, 0 violated, 9 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["deepseek-ai/DeepSeek-Coder"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":9,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"enhanced-labeling-technique-for-reddit-text","title":"Enhanced Labeling Technique for Reddit Text and Fine-Tuned Longformer Models for Classifying Depression Severity in English and Luganda","date":"2024-01-25","arxiv_id":"2401.14240","n_code_links":0,"syntology":null},{"paper":null,"slug":"evaluating-gpt-3-5-s-awareness-and","title":"Evaluating GPT-3.5's Awareness and Summarization Abilities for European Constitutional Texts with Shared Topics","date":"2024-01-25","arxiv_id":"2401.14524","n_code_links":0,"syntology":null},{"paper":null,"slug":"investigate-consolidate-exploit-a-general","title":"Investigate-Consolidate-Exploit: A General Strategy for Inter-Task Agent Self-Evolution","date":"2024-01-25","arxiv_id":"2401.13996","n_code_links":0,"syntology":null},{"paper":"/paper/longhealth-a-question-answering-benchmark","slug":"longhealth-a-question-answering-benchmark","title":"LongHealth: A Question Answering Benchmark with Long Clinical Documents","date":"2024-01-25","arxiv_id":"2401.14490","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":0,"n_instrument":3,"unverified":0,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 3 where Syntology's instrument failed) · 0 unverified","official":{"repos":["kbressem/longhealth"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"socially-aware-synthetic-data-generation-for","title":"Socially Aware Synthetic Data Generation for Suicidal Ideation Detection Using Large Language Models","date":"2024-01-25","arxiv_id":"2402.01712","n_code_links":0,"syntology":null},{"paper":"/paper/tricy-trigger-guided-data-to-text-generation-1","slug":"tricy-trigger-guided-data-to-text-generation-1","title":"TrICy: Trigger-guided Data-to-text Generation with Intent aware Attention-Copy","date":"2024-01-25","arxiv_id":"2402.01714","n_code_links":0,"syntology":null},{"paper":null,"slug":"unmasking-and-quantifying-racial-bias-of","title":"Unmasking and Quantifying Racial Bias of Large Language Models in Medical Report Generation","date":"2024-01-25","arxiv_id":"2401.13867","n_code_links":0,"syntology":null},{"paper":null,"slug":"zs4c-zero-shot-synthesis-of-compilable-code","title":"ZS4C: Zero-Shot Synthesis of Compilable Code for Incomplete Code Snippets using LLMs","date":"2024-01-25","arxiv_id":"2401.14279","n_code_links":0,"syntology":null},{"paper":"/paper/a-unified-approach-to-emotion-detection-and","slug":"a-unified-approach-to-emotion-detection-and","title":"A Unified Approach to Emotion Detection and Task-Oriented Dialogue Modeling","date":"2024-01-24","arxiv_id":"2401.13789","n_code_links":1,"syntology":null},{"paper":null,"slug":"automated-root-causing-of-cloud-incidents","title":"Automated Root Causing of Cloud Incidents using In-Context Learning with GPT-4","date":"2024-01-24","arxiv_id":"2401.13810","n_code_links":0,"syntology":null},{"paper":"/paper/can-gpt-3-5-generate-and-code-discharge","slug":"can-gpt-3-5-generate-and-code-discharge","title":"Can GPT-3.5 Generate and Code Discharge Summaries?","date":"2024-01-24","arxiv_id":"2401.13512","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":0,"n_instrument":3,"unverified":0,"pointer_only":3,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 3 where Syntology's instrument failed) · 0 unverified","official":{"repos":["edinburghclinicalnlp/chatgpt_icd_coding"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"discovering-mathematical-formulas-from-data","title":"Discovering Mathematical Formulas from Data via GPT-guided Monte Carlo Tree Search","date":"2024-01-24","arxiv_id":"2401.14424","n_code_links":0,"syntology":null},{"paper":null,"slug":"evaluation-of-general-large-language-models","title":"Evaluation of General Large Language Models in Contextually Assessing Semantic Concepts Extracted from Adult Critical Care Electronic Health Record Notes","date":"2024-01-24","arxiv_id":"2401.13588","n_code_links":0,"syntology":null},{"paper":null,"slug":"graph-guided-question-answer-generation-for","title":"Graph Guided Question Answer Generation for Procedural Question-Answering","date":"2024-01-24","arxiv_id":"2401.13594","n_code_links":0,"syntology":null},{"paper":"/paper/how-good-is-chatgpt-at-face-biometrics-a","slug":"how-good-is-chatgpt-at-face-biometrics-a","title":"How Good is ChatGPT at Face Biometrics? A First Look into Recognition, Soft Biometrics, and Explainability","date":"2024-01-24","arxiv_id":"2401.13641","n_code_links":1,"syntology":null},{"paper":null,"slug":"proactive-emotion-tracker-ai-driven","title":"Proactive Emotion Tracker: AI-Driven Continuous Mood and Emotion Monitoring","date":"2024-01-24","arxiv_id":"2401.13722","n_code_links":0,"syntology":null},{"paper":null,"slug":"scalable-link-prediction-on-large-scale","title":"LPNL: Scalable Link Prediction with Large Language Models","date":"2024-01-24","arxiv_id":"2401.13227","n_code_links":0,"syntology":null},{"paper":null,"slug":"segment-any-cell-a-sam-based-auto-prompting","title":"Segment Any Cell: A SAM-based Auto-prompting Fine-tuning Framework for Nuclei Segmentation","date":"2024-01-24","arxiv_id":"2401.13220","n_code_links":0,"syntology":null},{"paper":"/paper/contrastive-learning-in-distilled-models","slug":"contrastive-learning-in-distilled-models","title":"Contrastive Learning in Distilled Models","date":"2024-01-23","arxiv_id":"2401.12472","n_code_links":1,"syntology":null},{"paper":null,"slug":"detecting-and-recognizing-characters-in-greek","title":"Detecting and recognizing characters in Greek papyri with YOLOv8, DeiT and SimCLR","date":"2024-01-23","arxiv_id":"2401.12513","n_code_links":0,"syntology":null},{"paper":null,"slug":"fast-adversarial-training-against-textual","title":"Fast Adversarial Training against Textual Adversarial Attacks","date":"2024-01-23","arxiv_id":"2401.12461","n_code_links":0,"syntology":null},{"paper":null,"slug":"kam-cot-knowledge-augmented-multimodal-chain","title":"KAM-CoT: Knowledge Augmented Multimodal Chain-of-Thoughts Reasoning","date":"2024-01-23","arxiv_id":"2401.12863","n_code_links":0,"syntology":null},{"paper":null,"slug":"revolutionizing-retrieval-augmented","title":"Revolutionizing Retrieval-Augmented Generation with Enhanced PDF Structure Recognition","date":"2024-01-23","arxiv_id":"2401.12599","n_code_links":0,"syntology":null},{"paper":"/paper/trove-inducing-verifiable-and-efficient","slug":"trove-inducing-verifiable-and-efficient","title":"TroVE: Inducing Verifiable and Efficient Toolboxes for Solving Programmatic Tasks","date":"2024-01-23","arxiv_id":"2401.12869","n_code_links":1,"syntology":{"ran":16,"of":19,"n_ran_checked":16,"n_instrument":0,"unverified":3,"pointer_only":19,"phrase":"16 ran (of which 0 constructed an object rather than computing a result; 16 with no instrument failure: 0 honoured, 0 violated, 16 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":{"repos":["zorazrw/trove"],"state":"official (archive's flag): 16 ran","n_ran":16,"n_constructed":0,"n_ran_no_instrument_failure":16,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":"/paper/apt-adaptive-pruning-and-tuning-pretrained","slug":"apt-adaptive-pruning-and-tuning-pretrained","title":"APT: Adaptive Pruning and Tuning Pretrained Language Models for Efficient Training and Inference","date":"2024-01-22","arxiv_id":"2401.12200","n_code_links":1,"syntology":{"ran":6,"of":7,"n_ran_checked":5,"n_instrument":1,"unverified":1,"pointer_only":1,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":{"repos":["roim1998/apt"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/enhancing-in-context-learning-via-linear","slug":"enhancing-in-context-learning-via-linear","title":"Enhancing In-context Learning via Linear Probe Calibration","date":"2024-01-22","arxiv_id":"2401.12406","n_code_links":1,"syntology":{"ran":3,"of":6,"n_ran_checked":3,"n_instrument":0,"unverified":3,"pointer_only":6,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":{"repos":["mominabbass/linc"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"keep-decoding-parallel-with-effective","title":"Keep Decoding Parallel with Effective Knowledge Distillation from Language Models to End-to-end Speech Recognisers","date":"2024-01-22","arxiv_id":"2401.11700","n_code_links":0,"syntology":null},{"paper":null,"slug":"the-right-model-for-the-job-an-evaluation-of","title":"The Right Model for the Job: An Evaluation of Legal Multi-Label Classification Baselines","date":"2024-01-22","arxiv_id":"2401.11852","n_code_links":0,"syntology":null},{"paper":null,"slug":"zero-space-cost-fault-tolerance-for","title":"Zero-Space Cost Fault Tolerance for Transformer-based Language Models on ReRAM","date":"2024-01-22","arxiv_id":"2401.11664","n_code_links":0,"syntology":null},{"paper":"/paper/chex-gpt-harnessing-large-language-models-for","slug":"chex-gpt-harnessing-large-language-models-for","title":"CheX-GPT: Harnessing Large Language Models for Enhanced Chest X-ray Report Labeling","date":"2024-01-21","arxiv_id":"2401.11505","n_code_links":2,"syntology":null},{"paper":null,"slug":"confidence-preservation-property-in-knowledge","title":"Confidence Preservation Property in Knowledge Distillation Abstractions","date":"2024-01-21","arxiv_id":"2401.11365","n_code_links":0,"syntology":null},{"paper":null,"slug":"enhancing-recommendation-diversity-by-re","title":"Enhancing Recommendation Diversity by Re-ranking with Large Language Models","date":"2024-01-21","arxiv_id":"2401.11506","n_code_links":0,"syntology":null},{"paper":"/paper/finding-a-needle-in-the-adversarial-haystack","slug":"finding-a-needle-in-the-adversarial-haystack","title":"Finding a Needle in the Adversarial Haystack: A Targeted Paraphrasing Approach For Uncovering Edge Cases with Minimal Distribution Distortion","date":"2024-01-21","arxiv_id":"2401.11373","n_code_links":1,"syntology":null},{"paper":"/paper/sebertnets-sequence-enhanced-bert-networks","slug":"sebertnets-sequence-enhanced-bert-networks","title":"SEBERTNets: Sequence Enhanced BERT Networks for Event Entity Extraction Tasks Oriented to the Finance Field","date":"2024-01-21","arxiv_id":"2401.11408","n_code_links":1,"syntology":null},{"paper":"/paper/badchain-backdoor-chain-of-thought-prompting","slug":"badchain-backdoor-chain-of-thought-prompting","title":"BadChain: Backdoor Chain-of-Thought Prompting for Large Language Models","date":"2024-01-20","arxiv_id":"2401.12242","n_code_links":1,"syntology":{"ran":13,"of":15,"n_ran_checked":13,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"13 ran (of which 0 constructed an object rather than computing a result; 13 with no instrument failure: 0 honoured, 0 violated, 13 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["django-jiang/badchain"],"state":"official (archive's flag): 13 ran","n_ran":13,"n_constructed":0,"n_ran_no_instrument_failure":13,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/drop-your-decoder-pre-training-with-bag-of","slug":"drop-your-decoder-pre-training-with-bag-of","title":"Drop your Decoder: Pre-training with Bag-of-Word Prediction for Dense Passage Retrieval","date":"2024-01-20","arxiv_id":"2401.11248","n_code_links":3,"syntology":{"ran":2,"of":2,"n_ran_checked":1,"n_instrument":1,"unverified":0,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["ma787639046/bowdpr"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed","official"]}}},{"paper":null,"slug":"enhancing-large-language-models-for-clinical","title":"Enhancing Large Language Models for Clinical Decision Support by Incorporating Clinical Practice Guidelines","date":"2024-01-20","arxiv_id":"2401.11120","n_code_links":0,"syntology":null},{"paper":null,"slug":"evaluating-and-enhancing-large-language","title":"Evaluating and Enhancing Large Language Models Performance in Domain-specific Medicine: Osteoarthritis Management with DocOA","date":"2024-01-20","arxiv_id":"2401.12998","n_code_links":0,"syntology":null},{"paper":null,"slug":"lrp-qvit-mixed-precision-vision-transformer","title":"LRP-QViT: Mixed-Precision Vision Transformer Quantization via Layer-wise Relevance Propagation","date":"2024-01-20","arxiv_id":"2401.11243","n_code_links":0,"syntology":null},{"paper":null,"slug":"prompt-rag-pioneering-vector-embedding-free","title":"Prompt-RAG: Pioneering Vector Embedding-Free Retrieval-Augmented Generation in Niche Domains, Exemplified by Korean Medicine","date":"2024-01-20","arxiv_id":"2401.11246","n_code_links":0,"syntology":null},{"paper":null,"slug":"unfair-tos-an-automated-approach-using","title":"Unfair TOS: An Automated Approach using Customized BERT","date":"2024-01-20","arxiv_id":"2401.11207","n_code_links":0,"syntology":null}],"record_sha256":"efdac002d7f4abfae7e0962c5bcbf759f56072997e55e2b61e618402d3dc643c","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}