{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/attention-dropout/papers/42","list_of":"/method/attention-dropout","method":"Attention Dropout","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":42,"pages_in_order":109,"rows_per_page":100,"rows":[4101,4200],"of":10892,"counts":{"archive_papers_tagged":10892,"with_a_code_link":4634,"where_syntology_ran_a_sample":1270,"not_listed_spam_title":0,"listed":10892,"listed_where_code_ran":1270,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":1043,"every_run_a_failure_of_syntologys_instrument":227,"listed_with_a_run_with_no_instrument_failure":1043,"listed_every_run_a_failure_of_syntologys_instrument":227,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/attention-dropout","prev":"/method/attention-dropout/papers/41","next":"/method/attention-dropout/papers/43","papers":[{"paper":null,"slug":"real-customization-or-just-marketing-are","title":"Real Customization or Just Marketing: Are Customized Versions of Chat GPT Useful?","date":"2023-11-27","arxiv_id":"2312.03728","n_code_links":0,"syntology":null},{"paper":"/paper/ssin-self-supervised-learning-for-rainfall","slug":"ssin-self-supervised-learning-for-rainfall","title":"SSIN: Self-Supervised Learning for Rainfall Spatial Interpolation","date":"2023-11-27","arxiv_id":"2311.15530","n_code_links":1,"syntology":null},{"paper":null,"slug":"leveraging-ai-derived-data-for-carbon","title":"Leveraging AI-derived Data for Carbon Accounting: Information Extraction from Alternative Sources","date":"2023-11-26","arxiv_id":"2312.03722","n_code_links":0,"syntology":null},{"paper":"/paper/machine-generated-text-detection-using-deep","slug":"machine-generated-text-detection-using-deep","title":"Machine-Generated Text Detection using Deep Learning","date":"2023-11-26","arxiv_id":"2311.15425","n_code_links":1,"syntology":null},{"paper":"/paper/uhgeval-benchmarking-the-hallucination-of","slug":"uhgeval-benchmarking-the-hallucination-of","title":"UHGEval: Benchmarking the Hallucination of Chinese Large Language Models via Unconstrained Generation","date":"2023-11-26","arxiv_id":"2311.15296","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":3,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["IAAR-Shanghai/UHGEval"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"uncertainty-aware-language-modeling-for","title":"Uncertainty-aware Language Modeling for Selective Question Answering","date":"2023-11-26","arxiv_id":"2311.15451","n_code_links":0,"syntology":null},{"paper":null,"slug":"swiftlearn-a-data-efficient-training-method","title":"SwiftLearn: A Data-Efficient Training Method of Deep Learning Models using Importance Sampling","date":"2023-11-25","arxiv_id":"2311.15134","n_code_links":0,"syntology":null},{"paper":null,"slug":"cmed-gpt-prompt-tuning-for-entity-aware","title":"CMed-GPT: Prompt Tuning for Entity-Aware Chinese Medical Dialogue Generation","date":"2023-11-24","arxiv_id":"2311.14539","n_code_links":0,"syntology":null},{"paper":"/paper/data-to-text-bilingual-generation","slug":"data-to-text-bilingual-generation","title":"Data-to-Text Bilingual Generation","date":"2023-11-24","arxiv_id":"2311.14808","n_code_links":2,"syntology":null},{"paper":"/paper/gpt-struct-me-probing-gpt-models-on-narrative","slug":"gpt-struct-me-probing-gpt-models-on-narrative","title":"GPT Struct Me: Probing GPT Models on Narrative Entity Extraction","date":"2023-11-24","arxiv_id":"2311.14583","n_code_links":1,"syntology":null},{"paper":"/paper/image-super-resolution-with-text-prompt","slug":"image-super-resolution-with-text-prompt","title":"Image Super-Resolution with Text Prompt Diffusion","date":"2023-11-24","arxiv_id":"2311.14282","n_code_links":1,"syntology":null},{"paper":null,"slug":"large-language-models-as-automated-aligners","title":"Large Language Models as Automated Aligners for benchmarking Vision-Language Models","date":"2023-11-24","arxiv_id":"2311.14580","n_code_links":0,"syntology":null},{"paper":null,"slug":"machine-translation-for-ge-ez-language","title":"Machine Translation for Ge'ez Language","date":"2023-11-24","arxiv_id":"2311.14530","n_code_links":0,"syntology":null},{"paper":"/paper/a-cross-attention-approach-to-diagnostic","slug":"a-cross-attention-approach-to-diagnostic","title":"A Cross Attention Approach to Diagnostic Explainability using Clinical Practice Guidelines for Depression","date":"2023-11-23","arxiv_id":"2311.13852","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-multi-solution-study-on-gdpr-ai-enabled","title":"A Multi-solution Study on GDPR AI-enabled Completeness Checking of DPAs","date":"2023-11-23","arxiv_id":"2311.13881","n_code_links":0,"syntology":null},{"paper":"/paper/annotation-sensitivity-training-data","slug":"annotation-sensitivity-training-data","title":"Annotation Sensitivity: Training Data Collection Methods Affect Model Performance","date":"2023-11-23","arxiv_id":"2311.14212","n_code_links":1,"syntology":null},{"paper":"/paper/hardware-resilience-properties-of-text-guided-1","slug":"hardware-resilience-properties-of-text-guided-1","title":"Hardware Resilience Properties of Text-Guided Image Classifiers","date":"2023-11-23","arxiv_id":"2311.14062","n_code_links":1,"syntology":{"ran":3,"of":4,"n_ran_checked":1,"n_instrument":2,"unverified":1,"pointer_only":4,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 2 where Syntology's instrument failed) · 1 unverified","official":{"repos":["talalwasim/textguidedresilience"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"minimizing-factual-inconsistency-and","title":"Minimizing Factual Inconsistency and Hallucination in Large Language Models","date":"2023-11-23","arxiv_id":"2311.13878","n_code_links":0,"syntology":null},{"paper":null,"slug":"towards-auditing-large-language-models","title":"Towards Auditing Large Language Models: Improving Text-based Stereotype Detection","date":"2023-11-23","arxiv_id":"2311.14126","n_code_links":0,"syntology":null},{"paper":"/paper/comparison-of-pipeline-sequence-to-sequence","slug":"comparison-of-pipeline-sequence-to-sequence","title":"Comparison of pipeline, sequence-to-sequence, and GPT models for end-to-end relation extraction: experiments with the rare disease use-case","date":"2023-11-22","arxiv_id":"2311.13729","n_code_links":1,"syntology":null},{"paper":"/paper/compeft-compression-for-communicating","slug":"compeft-compression-for-communicating","title":"ComPEFT: Compression for Communicating Parameter Efficient Updates via Sparsification and Quantization","date":"2023-11-22","arxiv_id":"2311.13171","n_code_links":1,"syntology":{"ran":8,"of":10,"n_ran_checked":7,"n_instrument":1,"unverified":2,"pointer_only":1,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","official":{"repos":["prateeky2806/compeft"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":7,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"current-topological-and-machine-learning","title":"Current Topological and Machine Learning Applications for Bias Detection in Text","date":"2023-11-22","arxiv_id":"2311.13495","n_code_links":0,"syntology":null},{"paper":"/paper/detecting-out-of-distribution-text-using","slug":"detecting-out-of-distribution-text-using","title":"Detecting out-of-distribution text using topological features of transformer-based language models","date":"2023-11-22","arxiv_id":"2311.13102","n_code_links":1,"syntology":null},{"paper":null,"slug":"drilling-down-into-the-discourse-structure","title":"Drilling Down into the Discourse Structure with LLMs for Long Document Question Answering","date":"2023-11-22","arxiv_id":"2311.13565","n_code_links":0,"syntology":null},{"paper":null,"slug":"generation-of-explanations-for-logic","title":"Generation of Explanations for Logic Reasoning","date":"2023-11-22","arxiv_id":"2311.13455","n_code_links":0,"syntology":null},{"paper":null,"slug":"nova-generative-language-models-for-binaries","title":"Nova: Generative Language Models for Assembly Code with Hierarchical Attention and Contrastive Learning","date":"2023-11-22","arxiv_id":"2311.13721","n_code_links":0,"syntology":null},{"paper":"/paper/pg-video-llava-pixel-grounding-large-video","slug":"pg-video-llava-pixel-grounding-large-video","title":"PG-Video-LLaVA: Pixel Grounding Large Video-Language Models","date":"2023-11-22","arxiv_id":"2311.13435","n_code_links":1,"syntology":{"ran":3,"of":5,"n_ran_checked":2,"n_instrument":1,"unverified":2,"pointer_only":5,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","official":{"repos":["mbzuai-oryx/video-llava"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/speak-like-a-native-prompting-large-language","slug":"speak-like-a-native-prompting-large-language","title":"AlignedCoT: Prompting Large Language Models via Native-Speaking Demonstrations","date":"2023-11-22","arxiv_id":"2311.13538","n_code_links":1,"syntology":null},{"paper":null,"slug":"ve-a-chatbot-for-latin","title":"@ve: A Chatbot for Latin","date":"2023-11-22","arxiv_id":"2311.14741","n_code_links":0,"syntology":null},{"paper":"/paper/white-box-transformers-via-sparse-rate-1","slug":"white-box-transformers-via-sparse-rate-1","title":"White-Box Transformers via Sparse Rate Reduction: Compression Is All There Is?","date":"2023-11-22","arxiv_id":"2311.13110","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":1,"n_instrument":2,"unverified":0,"pointer_only":3,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 1 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":null}},{"paper":null,"slug":"a-survey-on-large-language-models-for-1","title":"A Survey on Large Language Models for Personalized and Explainable Recommendations","date":"2023-11-21","arxiv_id":"2311.12338","n_code_links":0,"syntology":null},{"paper":null,"slug":"academicgpt-empowering-academic-research","title":"AcademicGPT: Empowering Academic Research","date":"2023-11-21","arxiv_id":"2311.12315","n_code_links":0,"syntology":null},{"paper":"/paper/alpha-anomalous-physiological-health","slug":"alpha-anomalous-physiological-health","title":"ALPHA: AnomaLous Physiological Health Assessment Using Large Language Models","date":"2023-11-21","arxiv_id":"2311.12524","n_code_links":1,"syntology":null},{"paper":"/paper/descriptor-and-word-soups-overcoming-the","slug":"descriptor-and-word-soups-overcoming-the","title":"Descriptor and Word Soups: Overcoming the Parameter Efficiency Accuracy Tradeoff for Out-of-Distribution Few-shot Learning","date":"2023-11-21","arxiv_id":"2311.13612","n_code_links":1,"syntology":null},{"paper":"/paper/extracting-definienda-in-mathematical","slug":"extracting-definienda-in-mathematical","title":"Extracting Definienda in Mathematical Scholarly Articles with Transformers","date":"2023-11-21","arxiv_id":"2311.12448","n_code_links":2,"syntology":null},{"paper":null,"slug":"gpt4motion-scripting-physical-motions-in-text","title":"GPT4Motion: Scripting Physical Motions in Text-to-Video Generation via Blender-Oriented GPT Planning","date":"2023-11-21","arxiv_id":"2311.12631","n_code_links":0,"syntology":null},{"paper":null,"slug":"interprompt-interpretable-prompting-for","title":"InterPrompt: Interpretable Prompting for Interrelated Interpersonal Risk Factors in Reddit Posts","date":"2023-11-21","arxiv_id":"2311.12404","n_code_links":0,"syntology":null},{"paper":"/paper/lowresource-at-blp-2023-task-2-leveraging","slug":"lowresource-at-blp-2023-task-2-leveraging","title":"LowResource at BLP-2023 Task 2: Leveraging BanglaBert for Low Resource Sentiment Analysis of Bangla Language","date":"2023-11-21","arxiv_id":"2311.12735","n_code_links":1,"syntology":null},{"paper":null,"slug":"utilizing-language-models-for-tour-itinerary","title":"Utilizing Language Models for Tour Itinerary Recommendation","date":"2023-11-21","arxiv_id":"2311.12355","n_code_links":0,"syntology":null},{"paper":"/paper/assessing-prompt-injection-risks-in-200","slug":"assessing-prompt-injection-risks-in-200","title":"Assessing Prompt Injection Risks in 200+ Custom GPTs","date":"2023-11-20","arxiv_id":"2311.11538","n_code_links":1,"syntology":null},{"paper":"/paper/evil-geniuses-delving-into-the-safety-of-llm","slug":"evil-geniuses-delving-into-the-safety-of-llm","title":"Evil Geniuses: Delving into the Safety of LLM-based Agents","date":"2023-11-20","arxiv_id":"2311.11855","n_code_links":1,"syntology":{"ran":5,"of":5,"n_ran_checked":5,"n_instrument":0,"unverified":0,"pointer_only":5,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["t1ans1r/evil-geniuses"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/how-to-use-large-language-models-for-text","slug":"how-to-use-large-language-models-for-text","title":"Towards Human-Level Text Coding with LLMs: The Case of Fatherhood Roles in Public Policy Documents","date":"2023-11-20","arxiv_id":"2311.11844","n_code_links":1,"syntology":null},{"paper":"/paper/loglead-fast-and-integrated-log-loader","slug":"loglead-fast-and-integrated-log-loader","title":"LogLead -- Fast and Integrated Log Loader, Enhancer, and Anomaly Detector","date":"2023-11-20","arxiv_id":"2311.11809","n_code_links":1,"syntology":null},{"paper":"/paper/lq-lora-low-rank-plus-quantized-matrix","slug":"lq-lora-low-rank-plus-quantized-matrix","title":"LQ-LoRA: Low-rank Plus Quantized Matrix Decomposition for Efficient Language Model Finetuning","date":"2023-11-20","arxiv_id":"2311.12023","n_code_links":1,"syntology":{"ran":9,"of":12,"n_ran_checked":9,"n_instrument":0,"unverified":3,"pointer_only":0,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 9 with no instrument failure: 0 honoured, 0 violated, 9 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":{"repos":["hanguo97/lq-lora"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":9,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"memorycompanion-a-smart-healthcare-solution","title":"MemoryCompanion: A Smart Healthcare Solution to Empower Efficient Alzheimer's Care Via Unleashing Generative AI","date":"2023-11-20","arxiv_id":"2311.14730","n_code_links":0,"syntology":null},{"paper":null,"slug":"refactoring-programs-using-large-language","title":"Refactoring Programs Using Large Language Models with Few-Shot Examples","date":"2023-11-20","arxiv_id":"2311.11690","n_code_links":0,"syntology":null},{"paper":null,"slug":"spot-the-bot-distinguishing-human-written-and","title":"Spot the Bot: Distinguishing Human-Written and Bot-Generated Texts Using Clustering and Information Theory Techniques","date":"2023-11-19","arxiv_id":"2311.11441","n_code_links":0,"syntology":null},{"paper":"/paper/tensor-aware-energy-accounting","slug":"tensor-aware-energy-accounting","title":"Tensor-Aware Energy Accounting","date":"2023-11-19","arxiv_id":"2311.11424","n_code_links":1,"syntology":null},{"paper":null,"slug":"behavior-optimized-image-generation","title":"Behavior Optimized Image Generation","date":"2023-11-18","arxiv_id":"2311.10995","n_code_links":0,"syntology":null},{"paper":null,"slug":"bit-cipher-a-simple-yet-powerful-word","title":"Bit Cipher -- A Simple yet Powerful Word Representation System that Integrates Efficiently with Language Models","date":"2023-11-18","arxiv_id":"2311.11012","n_code_links":0,"syntology":null},{"paper":null,"slug":"compositional-fusion-of-signals-in-data","title":"Compositional Fusion of Signals in Data Embedding","date":"2023-11-18","arxiv_id":"2311.11085","n_code_links":0,"syntology":null},{"paper":"/paper/vashantor-a-large-scale-multilingual","slug":"vashantor-a-large-scale-multilingual","title":"Vashantor: A Large-scale Multilingual Benchmark Dataset for Automated Translation of Bangla Regional Dialects to Bangla Language","date":"2023-11-18","arxiv_id":"2311.11142","n_code_links":1,"syntology":null},{"paper":null,"slug":"advancements-in-generative-ai-a-comprehensive","title":"Advancements in Generative AI: A Comprehensive Review of GANs, GPT, Autoencoders, Diffusion Model, and Transformers","date":"2023-11-17","arxiv_id":"2311.10242","n_code_links":0,"syntology":null},{"paper":null,"slug":"bias-a-head-analyzing-bias-in-transformer","title":"Bias A-head? Analyzing Bias in Transformer-Based Language Model Attention Heads","date":"2023-11-17","arxiv_id":"2311.10395","n_code_links":0,"syntology":null},{"paper":"/paper/camels-in-a-changing-climate-enhancing-lm","slug":"camels-in-a-changing-climate-enhancing-lm","title":"Camels in a Changing Climate: Enhancing LM Adaptation with Tulu 2","date":"2023-11-17","arxiv_id":"2311.10702","n_code_links":3,"syntology":null},{"paper":"/paper/dynapipe-optimizing-multi-task-training","slug":"dynapipe-optimizing-multi-task-training","title":"DynaPipe: Optimizing Multi-task Training through Dynamic Pipelines","date":"2023-11-17","arxiv_id":"2311.10418","n_code_links":2,"syntology":null},{"paper":null,"slug":"extracting-periodontitis-diagnosis-in","title":"Extracting periodontitis diagnosis in clinical notes with RoBERTa and regular expression","date":"2023-11-17","arxiv_id":"2311.10809","n_code_links":0,"syntology":null},{"paper":"/paper/hashing-it-out-predicting-unhealthy","slug":"hashing-it-out-predicting-unhealthy","title":"Hashing it Out: Predicting Unhealthy Conversations on Twitter","date":"2023-11-17","arxiv_id":"2311.10596","n_code_links":1,"syntology":null},{"paper":null,"slug":"use-gpt-j-prompt-generation-with-roberta-for","title":"Use GPT-J Prompt Generation with RoBERTa for NER Models on Diagnosis Extraction of Periodontal Diagnosis from Electronic Dental Records","date":"2023-11-17","arxiv_id":"2311.10810","n_code_links":0,"syntology":null},{"paper":"/paper/ares-an-automated-evaluation-framework-for","slug":"ares-an-automated-evaluation-framework-for","title":"ARES: An Automated Evaluation Framework for Retrieval-Augmented Generation Systems","date":"2023-11-16","arxiv_id":"2311.09476","n_code_links":1,"syntology":{"ran":13,"of":15,"n_ran_checked":13,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"13 ran (of which 0 constructed an object rather than computing a result; 13 with no instrument failure: 0 honoured, 0 violated, 13 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["stanford-futuredata/ares"],"state":"official (archive's flag): 13 ran","n_ran":13,"n_constructed":0,"n_ran_no_instrument_failure":13,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/event-causality-is-key-to-computational-story","slug":"event-causality-is-key-to-computational-story","title":"Event Causality Is Key to Computational Story Understanding","date":"2023-11-16","arxiv_id":"2311.09648","n_code_links":1,"syntology":{"ran":6,"of":9,"n_ran_checked":6,"n_instrument":0,"unverified":3,"pointer_only":9,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":{"repos":["insundaycathy/event-causality-extraction"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"fumbling-in-babel-an-investigation-into","title":"Fumbling in Babel: An Investigation into ChatGPT's Language Identification Ability","date":"2023-11-16","arxiv_id":"2311.09696","n_code_links":0,"syntology":null},{"paper":null,"slug":"generative-ai-for-hate-speech-detection","title":"Generative AI for Hate Speech Detection: Evaluation and Findings","date":"2023-11-16","arxiv_id":"2311.09993","n_code_links":0,"syntology":null},{"paper":null,"slug":"human-still-wins-over-llm-an-empirical-study","title":"Human Still Wins over LLM: An Empirical Study of Active Learning on Domain-Specific Annotation Tasks","date":"2023-11-16","arxiv_id":"2311.09825","n_code_links":0,"syntology":null},{"paper":"/paper/intervenor-prompt-the-coding-ability-of-large","slug":"intervenor-prompt-the-coding-ability-of-large","title":"INTERVENOR: Prompting the Coding Ability of Large Language Models with the Interactive Chain of Repair","date":"2023-11-16","arxiv_id":"2311.09868","n_code_links":1,"syntology":null},{"paper":"/paper/knowledgemath-knowledge-intensive-math-word","slug":"knowledgemath-knowledge-intensive-math-word","title":"FinanceMath: Knowledge-Intensive Math Reasoning in Finance Domains","date":"2023-11-16","arxiv_id":"2311.09797","n_code_links":1,"syntology":{"ran":12,"of":15,"n_ran_checked":12,"n_instrument":0,"unverified":3,"pointer_only":15,"phrase":"12 ran (of which 0 constructed an object rather than computing a result; 12 with no instrument failure: 0 honoured, 0 violated, 12 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":{"repos":["yale-nlp/knowledgemath"],"state":"official (archive's flag): 12 ran","n_ran":12,"n_constructed":0,"n_ran_no_instrument_failure":12,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":"/paper/lifetox-unveiling-implicit-toxicity-in-life","slug":"lifetox-unveiling-implicit-toxicity-in-life","title":"LifeTox: Unveiling Implicit Toxicity in Life Advice","date":"2023-11-16","arxiv_id":"2311.09585","n_code_links":1,"syntology":null},{"paper":null,"slug":"on-retrieval-augmentation-and-the-limitations","title":"On Retrieval Augmentation and the Limitations of Language Model Training","date":"2023-11-16","arxiv_id":"2311.09615","n_code_links":0,"syntology":null},{"paper":null,"slug":"predictive-minds-llms-as-atypical-active","title":"Predictive Minds: LLMs As Atypical Active Inference Agents","date":"2023-11-16","arxiv_id":"2311.10215","n_code_links":0,"syntology":null},{"paper":null,"slug":"reducing-privacy-risks-in-online-self","title":"Reducing Privacy Risks in Online Self-Disclosures with Language Models","date":"2023-11-16","arxiv_id":"2311.09538","n_code_links":0,"syntology":null},{"paper":null,"slug":"sequencing-matters-a-generate-retrieve","title":"Sequencing Matters: A Generate-Retrieve-Generate Model for Building Conversational Agents","date":"2023-11-16","arxiv_id":"2311.09513","n_code_links":0,"syntology":null},{"paper":null,"slug":"towards-an-automatic-ai-agent-for-reaction","title":"Chemist-X: Large Language Model-empowered Agent for Reaction Condition Recommendation in Chemical Synthesis","date":"2023-11-16","arxiv_id":"2311.10776","n_code_links":0,"syntology":null},{"paper":"/paper/an-eye-on-clinical-bert-investigating","slug":"an-eye-on-clinical-bert-investigating","title":"An Eye on Clinical BERT: Investigating Language Model Generalization for Diabetic Eye Disease Phenotyping","date":"2023-11-15","arxiv_id":"2311.08687","n_code_links":1,"syntology":null},{"paper":"/paper/can-large-language-models-follow-concept","slug":"can-large-language-models-follow-concept","title":"Can Large Language Models Follow Concept Annotation Guidelines? A Case Study on Scientific and Financial Domains","date":"2023-11-15","arxiv_id":"2311.08704","n_code_links":1,"syntology":null},{"paper":null,"slug":"evaluating-gender-bias-in-the-translation-of","title":"Evaluating Gender Bias in the Translation of Gender-Neutral Languages into English","date":"2023-11-15","arxiv_id":"2311.08836","n_code_links":0,"syntology":null},{"paper":"/paper/exponentially-faster-language-modelling","slug":"exponentially-faster-language-modelling","title":"Exponentially Faster Language Modelling","date":"2023-11-15","arxiv_id":"2311.10770","n_code_links":3,"syntology":null},{"paper":null,"slug":"german-finbert-a-german-pre-trained-language","title":"German FinBERT: A German Pre-trained Language Model","date":"2023-11-15","arxiv_id":"2311.08793","n_code_links":0,"syntology":null},{"paper":null,"slug":"loke-linked-open-knowledge-extraction-for","title":"LOKE: Linked Open Knowledge Extraction for Automated Knowledge Graph Construction","date":"2023-11-15","arxiv_id":"2311.09366","n_code_links":0,"syntology":null},{"paper":null,"slug":"memory-augmented-language-models-through","title":"Memory Augmented Language Models through Mixture of Word Experts","date":"2023-11-15","arxiv_id":"2311.10768","n_code_links":0,"syntology":null},{"paper":"/paper/tooltalk-evaluating-tool-usage-in-a","slug":"tooltalk-evaluating-tool-usage-in-a","title":"ToolTalk: Evaluating Tool-Usage in a Conversational Setting","date":"2023-11-15","arxiv_id":"2311.10775","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":null}},{"paper":"/paper/we-demand-justice-towards-grounding-political","slug":"we-demand-justice-towards-grounding-political","title":"\"We Demand Justice!\": Towards Social Context Grounding of Political Texts","date":"2023-11-15","arxiv_id":"2311.09106","n_code_links":1,"syntology":null},{"paper":"/paper/xplainllm-a-qa-explanation-dataset-for","slug":"xplainllm-a-qa-explanation-dataset-for","title":"XplainLLM: A Knowledge-Augmented Dataset for Reliable Grounded Explanations in LLMs","date":"2023-11-15","arxiv_id":"2311.08614","n_code_links":1,"syntology":null},{"paper":"/paper/a-survey-on-language-models-for-code","slug":"a-survey-on-language-models-for-code","title":"Unifying the Perspectives of NLP and Software Engineering: A Survey on Language Models for Code","date":"2023-11-14","arxiv_id":"2311.07989","n_code_links":1,"syntology":null},{"paper":"/paper/artificial-text-boundary-detection-with","slug":"artificial-text-boundary-detection-with","title":"AI-generated text boundary detection with RoFT","date":"2023-11-14","arxiv_id":"2311.08349","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["silversolver/ai_boundary_detection"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"cpopqa-ranking-cultural-concept-popularity-by","title":"CPopQA: Ranking Cultural Concept Popularity by LLMs","date":"2023-11-14","arxiv_id":"2311.07897","n_code_links":0,"syntology":null},{"paper":null,"slug":"evaluating-llms-on-document-based-qa-exact","title":"Evaluating LLMs on Document-Based QA: Exact Answer Selection and Numerical Extraction using Cogtale dataset","date":"2023-11-14","arxiv_id":"2311.07878","n_code_links":0,"syntology":null},{"paper":"/paper/exploring-semi-supervised-hierarchical","slug":"exploring-semi-supervised-hierarchical","title":"Exploring Semi-supervised Hierarchical Stacked Encoder for Legal Judgement Prediction","date":"2023-11-14","arxiv_id":"2311.08103","n_code_links":1,"syntology":null},{"paper":"/paper/fair-abstractive-summarization-of-diverse","slug":"fair-abstractive-summarization-of-diverse","title":"Fair Abstractive Summarization of Diverse Perspectives","date":"2023-11-14","arxiv_id":"2311.07884","n_code_links":1,"syntology":null},{"paper":null,"slug":"investigating-the-encoding-of-words-in-bert-s","title":"Investigating the Encoding of Words in BERT's Neurons using Feature Textualization","date":"2023-11-14","arxiv_id":"2311.08240","n_code_links":0,"syntology":null},{"paper":"/paper/language-models-are-better-bug-detector","slug":"language-models-are-better-bug-detector","title":"Language Models are Better Bug Detector Through Code-Pair Classification","date":"2023-11-14","arxiv_id":"2311.07957","n_code_links":1,"syntology":null},{"paper":null,"slug":"large-language-model-driven-classroom","title":"Large Language Model-Driven Classroom Flipping: Empowering Student-Centric Peer Questioning with Flipped Interaction","date":"2023-11-14","arxiv_id":"2311.14708","n_code_links":0,"syntology":null},{"paper":"/paper/memory-efficient-stochastic-methods-for","slug":"memory-efficient-stochastic-methods-for","title":"Memory-efficient Stochastic methods for Memory-based Transformers","date":"2023-11-14","arxiv_id":"2311.08123","n_code_links":1,"syntology":null},{"paper":"/paper/spot-a-natural-language-interface-for","slug":"spot-a-natural-language-interface-for","title":"Spot: A Natural Language Interface for Geospatial Searches in OSM","date":"2023-11-14","arxiv_id":"2311.08093","n_code_links":1,"syntology":null},{"paper":null,"slug":"ut5-pretraining-non-autoregressive-t5-with","title":"UT5: Pretraining Non autoregressive T5 with unrolled denoising","date":"2023-11-14","arxiv_id":"2311.08552","n_code_links":0,"syntology":null},{"paper":"/paper/do-large-language-models-and-humans-have","slug":"do-large-language-models-and-humans-have","title":"Do large language models and humans have similar behaviors in causal inference with script knowledge?","date":"2023-11-13","arxiv_id":"2311.07311","n_code_links":1,"syntology":null},{"paper":"/paper/in-context-learning-generalizes-but-not","slug":"in-context-learning-generalizes-but-not","title":"In-context Learning Generalizes, But Not Always Robustly: The Case of Syntax","date":"2023-11-13","arxiv_id":"2311.07811","n_code_links":1,"syntology":null},{"paper":"/paper/it-s-not-easy-being-wrong-evaluating-process","slug":"it-s-not-easy-being-wrong-evaluating-process","title":"It's Not Easy Being Wrong: Large Language Models Struggle with Process of Elimination Reasoning","date":"2023-11-13","arxiv_id":"2311.07532","n_code_links":1,"syntology":null},{"paper":null,"slug":"language-model-in-the-loop-data-optimal","title":"Language Model-In-The-Loop: Data Optimal Approach to Learn-To-Recommend Actions in Text Games","date":"2023-11-13","arxiv_id":"2311.07687","n_code_links":0,"syntology":null},{"paper":null,"slug":"megaverse-benchmarking-large-language-models","title":"MEGAVERSE: Benchmarking Large Language Models Across Languages, Modalities, Models and Tasks","date":"2023-11-13","arxiv_id":"2311.07463","n_code_links":0,"syntology":null},{"paper":null,"slug":"on-the-truthfulness-of-surprisingly-likely","title":"On The Truthfulness of 'Surprisingly Likely' Responses of Large Language Models","date":"2023-11-13","arxiv_id":"2311.07692","n_code_links":0,"syntology":null}],"record_sha256":"f27c26ee5f9e99c6cee4fa28f588a592f581805d00fa3c7c3ae3f5337eb6ebaf","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}