{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/attention-dropout/papers/28","list_of":"/method/attention-dropout","method":"Attention Dropout","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":28,"pages_in_order":109,"rows_per_page":100,"rows":[2701,2800],"of":10892,"counts":{"archive_papers_tagged":10892,"with_a_code_link":4634,"where_syntology_ran_a_sample":1270,"not_listed_spam_title":0,"listed":10892,"listed_where_code_ran":1270,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":1043,"every_run_a_failure_of_syntologys_instrument":227,"listed_with_a_run_with_no_instrument_failure":1043,"listed_every_run_a_failure_of_syntologys_instrument":227,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/attention-dropout","prev":"/method/attention-dropout/papers/27","next":"/method/attention-dropout/papers/29","papers":[{"paper":"/paper/vid-gpt-introducing-gpt-style-autoregressive","slug":"vid-gpt-introducing-gpt-style-autoregressive","title":"ViD-GPT: Introducing GPT-style Autoregressive Generation in Video Diffusion Models","date":"2024-06-16","arxiv_id":"2406.10981","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-comprehensive-survey-of-foundation-models","title":"A Comprehensive Survey of Foundation Models in Medicine","date":"2024-06-15","arxiv_id":"2406.10729","n_code_links":0,"syntology":null},{"paper":"/paper/beyond-raw-videos-understanding-edited-videos","slug":"beyond-raw-videos-understanding-edited-videos","title":"Beyond Raw Videos: Understanding Edited Videos with Large Multimodal Model","date":"2024-06-15","arxiv_id":"2406.10484","n_code_links":1,"syntology":null},{"paper":"/paper/mint-a-multi-modal-image-and-narrative-text","slug":"mint-a-multi-modal-image-and-narrative-text","title":"MINT: a Multi-modal Image and Narrative Text Dubbing Dataset for Foley Audio Content Planning and Generation","date":"2024-06-15","arxiv_id":"2406.10591","n_code_links":1,"syntology":null},{"paper":null,"slug":"we-care-multimodal-depression-detection-and","title":"We Care: Multimodal Depression Detection and Knowledge Infused Mental Health Therapeutic Response Generation","date":"2024-06-15","arxiv_id":"2406.10561","n_code_links":0,"syntology":null},{"paper":null,"slug":"bag-of-lies-robustness-in-continuous-pre","title":"Bag of Lies: Robustness in Continuous Pre-training BERT","date":"2024-06-14","arxiv_id":"2406.09967","n_code_links":0,"syntology":null},{"paper":null,"slug":"exploring-the-correlation-between-human-and","title":"Exploring the Correlation between Human and Machine Evaluation of Simultaneous Speech Translation","date":"2024-06-14","arxiv_id":"2406.10091","n_code_links":0,"syntology":null},{"paper":"/paper/hiro-hierarchical-information-retrieval","slug":"hiro-hierarchical-information-retrieval","title":"HIRO: Hierarchical Information Retrieval Optimization","date":"2024-06-14","arxiv_id":"2406.09979","n_code_links":1,"syntology":null},{"paper":"/paper/liere-generalizing-rotary-position-encodings","slug":"liere-generalizing-rotary-position-encodings","title":"LieRE: Generalizing Rotary Position Encodings","date":"2024-06-14","arxiv_id":"2406.10322","n_code_links":1,"syntology":{"ran":1,"of":3,"n_ran_checked":1,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["stanford-aimi/liere"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/the-devil-is-in-the-neurons-interpreting-and","slug":"the-devil-is-in-the-neurons-interpreting-and","title":"The Devil is in the Neurons: Interpreting and Mitigating Social Biases in Pre-trained Language Models","date":"2024-06-14","arxiv_id":"2406.10130","n_code_links":1,"syntology":null},{"paper":"/paper/towards-efficient-pareto-set-approximation","slug":"towards-efficient-pareto-set-approximation","title":"Towards Efficient Pareto Set Approximation via Mixture of Experts Based Model Fusion","date":"2024-06-14","arxiv_id":"2406.09770","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-more-practical-approach-to-machine","title":"A More Practical Approach to Machine Unlearning","date":"2024-06-13","arxiv_id":"2406.09391","n_code_links":0,"syntology":null},{"paper":null,"slug":"analyzing-gender-polarity-in-short-social","title":"Analyzing Gender Polarity in Short Social Media Texts with BERT: The Role of Emojis and Emoticons","date":"2024-06-13","arxiv_id":"2406.09573","n_code_links":0,"syntology":null},{"paper":"/paper/bpe-knockout-pruning-pre-existing-bpe","slug":"bpe-knockout-pruning-pre-existing-bpe","title":"BPE-knockout: Pruning Pre-existing BPE Tokenisers with Backwards-compatible Morphological Semi-supervision","date":"2024-06-13","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"linguistic-bias-in-chatgpt-language-models","title":"Linguistic Bias in ChatGPT: Language Models Reinforce Dialect Discrimination","date":"2024-06-13","arxiv_id":"2406.08818","n_code_links":0,"syntology":null},{"paper":null,"slug":"optimizing-large-model-training-through","title":"Optimizing Large Model Training through Overlapped Activation Recomputation","date":"2024-06-13","arxiv_id":"2406.08756","n_code_links":0,"syntology":null},{"paper":null,"slug":"pc-lora-low-rank-adaptation-for-progressive","title":"PC-LoRA: Low-Rank Adaptation for Progressive Model Compression with Knowledge Distillation","date":"2024-06-13","arxiv_id":"2406.09117","n_code_links":0,"syntology":null},{"paper":null,"slug":"talking-heads-understanding-inter-layer","title":"Talking Heads: Understanding Inter-layer Communication in Transformer Language Models","date":"2024-06-13","arxiv_id":"2406.09519","n_code_links":0,"syntology":null},{"paper":null,"slug":"ad-auctions-for-llms-via-retrieval-augmented","title":"Ad Auctions for LLMs via Retrieval Augmented Generation","date":"2024-06-12","arxiv_id":"2406.09459","n_code_links":0,"syntology":null},{"paper":null,"slug":"exploring-fact-memorization-and-style","title":"Exploring Fact Memorization and Style Imitation in LLMs Using QLoRA: An Experimental Study and Quality Assessment Methods","date":"2024-06-12","arxiv_id":"2406.08582","n_code_links":0,"syntology":null},{"paper":null,"slug":"faithfill-faithful-inpainting-for-object","title":"FaithFill: Faithful Inpainting for Object Completion Using a Single Reference Image","date":"2024-06-12","arxiv_id":"2406.07865","n_code_links":0,"syntology":null},{"paper":"/paper/fine-tuned-small-llms-still-significantly","slug":"fine-tuned-small-llms-still-significantly","title":"Fine-Tuned 'Small' LLMs (Still) Significantly Outperform Zero-Shot Generative AI Models in Text Classification","date":"2024-06-12","arxiv_id":"2406.08660","n_code_links":1,"syntology":{"ran":5,"of":9,"n_ran_checked":5,"n_instrument":0,"unverified":4,"pointer_only":9,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","official":{"repos":["mnbucher/text-cls-llms"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"how-well-it-works-benchmarking-performance-of","title":"How well it works: Benchmarking performance of GPT models on medical natural language processing tasks","date":"2024-06-12","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/indirectrequests-making-task-oriented","slug":"indirectrequests-making-task-oriented","title":"Making Task-Oriented Dialogue Datasets More Natural by Synthetically Generating Indirect User Requests","date":"2024-06-12","arxiv_id":"2406.07794","n_code_links":0,"syntology":null},{"paper":"/paper/label-aware-hard-negative-sampling-strategies","slug":"label-aware-hard-negative-sampling-strategies","title":"Label-aware Hard Negative Sampling Strategies with Momentum Contrastive Learning for Implicit Hate Speech Detection","date":"2024-06-12","arxiv_id":"2406.07886","n_code_links":1,"syntology":null},{"paper":null,"slug":"leveraging-large-language-models-for-web","title":"Leveraging Large Language Models for Web Scraping","date":"2024-06-12","arxiv_id":"2406.08246","n_code_links":0,"syntology":null},{"paper":null,"slug":"multimodal-representation-loss-between-timed","title":"Multimodal Representation Loss Between Timed Text and Audio for Regularized Speech Separation","date":"2024-06-12","arxiv_id":"2406.08328","n_code_links":0,"syntology":null},{"paper":"/paper/tailoring-generative-ai-chatbots-for","slug":"tailoring-generative-ai-chatbots-for","title":"Tailoring Generative AI Chatbots for Multiethnic Communities in Disaster Preparedness Communication: Extending the CASA Paradigm","date":"2024-06-12","arxiv_id":"2406.08411","n_code_links":1,"syntology":null},{"paper":null,"slug":"beyond-words-on-large-language-models","title":"Beyond Words: On Large Language Models Actionability in Mission-Critical Risk Analysis","date":"2024-06-11","arxiv_id":"2406.10273","n_code_links":0,"syntology":null},{"paper":null,"slug":"bilingual-sexism-classification-fine-tuned","title":"Bilingual Sexism Classification: Fine-Tuned XLM-RoBERTa and GPT-3.5 Few-Shot Learning","date":"2024-06-11","arxiv_id":"2406.07287","n_code_links":0,"syntology":null},{"paper":null,"slug":"covid-19-twitter-sentiment-classification","title":"COVID-19 Twitter Sentiment Classification Using Hybrid Deep Learning Model Based on Grid Search Methodology","date":"2024-06-11","arxiv_id":"2406.10266","n_code_links":0,"syntology":null},{"paper":null,"slug":"dr-rag-applying-dynamic-document-relevance-to","title":"DR-RAG: Applying Dynamic Document Relevance to Retrieval-Augmented Generation for Question-Answering","date":"2024-06-11","arxiv_id":"2406.07348","n_code_links":0,"syntology":null},{"paper":null,"slug":"flextron-many-in-one-flexible-large-language","title":"Flextron: Many-in-One Flexible Large Language Model","date":"2024-06-11","arxiv_id":"2406.10260","n_code_links":0,"syntology":null},{"paper":"/paper/multi-objective-reinforcement-learning-from","slug":"multi-objective-reinforcement-learning-from","title":"Multi-objective Reinforcement learning from AI Feedback","date":"2024-06-11","arxiv_id":"2406.07295","n_code_links":1,"syntology":null},{"paper":"/paper/multimodal-belief-prediction","slug":"multimodal-belief-prediction","title":"Multimodal Belief Prediction","date":"2024-06-11","arxiv_id":"2406.07466","n_code_links":1,"syntology":null},{"paper":null,"slug":"question-answering-qa-model-for-a","title":"Question-Answering (QA) Model for a Personalized Learning Assistant for Arabic Language","date":"2024-06-11","arxiv_id":"2406.08519","n_code_links":0,"syntology":null},{"paper":"/paper/unused-information-in-token-probability","slug":"unused-information-in-token-probability","title":"Unused information in token probability distribution of generative LLM: improving LLM reading comprehension through calculation of expected values","date":"2024-06-11","arxiv_id":"2406.10267","n_code_links":1,"syntology":null},{"paper":"/paper/agb-de-a-corpus-for-the-automated-legal","slug":"agb-de-a-corpus-for-the-automated-legal","title":"AGB-DE: A Corpus for the Automated Legal Assessment of Clauses in German Consumer Contracts","date":"2024-06-10","arxiv_id":"2406.06809","n_code_links":1,"syntology":null},{"paper":"/paper/compute-better-spent-replacing-dense-layers","slug":"compute-better-spent-replacing-dense-layers","title":"Compute Better Spent: Replacing Dense Layers with Structured Matrices","date":"2024-06-10","arxiv_id":"2406.06248","n_code_links":1,"syntology":{"ran":4,"of":5,"n_ran_checked":4,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["shikaiqiu/compute-better-spent"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/in-context-learning-and-fine-tuning-gpt-for","slug":"in-context-learning-and-fine-tuning-gpt-for","title":"In-Context Learning and Fine-Tuning GPT for Argument Mining","date":"2024-06-10","arxiv_id":"2406.06699","n_code_links":1,"syntology":null},{"paper":null,"slug":"leveraging-large-language-models-for-12","title":"Leveraging Large Language Models for Knowledge-free Weak Supervision in Clinical Natural Language Processing","date":"2024-06-10","arxiv_id":"2406.06723","n_code_links":0,"syntology":null},{"paper":null,"slug":"llm-dcache-improving-tool-augmented-llms-with","title":"LLM-dCache: Improving Tool-Augmented LLMs with GPT-Driven Localized Data Caching","date":"2024-06-10","arxiv_id":"2406.06799","n_code_links":0,"syntology":null},{"paper":null,"slug":"securenet-a-comparative-study-of-deberta-and","title":"SecureNet: A Comparative Study of DeBERTa and Large Language Models for Phishing Detection","date":"2024-06-10","arxiv_id":"2406.06663","n_code_links":0,"syntology":null},{"paper":null,"slug":"the-impact-of-quantization-on-retrieval","title":"The Impact of Quantization on Retrieval-Augmented Generation: An Analysis of Small LLMs","date":"2024-06-10","arxiv_id":"2406.10251","n_code_links":0,"syntology":null},{"paper":"/paper/umbrela-umbrela-is-the-open-source","slug":"umbrela-umbrela-is-the-open-source","title":"UMBRELA: UMbrela is the (Open-Source Reproduction of the) Bing RELevance Assessor","date":"2024-06-10","arxiv_id":"2406.06519","n_code_links":1,"syntology":null},{"paper":"/paper/domainrag-a-chinese-benchmark-for-evaluating","slug":"domainrag-a-chinese-benchmark-for-evaluating","title":"DomainRAG: A Chinese Benchmark for Evaluating Domain-specific Retrieval-Augmented Generation","date":"2024-06-09","arxiv_id":"2406.05654","n_code_links":2,"syntology":{"ran":12,"of":12,"n_ran_checked":11,"n_instrument":1,"unverified":0,"pointer_only":12,"phrase":"12 ran (of which 0 constructed an object rather than computing a result; 11 with no instrument failure: 0 honoured, 0 violated, 11 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["ShootingWong/DomainRAG"],"state":"official (archive's flag): 12 ran","n_ran":12,"n_constructed":0,"n_ran_no_instrument_failure":11,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"hidden-holes-topological-aspects-of-language","title":"Hidden Holes: topological aspects of language models","date":"2024-06-09","arxiv_id":"2406.05798","n_code_links":0,"syntology":null},{"paper":null,"slug":"machine-against-the-rag-jamming-retrieval","title":"Machine Against the RAG: Jamming Retrieval-Augmented Generation with Blocker Documents","date":"2024-06-09","arxiv_id":"2406.05870","n_code_links":0,"syntology":null},{"paper":null,"slug":"medreqal-examining-medical-knowledge-recall","title":"MedREQAL: Examining Medical Knowledge Recall of Large Language Models via Question Answering","date":"2024-06-09","arxiv_id":"2406.05845","n_code_links":0,"syntology":null},{"paper":"/paper/re-rag-improving-open-domain-qa-performance","slug":"re-rag-improving-open-domain-qa-performance","title":"RE-RAG: Improving Open-Domain QA Performance and Interpretability with Relevance Estimator in Retrieval-Augmented Generation","date":"2024-06-09","arxiv_id":"2406.05794","n_code_links":1,"syntology":null},{"paper":null,"slug":"text2vp-generative-ai-for-visual-programming","title":"Text2VP: Generative AI for Visual Programming and Parametric Modeling","date":"2024-06-09","arxiv_id":"2407.07732","n_code_links":0,"syntology":null},{"paper":"/paper/advancing-semantic-textual-similarity","slug":"advancing-semantic-textual-similarity","title":"Advancing Semantic Textual Similarity Modeling: A Regression Framework with Translated ReLU and Smooth K2 Loss","date":"2024-06-08","arxiv_id":"2406.05326","n_code_links":2,"syntology":null},{"paper":null,"slug":"concept-formation-and-alignment-in-language","title":"Concept Formation and Alignment in Language Models: Bridging Statistical Patterns in Latent Space to Concept Taxonomy","date":"2024-06-08","arxiv_id":"2406.05315","n_code_links":0,"syntology":null},{"paper":null,"slug":"critical-phase-transition-in-a-large-language","title":"Critical Phase Transition in Large Language Models","date":"2024-06-08","arxiv_id":"2406.05335","n_code_links":0,"syntology":null},{"paper":null,"slug":"do-llms-recognize-me-when-i-is-not-me","title":"Do LLMs Recognize me, When I is not me: Assessment of LLMs Understanding of Turkish Indexical Pronouns in Indexical Shift Contexts","date":"2024-06-08","arxiv_id":"2406.05569","n_code_links":0,"syntology":null},{"paper":null,"slug":"matablegpt-gpt-based-table-data-extractor","title":"MaTableGPT: GPT-based Table Data Extractor from Materials Science Literature","date":"2024-06-08","arxiv_id":"2406.05431","n_code_links":0,"syntology":null},{"paper":null,"slug":"selfdefend-llms-can-defend-themselves-against","title":"SelfDefend: LLMs Can Defend Themselves against Jailbreaking in a Practical Manner","date":"2024-06-08","arxiv_id":"2406.05498","n_code_links":0,"syntology":null},{"paper":null,"slug":"vp-llm-text-driven-3d-volume-completion-with","title":"VP-LLM: Text-Driven 3D Volume Completion with Large Language Models through Patchification","date":"2024-06-08","arxiv_id":"2406.05543","n_code_links":0,"syntology":null},{"paper":"/paper/bamo-at-semeval-2024-task-9-brainteaser-a","slug":"bamo-at-semeval-2024-task-9-brainteaser-a","title":"BAMO at SemEval-2024 Task 9: BRAINTEASER: A Novel Task Defying Common Sense","date":"2024-06-07","arxiv_id":"2406.04947","n_code_links":1,"syntology":null},{"paper":"/paper/berts-are-generative-in-context-learners","slug":"berts-are-generative-in-context-learners","title":"BERTs are Generative In-Context Learners","date":"2024-06-07","arxiv_id":"2406.04823","n_code_links":1,"syntology":{"ran":13,"of":26,"n_ran_checked":12,"n_instrument":1,"unverified":13,"pointer_only":0,"phrase":"13 ran (of which 0 constructed an object rather than computing a result; 12 with no instrument failure: 0 honoured, 1 violated, 11 with no contract checked; 1 where Syntology's instrument failed) · 13 unverified","official":{"repos":["ltgoslo/bert-in-context"],"state":"official (archive's flag): 13 ran","n_ran":13,"n_constructed":0,"n_ran_no_instrument_failure":12,"n_unverified":13,"ran_from_kinds":["official"]}}},{"paper":"/paper/corpus-poisoning-via-approximate-greedy","slug":"corpus-poisoning-via-approximate-greedy","title":"Corpus Poisoning via Approximate Greedy Gradient Descent","date":"2024-06-07","arxiv_id":"2406.05087","n_code_links":1,"syntology":null},{"paper":"/paper/crag-comprehensive-rag-benchmark","slug":"crag-comprehensive-rag-benchmark","title":"CRAG -- Comprehensive RAG Benchmark","date":"2024-06-07","arxiv_id":"2406.04744","n_code_links":2,"syntology":{"ran":4,"of":4,"n_ran_checked":3,"n_instrument":1,"unverified":0,"pointer_only":4,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["facebookresearch/crag"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/diner-a-large-realistic-dataset-for","slug":"diner-a-large-realistic-dataset-for","title":"DiNeR: a Large Realistic Dataset for Evaluating Compositional Generalization","date":"2024-06-07","arxiv_id":"2406.04669","n_code_links":1,"syntology":null},{"paper":"/paper/gamebench-evaluating-strategic-reasoning","slug":"gamebench-evaluating-strategic-reasoning","title":"GameBench: Evaluating Strategic Reasoning Abilities of LLM Agents","date":"2024-06-07","arxiv_id":"2406.06613","n_code_links":2,"syntology":{"ran":2,"of":2,"n_ran_checked":2,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["Joshuaclymer/GameBench"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"large-generative-graph-models","title":"Large Generative Graph Models","date":"2024-06-07","arxiv_id":"2406.05109","n_code_links":0,"syntology":null},{"paper":null,"slug":"low-resource-cross-lingual-summarization","title":"Low-Resource Cross-Lingual Summarization through Few-Shot Learning with Large Language Models","date":"2024-06-07","arxiv_id":"2406.04630","n_code_links":0,"syntology":null},{"paper":"/paper/multi-head-rag-solving-multi-aspect-problems","slug":"multi-head-rag-solving-multi-aspect-problems","title":"Multi-Head RAG: Solving Multi-Aspect Problems with LLMs","date":"2024-06-07","arxiv_id":"2406.05085","n_code_links":2,"syntology":null},{"paper":"/paper/on-subjective-uncertainty-quantification-and","slug":"on-subjective-uncertainty-quantification-and","title":"On Subjective Uncertainty Quantification and Calibration in Natural Language Generation","date":"2024-06-07","arxiv_id":"2406.05213","n_code_links":1,"syntology":{"ran":12,"of":14,"n_ran_checked":12,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"12 ran (of which 0 constructed an object rather than computing a result; 12 with no instrument failure: 0 honoured, 0 violated, 12 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["meta-inf/suq-nlg"],"state":"official (archive's flag): 12 ran","n_ran":12,"n_constructed":0,"n_ran_no_instrument_failure":12,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"vtrans-accelerating-transformer-compression","title":"VTrans: Accelerating Transformer Compression with Variational Information Bottleneck based Pruning","date":"2024-06-07","arxiv_id":"2406.05276","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-survey-on-medical-large-language-models","title":"A Survey on Medical Large Language Models: Technology, Application, Trustworthiness, and Future Directions","date":"2024-06-06","arxiv_id":"2406.03712","n_code_links":0,"syntology":null},{"paper":"/paper/do-language-models-understand-morality","slug":"do-language-models-understand-morality","title":"Do Language Models Understand Morality? Towards a Robust Detection of Moral Content","date":"2024-06-06","arxiv_id":"2406.04143","n_code_links":1,"syntology":null},{"paper":null,"slug":"empirical-guidelines-for-deploying-llms-onto","title":"Empirical Guidelines for Deploying LLMs onto Resource-constrained Edge Devices","date":"2024-06-06","arxiv_id":"2406.03777","n_code_links":0,"syntology":null},{"paper":"/paper/horae-a-domain-agnostic-modeling-language-for","slug":"horae-a-domain-agnostic-modeling-language-for","title":"HORAE: A Domain-Agnostic Language for Automated Service Regulation","date":"2024-06-06","arxiv_id":"2406.06600","n_code_links":1,"syntology":null},{"paper":"/paper/llmembed-rethinking-lightweight-llm-s-genuine","slug":"llmembed-rethinking-lightweight-llm-s-genuine","title":"LLMEmbed: Rethinking Lightweight LLM's Genuine Function in Text Classification","date":"2024-06-06","arxiv_id":"2406.03725","n_code_links":1,"syntology":null},{"paper":"/paper/simplified-and-generalized-masked-diffusion","slug":"simplified-and-generalized-masked-diffusion","title":"Simplified and Generalized Masked Diffusion for Discrete Data","date":"2024-06-06","arxiv_id":"2406.04329","n_code_links":1,"syntology":{"ran":6,"of":21,"n_ran_checked":2,"n_instrument":4,"unverified":15,"pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 1 honoured, 0 violated, 1 with no contract checked; 4 where Syntology's instrument failed) · 15 unverified","official":{"repos":["google-deepmind/md4"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":15,"ran_from_kinds":["official"]}}},{"paper":"/paper/tox-bart-leveraging-toxicity-attributes-for","slug":"tox-bart-leveraging-toxicity-attributes-for","title":"Tox-BART: Leveraging Toxicity Attributes for Explanation Generation of Implicit Hate Speech","date":"2024-06-06","arxiv_id":"2406.03953","n_code_links":1,"syntology":null},{"paper":"/paper/your-absorbing-discrete-diffusion-secretly","slug":"your-absorbing-discrete-diffusion-secretly","title":"Your Absorbing Discrete Diffusion Secretly Models the Conditional Distributions of Clean Data","date":"2024-06-06","arxiv_id":"2406.03736","n_code_links":2,"syntology":{"ran":13,"of":13,"n_ran_checked":12,"n_instrument":1,"unverified":0,"pointer_only":13,"phrase":"13 ran (of which 0 constructed an object rather than computing a result; 12 with no instrument failure: 1 honoured, 0 violated, 11 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["ml-gsai/radd"],"state":"official (archive's flag): 13 ran","n_ran":13,"n_constructed":0,"n_ran_no_instrument_failure":12,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/automating-turkish-educational-quiz","slug":"automating-turkish-educational-quiz","title":"Automating Turkish Educational Quiz Generation Using Large Language Models","date":"2024-06-05","arxiv_id":"2406.03397","n_code_links":4,"syntology":null},{"paper":null,"slug":"exact-conversion-of-in-context-learning-to","title":"Exact Conversion of In-Context Learning to Model Weights in Linearized-Attention Transformers","date":"2024-06-05","arxiv_id":"2406.02847","n_code_links":0,"syntology":null},{"paper":null,"slug":"exploring-multilingual-large-language-models","title":"Exploring Multilingual Large Language Models for Enhanced TNM classification of Radiology Report in lung cancer staging","date":"2024-06-05","arxiv_id":"2406.06591","n_code_links":0,"syntology":null},{"paper":"/paper/missci-reconstructing-fallacies-in","slug":"missci-reconstructing-fallacies-in","title":"Missci: Reconstructing Fallacies in Misrepresented Science","date":"2024-06-05","arxiv_id":"2406.03181","n_code_links":2,"syntology":null},{"paper":"/paper/polytc-a-novel-bert-based-classifier-to","slug":"polytc-a-novel-bert-based-classifier-to","title":"PoLYTC: a novel BERT-based classifier to detect political leaning of YouTube videos based on their titles","date":"2024-06-05","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/rico-reddit-ideological-communities","slug":"rico-reddit-ideological-communities","title":"RICo: Reddit ideological communities","date":"2024-06-05","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"statbot-swiss-bilingual-open-data-exploration","title":"StatBot.Swiss: Bilingual Open Data Exploration in Natural Language","date":"2024-06-05","arxiv_id":"2406.03170","n_code_links":0,"syntology":null},{"paper":null,"slug":"the-good-the-bad-and-the-hulk-like-gpt","title":"The Good, the Bad, and the Hulk-like GPT: Analyzing Emotional Decisions of Large Language Models in Cooperation and Bargaining Games","date":"2024-06-05","arxiv_id":"2406.03299","n_code_links":0,"syntology":null},{"paper":"/paper/too-big-to-fail-larger-language-models-are","slug":"too-big-to-fail-larger-language-models-are","title":"Too Big to Fail: Larger Language Models are Disproportionately Resilient to Induction of Dementia-Related Linguistic Anomalies","date":"2024-06-05","arxiv_id":"2406.02830","n_code_links":1,"syntology":null},{"paper":"/paper/chain-of-agents-large-language-models","slug":"chain-of-agents-large-language-models","title":"Chain of Agents: Large Language Models Collaborating on Long-Context Tasks","date":"2024-06-04","arxiv_id":"2406.02818","n_code_links":0,"syntology":null},{"paper":"/paper/checkembed-effective-verification-of-llm","slug":"checkembed-effective-verification-of-llm","title":"CheckEmbed: Effective Verification of LLM Solutions to Open-Ended Tasks","date":"2024-06-04","arxiv_id":"2406.02524","n_code_links":1,"syntology":null},{"paper":null,"slug":"large-language-model-enabled-multi-agent","title":"Large Language Model-Enabled Multi-Agent Manufacturing Systems","date":"2024-06-04","arxiv_id":"2406.01893","n_code_links":0,"syntology":null},{"paper":null,"slug":"occamllm-fast-and-exact-language-model","title":"OccamLLM: Fast and Exact Language Model Arithmetic in a Single Step","date":"2024-06-04","arxiv_id":"2406.06576","n_code_links":0,"syntology":null},{"paper":null,"slug":"probing-the-category-of-verbal-aspect-in","title":"Probing the Category of Verbal Aspect in Transformer Language Models","date":"2024-06-04","arxiv_id":"2406.02335","n_code_links":0,"syntology":null},{"paper":"/paper/randomized-geometric-algebra-methods-for","slug":"randomized-geometric-algebra-methods-for","title":"Randomized Geometric Algebra Methods for Convex Neural Networks","date":"2024-06-04","arxiv_id":"2406.02806","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["pilancilab/Randomized-Geometric-Algebra-Methods-for-Convex-Neural-Networks"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"sms-spam-detection-and-classification-to","title":"SMS Spam Detection and Classification to Combat Abuse in Telephone Networks Using Natural Language Processing","date":"2024-06-04","arxiv_id":"2406.06578","n_code_links":0,"syntology":null},{"paper":"/paper/synergetic-event-understanding-a","slug":"synergetic-event-understanding-a","title":"Synergetic Event Understanding: A Collaborative Approach to Cross-Document Event Coreference Resolution with Large Language Models","date":"2024-06-04","arxiv_id":"2406.02148","n_code_links":1,"syntology":null},{"paper":null,"slug":"towards-effective-time-aware-language","title":"Towards Effective Time-Aware Language Representation: Exploring Enhanced Temporal Understanding in Language Models","date":"2024-06-04","arxiv_id":"2406.01863","n_code_links":0,"syntology":null},{"paper":null,"slug":"annotation-guidelines-based-knowledge","title":"Annotation Guidelines-Based Knowledge Augmentation: Towards Enhancing Large Language Models for Educational Text Classification","date":"2024-06-03","arxiv_id":"2406.00954","n_code_links":0,"syntology":null},{"paper":null,"slug":"ask-eda-a-design-assistant-empowered-by-llm","title":"Ask-EDA: A Design Assistant Empowered by LLM, Hybrid RAG and Abbreviation De-hallucination","date":"2024-06-03","arxiv_id":"2406.06575","n_code_links":0,"syntology":null},{"paper":null,"slug":"badrag-identifying-vulnerabilities-in","title":"BadRAG: Identifying Vulnerabilities in Retrieval Augmented Generation of Large Language Models","date":"2024-06-03","arxiv_id":"2406.00083","n_code_links":0,"syntology":null},{"paper":"/paper/decomposing-and-interpreting-image","slug":"decomposing-and-interpreting-image","title":"Decomposing and Interpreting Image Representations via Text in ViTs Beyond CLIP","date":"2024-06-03","arxiv_id":"2406.01583","n_code_links":1,"syntology":{"ran":10,"of":13,"n_ran_checked":10,"n_instrument":0,"unverified":3,"pointer_only":13,"phrase":"10 ran (of which 0 constructed an object rather than computing a result; 10 with no instrument failure: 1 honoured, 0 violated, 9 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":{"repos":["sriramb-98/vit-decompose"],"state":"official (archive's flag): 10 ran","n_ran":10,"n_constructed":0,"n_ran_no_instrument_failure":10,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":"/paper/factgenius-combining-zero-shot-prompting-and","slug":"factgenius-combining-zero-shot-prompting-and","title":"FactGenius: Combining Zero-Shot Prompting and Fuzzy Relation Mining to Improve Fact Verification with Knowledge Graphs","date":"2024-06-03","arxiv_id":"2406.01311","n_code_links":1,"syntology":null}],"record_sha256":"32f0186416cee6cbdb3affdca534ac013d11f96e040c667296d696d25fd16ae7","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}