{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/discriminative-fine-tuning/papers/7","list_of":"/method/discriminative-fine-tuning","method":"Discriminative Fine-Tuning","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":7,"pages_in_order":20,"rows_per_page":100,"rows":[601,700],"of":1990,"counts":{"archive_papers_tagged":1990,"with_a_code_link":794,"where_syntology_ran_a_sample":271,"not_listed_spam_title":0,"listed":1990,"listed_where_code_ran":271,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":223,"every_run_a_failure_of_syntologys_instrument":48,"listed_with_a_run_with_no_instrument_failure":223,"listed_every_run_a_failure_of_syntologys_instrument":48,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/discriminative-fine-tuning","prev":"/method/discriminative-fine-tuning/papers/6","next":"/method/discriminative-fine-tuning/papers/8","papers":[{"paper":null,"slug":"bert-neural-information-retrieval-boolean","title":"SetBERT: Enhancing Retrieval Performance for Boolean Logic and Set Operation Queries","date":"2024-06-25","arxiv_id":"2406.17282","n_code_links":0,"syntology":null},{"paper":"/paper/interpreting-attention-layer-outputs-with","slug":"interpreting-attention-layer-outputs-with","title":"Interpreting Attention Layer Outputs with Sparse Autoencoders","date":"2024-06-25","arxiv_id":"2406.17759","n_code_links":1,"syntology":{"ran":0,"of":2,"n_ran_checked":0,"n_instrument":0,"unverified":2,"pointer_only":2,"phrase":"0 ran · 2 unverified","official":{"repos":["ckkissane/attention-output-saes"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":2,"ran_from_kinds":[]}}},{"paper":null,"slug":"what-do-the-circuits-mean-a-knowledge-edit","title":"Understanding Language Model Circuits through Knowledge Editing","date":"2024-06-25","arxiv_id":"2406.17241","n_code_links":0,"syntology":null},{"paper":"/paper/dreambench-a-human-aligned-benchmark-for","slug":"dreambench-a-human-aligned-benchmark-for","title":"DreamBench++: A Human-Aligned Benchmark for Personalized Image Generation","date":"2024-06-24","arxiv_id":"2406.16855","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["yuangpeng/dreambench_plus"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/finding-transformer-circuits-with-edge","slug":"finding-transformer-circuits-with-edge","title":"Finding Transformer Circuits with Edge Pruning","date":"2024-06-24","arxiv_id":"2406.16778","n_code_links":1,"syntology":{"ran":8,"of":8,"n_ran_checked":3,"n_instrument":5,"unverified":0,"pointer_only":0,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 3 honoured, 0 violated, 0 with no contract checked; 5 where Syntology's instrument failed) · 0 unverified","official":{"repos":["princeton-nlp/edge-pruning"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"modeling-a-novel-dataset-for-testing","title":"modeLing: A Novel Dataset for Testing Linguistic Reasoning in Language Models","date":"2024-06-24","arxiv_id":"2406.17038","n_code_links":0,"syntology":null},{"paper":"/paper/ss-bench-a-benchmark-for-social-story","slug":"ss-bench-a-benchmark-for-social-story","title":"SS-GEN: A Social Story Generation Framework with Large Language Models","date":"2024-06-22","arxiv_id":"2406.15695","n_code_links":2,"syntology":null},{"paper":null,"slug":"a-gpt-based-code-review-system-for","title":"A GPT-based Code Review System for Programming Language Learning","date":"2024-06-21","arxiv_id":"2407.04722","n_code_links":0,"syntology":null},{"paper":null,"slug":"anime-popularity-prediction-before-huge","title":"Anime Popularity Prediction Before Huge Investments: a Multimodal Approach Using Deep Learning","date":"2024-06-21","arxiv_id":"2406.16961","n_code_links":0,"syntology":null},{"paper":null,"slug":"chatgpt-as-research-scientist-probing-gpt-s","title":"ChatGPT as Research Scientist: Probing GPT's Capabilities as a Research Librarian, Research Ethicist, Data Generator and Data Predictor","date":"2024-06-20","arxiv_id":"2406.14765","n_code_links":0,"syntology":null},{"paper":"/paper/how-to-compute-the-probability-of-a-word","slug":"how-to-compute-the-probability-of-a-word","title":"How to Compute the Probability of a Word","date":"2024-06-20","arxiv_id":"2406.14561","n_code_links":2,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["tpimentelms/probability-of-a-word"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"identifying-user-goals-from-ui-trajectories","title":"Identifying User Goals from UI Trajectories","date":"2024-06-20","arxiv_id":"2406.14314","n_code_links":0,"syntology":null},{"paper":"/paper/persuasiveness-of-generated-free-text","slug":"persuasiveness-of-generated-free-text","title":"Persuasiveness of Generated Free-Text Rationales in Subjective Decisions: A Case Study on Pairwise Argument Ranking","date":"2024-06-20","arxiv_id":"2406.13905","n_code_links":1,"syntology":null},{"paper":"/paper/on-ai-inspired-ui-design","slug":"on-ai-inspired-ui-design","title":"On AI-Inspired UI-Design","date":"2024-06-19","arxiv_id":"2406.13631","n_code_links":1,"syntology":null},{"paper":null,"slug":"open-generative-large-language-models-for","title":"Open Generative Large Language Models for Galician","date":"2024-06-19","arxiv_id":"2406.13893","n_code_links":0,"syntology":null},{"paper":"/paper/ipeval-a-bilingual-intellectual-property","slug":"ipeval-a-bilingual-intellectual-property","title":"IPEval: A Bilingual Intellectual Property Agency Consultation Evaluation Benchmark for Large Language Models","date":"2024-06-18","arxiv_id":"2406.12386","n_code_links":1,"syntology":null},{"paper":null,"slug":"urbanllm-autonomous-urban-activity-planning","title":"UrbanLLM: Autonomous Urban Activity Planning and Management with Large Language Models","date":"2024-06-18","arxiv_id":"2406.12360","n_code_links":0,"syntology":null},{"paper":null,"slug":"what-makes-two-models-think-alike","title":"What Makes Two Language Models Think Alike?","date":"2024-06-18","arxiv_id":"2406.12620","n_code_links":0,"syntology":null},{"paper":null,"slug":"breaking-boundaries-investigating-the-effects","title":"Breaking Boundaries: Investigating the Effects of Model Editing on Cross-linguistic Performance","date":"2024-06-17","arxiv_id":"2406.11139","n_code_links":0,"syntology":null},{"paper":null,"slug":"gpt-powered-elicitation-interview-script","title":"GPT-Powered Elicitation Interview Script Generator for Requirements Engineering Training","date":"2024-06-17","arxiv_id":"2406.11439","n_code_links":0,"syntology":null},{"paper":null,"slug":"improving-multi-agent-debate-with-sparse","title":"Improving Multi-Agent Debate with Sparse Communication Topology","date":"2024-06-17","arxiv_id":"2406.11776","n_code_links":0,"syntology":null},{"paper":"/paper/investigating-annotator-bias-in-large","slug":"investigating-annotator-bias-in-large","title":"Investigating Annotator Bias in Large Language Models for Hate Speech Detection","date":"2024-06-17","arxiv_id":"2406.11109","n_code_links":3,"syntology":null},{"paper":null,"slug":"promises-outlooks-and-challenges-of-diffusion","title":"Promises, Outlooks and Challenges of Diffusion Language Modeling","date":"2024-06-17","arxiv_id":"2406.11473","n_code_links":0,"syntology":null},{"paper":"/paper/scaling-the-codebook-size-of-vqgan-to-100000","slug":"scaling-the-codebook-size-of-vqgan-to-100000","title":"Scaling the Codebook Size of VQGAN to 100,000 with a Utilization Rate of 99%","date":"2024-06-17","arxiv_id":"2406.11837","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":1,"n_instrument":1,"unverified":0,"pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["zh460045050/vqgan-lc"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official","unlocated"]}}},{"paper":null,"slug":"self-and-cross-model-distillation-for-llms","title":"Self and Cross-Model Distillation for LLMs: Effective Methods for Refusal Pattern Alignment","date":"2024-06-17","arxiv_id":"2406.11285","n_code_links":0,"syntology":null},{"paper":null,"slug":"enhancing-supermarket-robot-interaction-a","title":"Enhancing Supermarket Robot Interaction: A Multi-Level LLM Conversational Interface for Handling Diverse Customer Intents","date":"2024-06-16","arxiv_id":"2406.11047","n_code_links":0,"syntology":null},{"paper":null,"slug":"large-language-models-for-automatic-milestone","title":"Large Language Models for Automatic Milestone Detection in Group Discussions","date":"2024-06-16","arxiv_id":"2406.10842","n_code_links":0,"syntology":null},{"paper":"/paper/sharelora-parameter-efficient-and-robust","slug":"sharelora-parameter-efficient-and-robust","title":"ShareLoRA: Parameter Efficient and Robust Large Language Model Fine-tuning via Shared Low-Rank Adaptation","date":"2024-06-16","arxiv_id":"2406.10785","n_code_links":1,"syntology":{"ran":5,"of":6,"n_ran_checked":0,"n_instrument":5,"unverified":1,"pointer_only":6,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 5 where Syntology's instrument failed) · 1 unverified","official":{"repos":["Rain9876/ShareLoRA"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/vid-gpt-introducing-gpt-style-autoregressive","slug":"vid-gpt-introducing-gpt-style-autoregressive","title":"ViD-GPT: Introducing GPT-style Autoregressive Generation in Video Diffusion Models","date":"2024-06-16","arxiv_id":"2406.10981","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-comprehensive-survey-of-foundation-models","title":"A Comprehensive Survey of Foundation Models in Medicine","date":"2024-06-15","arxiv_id":"2406.10729","n_code_links":0,"syntology":null},{"paper":"/paper/beyond-raw-videos-understanding-edited-videos","slug":"beyond-raw-videos-understanding-edited-videos","title":"Beyond Raw Videos: Understanding Edited Videos with Large Multimodal Model","date":"2024-06-15","arxiv_id":"2406.10484","n_code_links":1,"syntology":null},{"paper":"/paper/mint-a-multi-modal-image-and-narrative-text","slug":"mint-a-multi-modal-image-and-narrative-text","title":"MINT: a Multi-modal Image and Narrative Text Dubbing Dataset for Foley Audio Content Planning and Generation","date":"2024-06-15","arxiv_id":"2406.10591","n_code_links":1,"syntology":null},{"paper":null,"slug":"exploring-the-correlation-between-human-and","title":"Exploring the Correlation between Human and Machine Evaluation of Simultaneous Speech Translation","date":"2024-06-14","arxiv_id":"2406.10091","n_code_links":0,"syntology":null},{"paper":"/paper/towards-efficient-pareto-set-approximation","slug":"towards-efficient-pareto-set-approximation","title":"Towards Efficient Pareto Set Approximation via Mixture of Experts Based Model Fusion","date":"2024-06-14","arxiv_id":"2406.09770","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-more-practical-approach-to-machine","title":"A More Practical Approach to Machine Unlearning","date":"2024-06-13","arxiv_id":"2406.09391","n_code_links":0,"syntology":null},{"paper":null,"slug":"optimizing-large-model-training-through","title":"Optimizing Large Model Training through Overlapped Activation Recomputation","date":"2024-06-13","arxiv_id":"2406.08756","n_code_links":0,"syntology":null},{"paper":null,"slug":"talking-heads-understanding-inter-layer","title":"Talking Heads: Understanding Inter-layer Communication in Transformer Language Models","date":"2024-06-13","arxiv_id":"2406.09519","n_code_links":0,"syntology":null},{"paper":null,"slug":"faithfill-faithful-inpainting-for-object","title":"FaithFill: Faithful Inpainting for Object Completion Using a Single Reference Image","date":"2024-06-12","arxiv_id":"2406.07865","n_code_links":0,"syntology":null},{"paper":null,"slug":"how-well-it-works-benchmarking-performance-of","title":"How well it works: Benchmarking performance of GPT models on medical natural language processing tasks","date":"2024-06-12","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/tailoring-generative-ai-chatbots-for","slug":"tailoring-generative-ai-chatbots-for","title":"Tailoring Generative AI Chatbots for Multiethnic Communities in Disaster Preparedness Communication: Extending the CASA Paradigm","date":"2024-06-12","arxiv_id":"2406.08411","n_code_links":1,"syntology":null},{"paper":"/paper/unused-information-in-token-probability","slug":"unused-information-in-token-probability","title":"Unused information in token probability distribution of generative LLM: improving LLM reading comprehension through calculation of expected values","date":"2024-06-11","arxiv_id":"2406.10267","n_code_links":1,"syntology":null},{"paper":"/paper/compute-better-spent-replacing-dense-layers","slug":"compute-better-spent-replacing-dense-layers","title":"Compute Better Spent: Replacing Dense Layers with Structured Matrices","date":"2024-06-10","arxiv_id":"2406.06248","n_code_links":1,"syntology":{"ran":4,"of":5,"n_ran_checked":4,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["shikaiqiu/compute-better-spent"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"llm-dcache-improving-tool-augmented-llms-with","title":"LLM-dCache: Improving Tool-Augmented LLMs with GPT-Driven Localized Data Caching","date":"2024-06-10","arxiv_id":"2406.06799","n_code_links":0,"syntology":null},{"paper":null,"slug":"securenet-a-comparative-study-of-deberta-and","title":"SecureNet: A Comparative Study of DeBERTa and Large Language Models for Phishing Detection","date":"2024-06-10","arxiv_id":"2406.06663","n_code_links":0,"syntology":null},{"paper":"/paper/domainrag-a-chinese-benchmark-for-evaluating","slug":"domainrag-a-chinese-benchmark-for-evaluating","title":"DomainRAG: A Chinese Benchmark for Evaluating Domain-specific Retrieval-Augmented Generation","date":"2024-06-09","arxiv_id":"2406.05654","n_code_links":2,"syntology":{"ran":12,"of":12,"n_ran_checked":11,"n_instrument":1,"unverified":0,"pointer_only":12,"phrase":"12 ran (of which 0 constructed an object rather than computing a result; 11 with no instrument failure: 0 honoured, 0 violated, 11 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["ShootingWong/DomainRAG"],"state":"official (archive's flag): 12 ran","n_ran":12,"n_constructed":0,"n_ran_no_instrument_failure":11,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"hidden-holes-topological-aspects-of-language","title":"Hidden Holes: topological aspects of language models","date":"2024-06-09","arxiv_id":"2406.05798","n_code_links":0,"syntology":null},{"paper":null,"slug":"medreqal-examining-medical-knowledge-recall","title":"MedREQAL: Examining Medical Knowledge Recall of Large Language Models via Question Answering","date":"2024-06-09","arxiv_id":"2406.05845","n_code_links":0,"syntology":null},{"paper":null,"slug":"text2vp-generative-ai-for-visual-programming","title":"Text2VP: Generative AI for Visual Programming and Parametric Modeling","date":"2024-06-09","arxiv_id":"2407.07732","n_code_links":0,"syntology":null},{"paper":"/paper/which-backbone-to-use-a-resource-efficient","slug":"which-backbone-to-use-a-resource-efficient","title":"Which Backbone to Use: A Resource-efficient Domain Specific Comparison for Computer Vision","date":"2024-06-09","arxiv_id":"2406.05612","n_code_links":1,"syntology":null},{"paper":null,"slug":"critical-phase-transition-in-a-large-language","title":"Critical Phase Transition in Large Language Models","date":"2024-06-08","arxiv_id":"2406.05335","n_code_links":0,"syntology":null},{"paper":null,"slug":"matablegpt-gpt-based-table-data-extractor","title":"MaTableGPT: GPT-based Table Data Extractor from Materials Science Literature","date":"2024-06-08","arxiv_id":"2406.05431","n_code_links":0,"syntology":null},{"paper":"/paper/berts-are-generative-in-context-learners","slug":"berts-are-generative-in-context-learners","title":"BERTs are Generative In-Context Learners","date":"2024-06-07","arxiv_id":"2406.04823","n_code_links":1,"syntology":{"ran":13,"of":26,"n_ran_checked":12,"n_instrument":1,"unverified":13,"pointer_only":0,"phrase":"13 ran (of which 0 constructed an object rather than computing a result; 12 with no instrument failure: 0 honoured, 1 violated, 11 with no contract checked; 1 where Syntology's instrument failed) · 13 unverified","official":{"repos":["ltgoslo/bert-in-context"],"state":"official (archive's flag): 13 ran","n_ran":13,"n_constructed":0,"n_ran_no_instrument_failure":12,"n_unverified":13,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"large-generative-graph-models","title":"Large Generative Graph Models","date":"2024-06-07","arxiv_id":"2406.05109","n_code_links":0,"syntology":null},{"paper":"/paper/on-subjective-uncertainty-quantification-and","slug":"on-subjective-uncertainty-quantification-and","title":"On Subjective Uncertainty Quantification and Calibration in Natural Language Generation","date":"2024-06-07","arxiv_id":"2406.05213","n_code_links":1,"syntology":{"ran":12,"of":14,"n_ran_checked":12,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"12 ran (of which 0 constructed an object rather than computing a result; 12 with no instrument failure: 0 honoured, 0 violated, 12 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["meta-inf/suq-nlg"],"state":"official (archive's flag): 12 ran","n_ran":12,"n_constructed":0,"n_ran_no_instrument_failure":12,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"vtrans-accelerating-transformer-compression","title":"VTrans: Accelerating Transformer Compression with Variational Information Bottleneck based Pruning","date":"2024-06-07","arxiv_id":"2406.05276","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-survey-on-medical-large-language-models","title":"A Survey on Medical Large Language Models: Technology, Application, Trustworthiness, and Future Directions","date":"2024-06-06","arxiv_id":"2406.03712","n_code_links":0,"syntology":null},{"paper":"/paper/do-language-models-understand-morality","slug":"do-language-models-understand-morality","title":"Do Language Models Understand Morality? Towards a Robust Detection of Moral Content","date":"2024-06-06","arxiv_id":"2406.04143","n_code_links":1,"syntology":null},{"paper":"/paper/simplified-and-generalized-masked-diffusion","slug":"simplified-and-generalized-masked-diffusion","title":"Simplified and Generalized Masked Diffusion for Discrete Data","date":"2024-06-06","arxiv_id":"2406.04329","n_code_links":1,"syntology":{"ran":6,"of":21,"n_ran_checked":2,"n_instrument":4,"unverified":15,"pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 1 honoured, 0 violated, 1 with no contract checked; 4 where Syntology's instrument failed) · 15 unverified","official":{"repos":["google-deepmind/md4"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":15,"ran_from_kinds":["official"]}}},{"paper":"/paper/your-absorbing-discrete-diffusion-secretly","slug":"your-absorbing-discrete-diffusion-secretly","title":"Your Absorbing Discrete Diffusion Secretly Models the Conditional Distributions of Clean Data","date":"2024-06-06","arxiv_id":"2406.03736","n_code_links":2,"syntology":{"ran":13,"of":13,"n_ran_checked":12,"n_instrument":1,"unverified":0,"pointer_only":13,"phrase":"13 ran (of which 0 constructed an object rather than computing a result; 12 with no instrument failure: 1 honoured, 0 violated, 11 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["ml-gsai/radd"],"state":"official (archive's flag): 13 ran","n_ran":13,"n_constructed":0,"n_ran_no_instrument_failure":12,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"exact-conversion-of-in-context-learning-to","title":"Exact Conversion of In-Context Learning to Model Weights in Linearized-Attention Transformers","date":"2024-06-05","arxiv_id":"2406.02847","n_code_links":0,"syntology":null},{"paper":"/paper/missci-reconstructing-fallacies-in","slug":"missci-reconstructing-fallacies-in","title":"Missci: Reconstructing Fallacies in Misrepresented Science","date":"2024-06-05","arxiv_id":"2406.03181","n_code_links":2,"syntology":null},{"paper":"/paper/rico-reddit-ideological-communities","slug":"rico-reddit-ideological-communities","title":"RICo: Reddit ideological communities","date":"2024-06-05","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/too-big-to-fail-larger-language-models-are","slug":"too-big-to-fail-larger-language-models-are","title":"Too Big to Fail: Larger Language Models are Disproportionately Resilient to Induction of Dementia-Related Linguistic Anomalies","date":"2024-06-05","arxiv_id":"2406.02830","n_code_links":1,"syntology":null},{"paper":"/paper/checkembed-effective-verification-of-llm","slug":"checkembed-effective-verification-of-llm","title":"CheckEmbed: Effective Verification of LLM Solutions to Open-Ended Tasks","date":"2024-06-04","arxiv_id":"2406.02524","n_code_links":1,"syntology":null},{"paper":null,"slug":"occamllm-fast-and-exact-language-model","title":"OccamLLM: Fast and Exact Language Model Arithmetic in a Single Step","date":"2024-06-04","arxiv_id":"2406.06576","n_code_links":0,"syntology":null},{"paper":"/paper/randomized-geometric-algebra-methods-for","slug":"randomized-geometric-algebra-methods-for","title":"Randomized Geometric Algebra Methods for Convex Neural Networks","date":"2024-06-04","arxiv_id":"2406.02806","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["pilancilab/Randomized-Geometric-Algebra-Methods-for-Convex-Neural-Networks"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"annotation-guidelines-based-knowledge","title":"Annotation Guidelines-Based Knowledge Augmentation: Towards Enhancing Large Language Models for Educational Text Classification","date":"2024-06-03","arxiv_id":"2406.00954","n_code_links":0,"syntology":null},{"paper":null,"slug":"in-context-learning-of-physical-properties","title":"In-Context Learning of Physical Properties: Few-Shot Adaptation to Out-of-Distribution Molecular Graphs","date":"2024-06-03","arxiv_id":"2406.01808","n_code_links":0,"syntology":null},{"paper":"/paper/spatialrgpt-grounded-spatial-reasoning-in","slug":"spatialrgpt-grounded-spatial-reasoning-in","title":"SpatialRGPT: Grounded Spatial Reasoning in Vision Language Models","date":"2024-06-03","arxiv_id":"2406.01584","n_code_links":1,"syntology":null},{"paper":null,"slug":"focus-forging-originality-through-contrastive","title":"FOCUS: Forging Originality through Contrastive Use in Self-Plagiarism for Language Models","date":"2024-06-02","arxiv_id":"2406.00839","n_code_links":0,"syntology":null},{"paper":null,"slug":"pretrained-hybrids-with-mad-skills","title":"Pretrained Hybrids with MAD Skills","date":"2024-06-02","arxiv_id":"2406.00894","n_code_links":0,"syntology":null},{"paper":null,"slug":"2406-07572","title":"Domain-specific ReAct for physics-integrated iterative modeling: A case study of LLM agents for gas path analysis of gas turbines","date":"2024-06-01","arxiv_id":"2406.07572","n_code_links":0,"syntology":null},{"paper":"/paper/generative-ai-voting-fair-collective-choice","slug":"generative-ai-voting-fair-collective-choice","title":"Generative AI Voting: Fair Collective Choice is Resilient to LLM Biases and Inconsistencies","date":"2024-05-31","arxiv_id":"2406.11871","n_code_links":1,"syntology":null},{"paper":"/paper/hard-cases-detection-in-motion-prediction-by","slug":"hard-cases-detection-in-motion-prediction-by","title":"Hard Cases Detection in Motion Prediction by Vision-Language Foundation Models","date":"2024-05-31","arxiv_id":"2405.20991","n_code_links":1,"syntology":null},{"paper":null,"slug":"lolameme-logic-language-memory-mechanistic","title":"LOLAMEME: Logic, Language, Memory, Mechanistic Framework","date":"2024-05-31","arxiv_id":"2406.02592","n_code_links":0,"syntology":null},{"paper":null,"slug":"knowledge-graph-tuning-real-time-large","title":"Knowledge Graph Tuning: Real-time Large Language Model Personalization based on Human Feedback","date":"2024-05-30","arxiv_id":"2405.19686","n_code_links":0,"syntology":null},{"paper":"/paper/llamea-a-large-language-model-evolutionary","slug":"llamea-a-large-language-model-evolutionary","title":"LLaMEA: A Large Language Model Evolutionary Algorithm for Automatically Generating Metaheuristics","date":"2024-05-30","arxiv_id":"2405.20132","n_code_links":2,"syntology":{"ran":6,"of":6,"n_ran_checked":6,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["nikivanstein/LLaMEA"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["listed","official"]}}},{"paper":null,"slug":"significance-of-chain-of-thought-in-gender","title":"Significance of Chain of Thought in Gender Bias Mitigation for English-Dravidian Machine Translation","date":"2024-05-30","arxiv_id":"2405.19701","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-multi-source-retrieval-question-answering","title":"A Multi-Source Retrieval Question Answering Framework Based on RAG","date":"2024-05-29","arxiv_id":"2405.19207","n_code_links":0,"syntology":null},{"paper":null,"slug":"can-gpt-redefine-medical-understanding","title":"Can GPT Redefine Medical Understanding? Evaluating GPT on Biomedical Machine Reading Comprehension","date":"2024-05-29","arxiv_id":"2405.18682","n_code_links":0,"syntology":null},{"paper":null,"slug":"efficient-model-agnostic-alignment-via","title":"Efficient Model-agnostic Alignment via Bayesian Persuasion","date":"2024-05-29","arxiv_id":"2405.18718","n_code_links":0,"syntology":null},{"paper":null,"slug":"lmo-dp-optimizing-the-randomization-mechanism","title":"LMO-DP: Optimizing the Randomization Mechanism for Differentially Private Fine-Tuning (Large) Language Models","date":"2024-05-29","arxiv_id":"2405.18776","n_code_links":0,"syntology":null},{"paper":"/paper/map-neo-highly-capable-and-transparent","slug":"map-neo-highly-capable-and-transparent","title":"MAP-Neo: Highly Capable and Transparent Bilingual Large Language Model Series","date":"2024-05-29","arxiv_id":"2405.19327","n_code_links":1,"syntology":{"ran":12,"of":14,"n_ran_checked":12,"n_instrument":0,"unverified":2,"pointer_only":14,"phrase":"12 ran (of which 0 constructed an object rather than computing a result; 12 with no instrument failure: 0 honoured, 0 violated, 12 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["multimodal-art-projection/map-neo"],"state":"official (archive's flag): 12 ran","n_ran":12,"n_constructed":0,"n_ran_no_instrument_failure":12,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"are-ppo-ed-language-models-hackable","title":"Are PPO-ed Language Models Hackable?","date":"2024-05-28","arxiv_id":"2406.02577","n_code_links":0,"syntology":null},{"paper":"/paper/llms-and-memorization-on-quality-and","slug":"llms-and-memorization-on-quality-and","title":"LLMs and Memorization: On Quality and Specificity of Copyright Compliance","date":"2024-05-28","arxiv_id":"2405.18492","n_code_links":1,"syntology":null},{"paper":null,"slug":"understanding-intrinsic-socioeconomic-biases","title":"Understanding Intrinsic Socioeconomic Biases in Large Language Models","date":"2024-05-28","arxiv_id":"2405.18662","n_code_links":0,"syntology":null},{"paper":"/paper/inversionview-a-general-purpose-method-for","slug":"inversionview-a-general-purpose-method-for","title":"InversionView: A General-Purpose Method for Reading Information from Neural Activations","date":"2024-05-27","arxiv_id":"2405.17653","n_code_links":1,"syntology":{"ran":1,"of":6,"n_ran_checked":1,"n_instrument":0,"unverified":5,"pointer_only":6,"phrase":"1 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 5 unverified; the one sample that ran constructed an object rather than computing a result","official":{"repos":["huangxt39/inversionview"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":5,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"performance-evaluation-of-reddit-comments","title":"Performance evaluation of Reddit Comments using Machine Learning and Natural Language Processing methods in Sentiment Analysis","date":"2024-05-27","arxiv_id":"2405.16810","n_code_links":0,"syntology":null},{"paper":null,"slug":"the-scaling-law-in-stellar-light-curves","title":"The Scaling Law in Stellar Light Curves","date":"2024-05-27","arxiv_id":"2405.17156","n_code_links":0,"syntology":null},{"paper":"/paper/accelerating-transformers-with-spectrum-1","slug":"accelerating-transformers-with-spectrum-1","title":"Accelerating Transformers with Spectrum-Preserving Token Merging","date":"2024-05-25","arxiv_id":"2405.16148","n_code_links":1,"syntology":{"ran":6,"of":7,"n_ran_checked":2,"n_instrument":4,"unverified":1,"pointer_only":7,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 2 honoured, 0 violated, 0 with no contract checked; 4 where Syntology's instrument failed) · 1 unverified","official":{"repos":["hchautran/PiToMe"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"incremental-comprehension-of-garden-path","title":"Incremental Comprehension of Garden-Path Sentences by Large Language Models: Semantic Interpretation, Syntactic Re-Analysis, and Attention","date":"2024-05-25","arxiv_id":"2405.16042","n_code_links":0,"syntology":null},{"paper":"/paper/learning-the-language-of-protein-structure","slug":"learning-the-language-of-protein-structure","title":"Learning the Language of Protein Structure","date":"2024-05-24","arxiv_id":"2405.15840","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["instadeepai/protein-structure-tokenizer"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"the-impact-and-opportunities-of-generative-ai","title":"The Impact and Opportunities of Generative AI in Fact-Checking","date":"2024-05-24","arxiv_id":"2405.15985","n_code_links":0,"syntology":null},{"paper":null,"slug":"towards-understanding-how-transformer-perform","title":"The Buffer Mechanism for Multi-Step Information Reasoning in Language Models","date":"2024-05-24","arxiv_id":"2405.15302","n_code_links":0,"syntology":null},{"paper":"/paper/from-explicit-cot-to-implicit-cot-learning-to","slug":"from-explicit-cot-to-implicit-cot-learning-to","title":"From Explicit CoT to Implicit CoT: Learning to Internalize CoT Step by Step","date":"2024-05-23","arxiv_id":"2405.14838","n_code_links":1,"syntology":{"ran":4,"of":4,"n_ran_checked":4,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["da03/internalize_cot_step_by_step"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/not-all-language-model-features-are-linear","slug":"not-all-language-model-features-are-linear","title":"Not All Language Model Features Are Linear","date":"2024-05-23","arxiv_id":"2405.14860","n_code_links":1,"syntology":{"ran":6,"of":10,"n_ran_checked":6,"n_instrument":0,"unverified":4,"pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","official":{"repos":["joshengels/multidimensionalfeatures"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":"/paper/wise-rethinking-the-knowledge-memory-for","slug":"wise-rethinking-the-knowledge-memory-for","title":"WISE: Rethinking the Knowledge Memory for Lifelong Model Editing of Large Language Models","date":"2024-05-23","arxiv_id":"2405.14768","n_code_links":1,"syntology":{"ran":13,"of":16,"n_ran_checked":4,"n_instrument":9,"unverified":3,"pointer_only":0,"phrase":"13 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 1 honoured, 0 violated, 3 with no contract checked; 9 where Syntology's instrument failed) · 3 unverified","official":{"repos":["zjunlp/easyedit"],"state":"official (archive's flag): 13 ran","n_ran":13,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"automatically-identifying-local-and-global","title":"Automatically Identifying Local and Global Circuits with Linear Computation Graphs","date":"2024-05-22","arxiv_id":"2405.13868","n_code_links":0,"syntology":null},{"paper":null,"slug":"ku-dmis-at-ehrsql-2024-generating-sql-query","title":"KU-DMIS at EHRSQL 2024:Generating SQL query via question templatization in EHR","date":"2024-05-22","arxiv_id":"2406.00014","n_code_links":0,"syntology":null},{"paper":null,"slug":"how-reliable-ai-chatbots-are-for-disease","title":"How Reliable AI Chatbots are for Disease Prediction from Patient Complaints?","date":"2024-05-21","arxiv_id":"2405.13219","n_code_links":0,"syntology":null}],"record_sha256":"bc72e98782da46e5f12cfbd6e594c9c0ba9bb15d92f85cd950836e1671ecf775","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}