{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/gpt/papers/5","list_of":"/method/gpt","method":"GPT","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":5,"pages_in_order":13,"rows_per_page":100,"rows":[401,500],"of":1212,"counts":{"archive_papers_tagged":1212,"with_a_code_link":453,"where_syntology_ran_a_sample":152,"not_listed_spam_title":0,"listed":1212,"listed_where_code_ran":152,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":130,"every_run_a_failure_of_syntologys_instrument":22,"listed_with_a_run_with_no_instrument_failure":130,"listed_every_run_a_failure_of_syntologys_instrument":22,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/gpt","prev":"/method/gpt/papers/4","next":"/method/gpt/papers/6","papers":[{"paper":"/paper/trust-no-bot-discovering-personal-disclosures","slug":"trust-no-bot-discovering-personal-disclosures","title":"Trust No Bot: Discovering Personal Disclosures in Human-LLM Conversations in the Wild","date":"2024-07-16","arxiv_id":"2407.11438","n_code_links":1,"syntology":null},{"paper":null,"slug":"empowering-llms-for-verilog-generation","title":"CodeV: Empowering LLMs with HDL Generation through Multi-Level Summarization","date":"2024-07-15","arxiv_id":"2407.10424","n_code_links":0,"syntology":null},{"paper":null,"slug":"making-new-connections-llms-as-puzzle","title":"Making New Connections: LLMs as Puzzle Generators for The New York Times' Connections Word Game","date":"2024-07-15","arxiv_id":"2407.11240","n_code_links":0,"syntology":null},{"paper":"/paper/metallm-a-high-performant-and-cost-efficient","slug":"metallm-a-high-performant-and-cost-efficient","title":"MetaLLM: A High-performant and Cost-efficient Dynamic Framework for Wrapping LLMs","date":"2024-07-15","arxiv_id":"2407.10834","n_code_links":1,"syntology":{"ran":6,"of":7,"n_ran_checked":6,"n_instrument":0,"unverified":1,"pointer_only":7,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["mail-research/metallm-wrapper"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"curriculum-learning-for-small-code-language","title":"Curriculum Learning for Small Code Language Models","date":"2024-07-14","arxiv_id":"2407.10194","n_code_links":0,"syntology":null},{"paper":null,"slug":"document-level-clinical-entity-and-relation","title":"Document-level Clinical Entity and Relation Extraction via Knowledge Base-Guided Generation","date":"2024-07-13","arxiv_id":"2407.10021","n_code_links":0,"syntology":null},{"paper":"/paper/show-don-t-tell-evaluating-large-language","slug":"show-don-t-tell-evaluating-large-language","title":"Show, Don't Tell: Evaluating Large Language Models Beyond Textual Understanding with ChildPlay","date":"2024-07-12","arxiv_id":"2407.11068","n_code_links":1,"syntology":null},{"paper":null,"slug":"toward-automatic-group-membership-annotation","title":"Toward Automatic Group Membership Annotation for Group Fairness Evaluation","date":"2024-07-12","arxiv_id":"2407.08926","n_code_links":0,"syntology":null},{"paper":"/paper/mavis-mathematical-visual-instruction-tuning","slug":"mavis-mathematical-visual-instruction-tuning","title":"MAVIS: Mathematical Visual Instruction Tuning with an Automatic Data Engine","date":"2024-07-11","arxiv_id":"2407.08739","n_code_links":3,"syntology":null},{"paper":null,"slug":"on-the-in-security-of-llm-app-stores","title":"On the (In)Security of LLM App Stores","date":"2024-07-11","arxiv_id":"2407.08422","n_code_links":0,"syntology":null},{"paper":null,"slug":"real-time-anomaly-detection-and-reactive","title":"Real-Time Anomaly Detection and Reactive Planning with Large Language Models","date":"2024-07-11","arxiv_id":"2407.08735","n_code_links":0,"syntology":null},{"paper":"/paper/kpopmt-translation-dataset-with-terminology","slug":"kpopmt-translation-dataset-with-terminology","title":"KpopMT: Translation Dataset with Terminology for Kpop Fandom","date":"2024-07-10","arxiv_id":"2407.07413","n_code_links":1,"syntology":null},{"paper":null,"slug":"measuring-sustainability-intention-of-esg","title":"Measuring Sustainability Intention of ESG Fund Disclosure using Few-Shot Learning","date":"2024-07-09","arxiv_id":"2407.06893","n_code_links":0,"syntology":null},{"paper":null,"slug":"solving-general-natural-language-description","title":"Solving General Natural-Language-Description Optimization Problems with Large Language Models","date":"2024-07-09","arxiv_id":"2407.07924","n_code_links":0,"syntology":null},{"paper":null,"slug":"using-pretrained-large-language-model-with","title":"Using Pretrained Large Language Model with Prompt Engineering to Answer Biomedical Questions","date":"2024-07-09","arxiv_id":"2407.06779","n_code_links":0,"syntology":null},{"paper":null,"slug":"potential-of-multimodal-large-language-models","title":"Potential of Multimodal Large Language Models for Data Mining of Medical Images and Free-text Reports","date":"2024-07-08","arxiv_id":"2407.05758","n_code_links":0,"syntology":null},{"paper":null,"slug":"surprising-gender-biases-in-gpt","title":"Surprising gender biases in GPT","date":"2024-07-08","arxiv_id":"2407.06003","n_code_links":0,"syntology":null},{"paper":null,"slug":"gpt-vs-retro-exploring-the-intersection-of","title":"GPT vs RETRO: Exploring the Intersection of Retrieval and Parameter-Efficient Fine-Tuning","date":"2024-07-05","arxiv_id":"2407.04528","n_code_links":0,"syntology":null},{"paper":null,"slug":"question-analysis-prompting-improves-llm","title":"Question-Analysis Prompting Improves LLM Performance in Reasoning Tasks","date":"2024-07-04","arxiv_id":"2407.03624","n_code_links":0,"syntology":null},{"paper":"/paper/automatic-gradient-descent-with-generalized","slug":"automatic-gradient-descent-with-generalized","title":"Gradient descent with generalized Newton's method","date":"2024-07-03","arxiv_id":"2407.02772","n_code_links":1,"syntology":{"ran":1,"of":2,"n_ran_checked":1,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["shiyunxu/autogen"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"improving-llm-abilities-in-idiomatic","title":"Improving LLM Abilities in Idiomatic Translation","date":"2024-07-03","arxiv_id":"2407.03518","n_code_links":0,"syntology":null},{"paper":null,"slug":"ospc-artificial-vlm-features-for-hateful-meme","title":"OSPC: Artificial VLM Features for Hateful Meme Detection","date":"2024-07-03","arxiv_id":"2407.12836","n_code_links":0,"syntology":null},{"paper":null,"slug":"assessing-the-code-clone-detection-capability","title":"Assessing the Code Clone Detection Capability of Large Language Models","date":"2024-07-02","arxiv_id":"2407.02402","n_code_links":0,"syntology":null},{"paper":"/paper/gptcast-a-weather-language-model-for","slug":"gptcast-a-weather-language-model-for","title":"GPTCast: a weather language model for precipitation nowcasting","date":"2024-07-02","arxiv_id":"2407.02089","n_code_links":1,"syntology":null},{"paper":null,"slug":"predicting-dc-link-capacitor-current-ripple","title":"Predicting DC-Link Capacitor Current Ripple in AC-DC Rectifier Circuits Using Fine-Tuned Large Language Models","date":"2024-07-01","arxiv_id":"2407.01724","n_code_links":0,"syntology":null},{"paper":null,"slug":"granite-function-calling-model-introducing","title":"Granite-Function Calling Model: Introducing Function Calling Abilities via Multi-task Learning of Granular Tasks","date":"2024-06-27","arxiv_id":"2407.00121","n_code_links":0,"syntology":null},{"paper":"/paper/factfinders-at-checkthat-2024-refining-check","slug":"factfinders-at-checkthat-2024-refining-check","title":"FactFinders at CheckThat! 2024: Refining Check-worthy Statement Detection with LLMs through Data Pruning","date":"2024-06-26","arxiv_id":"2406.18297","n_code_links":1,"syntology":null},{"paper":"/paper/mathodyssey-benchmarking-mathematical-problem","slug":"mathodyssey-benchmarking-mathematical-problem","title":"MathOdyssey: Benchmarking Mathematical Problem-Solving Skills in Large Language Models Using Odyssey Math Data","date":"2024-06-26","arxiv_id":"2406.18321","n_code_links":3,"syntology":{"ran":5,"of":5,"n_ran_checked":4,"n_instrument":1,"unverified":0,"pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":null}},{"paper":null,"slug":"bert-neural-information-retrieval-boolean","title":"SetBERT: Enhancing Retrieval Performance for Boolean Logic and Set Operation Queries","date":"2024-06-25","arxiv_id":"2406.17282","n_code_links":0,"syntology":null},{"paper":"/paper/dreambench-a-human-aligned-benchmark-for","slug":"dreambench-a-human-aligned-benchmark-for","title":"DreamBench++: A Human-Aligned Benchmark for Personalized Image Generation","date":"2024-06-24","arxiv_id":"2406.16855","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["yuangpeng/dreambench_plus"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"modeling-a-novel-dataset-for-testing","title":"modeLing: A Novel Dataset for Testing Linguistic Reasoning in Language Models","date":"2024-06-24","arxiv_id":"2406.17038","n_code_links":0,"syntology":null},{"paper":"/paper/ss-bench-a-benchmark-for-social-story","slug":"ss-bench-a-benchmark-for-social-story","title":"SS-GEN: A Social Story Generation Framework with Large Language Models","date":"2024-06-22","arxiv_id":"2406.15695","n_code_links":2,"syntology":null},{"paper":null,"slug":"a-gpt-based-code-review-system-for","title":"A GPT-based Code Review System for Programming Language Learning","date":"2024-06-21","arxiv_id":"2407.04722","n_code_links":0,"syntology":null},{"paper":null,"slug":"chatgpt-as-research-scientist-probing-gpt-s","title":"ChatGPT as Research Scientist: Probing GPT's Capabilities as a Research Librarian, Research Ethicist, Data Generator and Data Predictor","date":"2024-06-20","arxiv_id":"2406.14765","n_code_links":0,"syntology":null},{"paper":"/paper/how-to-compute-the-probability-of-a-word","slug":"how-to-compute-the-probability-of-a-word","title":"How to Compute the Probability of a Word","date":"2024-06-20","arxiv_id":"2406.14561","n_code_links":2,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["tpimentelms/probability-of-a-word"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"identifying-user-goals-from-ui-trajectories","title":"Identifying User Goals from UI Trajectories","date":"2024-06-20","arxiv_id":"2406.14314","n_code_links":0,"syntology":null},{"paper":"/paper/persuasiveness-of-generated-free-text","slug":"persuasiveness-of-generated-free-text","title":"Persuasiveness of Generated Free-Text Rationales in Subjective Decisions: A Case Study on Pairwise Argument Ranking","date":"2024-06-20","arxiv_id":"2406.13905","n_code_links":1,"syntology":null},{"paper":"/paper/on-ai-inspired-ui-design","slug":"on-ai-inspired-ui-design","title":"On AI-Inspired UI-Design","date":"2024-06-19","arxiv_id":"2406.13631","n_code_links":1,"syntology":null},{"paper":null,"slug":"open-generative-large-language-models-for","title":"Open Generative Large Language Models for Galician","date":"2024-06-19","arxiv_id":"2406.13893","n_code_links":0,"syntology":null},{"paper":"/paper/ipeval-a-bilingual-intellectual-property","slug":"ipeval-a-bilingual-intellectual-property","title":"IPEval: A Bilingual Intellectual Property Agency Consultation Evaluation Benchmark for Large Language Models","date":"2024-06-18","arxiv_id":"2406.12386","n_code_links":1,"syntology":null},{"paper":null,"slug":"urbanllm-autonomous-urban-activity-planning","title":"UrbanLLM: Autonomous Urban Activity Planning and Management with Large Language Models","date":"2024-06-18","arxiv_id":"2406.12360","n_code_links":0,"syntology":null},{"paper":null,"slug":"breaking-boundaries-investigating-the-effects","title":"Breaking Boundaries: Investigating the Effects of Model Editing on Cross-linguistic Performance","date":"2024-06-17","arxiv_id":"2406.11139","n_code_links":0,"syntology":null},{"paper":null,"slug":"gpt-powered-elicitation-interview-script","title":"GPT-Powered Elicitation Interview Script Generator for Requirements Engineering Training","date":"2024-06-17","arxiv_id":"2406.11439","n_code_links":0,"syntology":null},{"paper":null,"slug":"improving-multi-agent-debate-with-sparse","title":"Improving Multi-Agent Debate with Sparse Communication Topology","date":"2024-06-17","arxiv_id":"2406.11776","n_code_links":0,"syntology":null},{"paper":"/paper/investigating-annotator-bias-in-large","slug":"investigating-annotator-bias-in-large","title":"Investigating Annotator Bias in Large Language Models for Hate Speech Detection","date":"2024-06-17","arxiv_id":"2406.11109","n_code_links":3,"syntology":null},{"paper":"/paper/scaling-the-codebook-size-of-vqgan-to-100000","slug":"scaling-the-codebook-size-of-vqgan-to-100000","title":"Scaling the Codebook Size of VQGAN to 100,000 with a Utilization Rate of 99%","date":"2024-06-17","arxiv_id":"2406.11837","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":1,"n_instrument":1,"unverified":0,"pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["zh460045050/vqgan-lc"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official","unlocated"]}}},{"paper":null,"slug":"self-and-cross-model-distillation-for-llms","title":"Self and Cross-Model Distillation for LLMs: Effective Methods for Refusal Pattern Alignment","date":"2024-06-17","arxiv_id":"2406.11285","n_code_links":0,"syntology":null},{"paper":null,"slug":"enhancing-supermarket-robot-interaction-a","title":"Enhancing Supermarket Robot Interaction: A Multi-Level LLM Conversational Interface for Handling Diverse Customer Intents","date":"2024-06-16","arxiv_id":"2406.11047","n_code_links":0,"syntology":null},{"paper":null,"slug":"large-language-models-for-automatic-milestone","title":"Large Language Models for Automatic Milestone Detection in Group Discussions","date":"2024-06-16","arxiv_id":"2406.10842","n_code_links":0,"syntology":null},{"paper":"/paper/vid-gpt-introducing-gpt-style-autoregressive","slug":"vid-gpt-introducing-gpt-style-autoregressive","title":"ViD-GPT: Introducing GPT-style Autoregressive Generation in Video Diffusion Models","date":"2024-06-16","arxiv_id":"2406.10981","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-comprehensive-survey-of-foundation-models","title":"A Comprehensive Survey of Foundation Models in Medicine","date":"2024-06-15","arxiv_id":"2406.10729","n_code_links":0,"syntology":null},{"paper":"/paper/beyond-raw-videos-understanding-edited-videos","slug":"beyond-raw-videos-understanding-edited-videos","title":"Beyond Raw Videos: Understanding Edited Videos with Large Multimodal Model","date":"2024-06-15","arxiv_id":"2406.10484","n_code_links":1,"syntology":null},{"paper":null,"slug":"exploring-the-correlation-between-human-and","title":"Exploring the Correlation between Human and Machine Evaluation of Simultaneous Speech Translation","date":"2024-06-14","arxiv_id":"2406.10091","n_code_links":0,"syntology":null},{"paper":null,"slug":"optimizing-large-model-training-through","title":"Optimizing Large Model Training through Overlapped Activation Recomputation","date":"2024-06-13","arxiv_id":"2406.08756","n_code_links":0,"syntology":null},{"paper":null,"slug":"faithfill-faithful-inpainting-for-object","title":"FaithFill: Faithful Inpainting for Object Completion Using a Single Reference Image","date":"2024-06-12","arxiv_id":"2406.07865","n_code_links":0,"syntology":null},{"paper":null,"slug":"how-well-it-works-benchmarking-performance-of","title":"How well it works: Benchmarking performance of GPT models on medical natural language processing tasks","date":"2024-06-12","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/tailoring-generative-ai-chatbots-for","slug":"tailoring-generative-ai-chatbots-for","title":"Tailoring Generative AI Chatbots for Multiethnic Communities in Disaster Preparedness Communication: Extending the CASA Paradigm","date":"2024-06-12","arxiv_id":"2406.08411","n_code_links":1,"syntology":null},{"paper":"/paper/unused-information-in-token-probability","slug":"unused-information-in-token-probability","title":"Unused information in token probability distribution of generative LLM: improving LLM reading comprehension through calculation of expected values","date":"2024-06-11","arxiv_id":"2406.10267","n_code_links":1,"syntology":null},{"paper":null,"slug":"llm-dcache-improving-tool-augmented-llms-with","title":"LLM-dCache: Improving Tool-Augmented LLMs with GPT-Driven Localized Data Caching","date":"2024-06-10","arxiv_id":"2406.06799","n_code_links":0,"syntology":null},{"paper":null,"slug":"securenet-a-comparative-study-of-deberta-and","title":"SecureNet: A Comparative Study of DeBERTa and Large Language Models for Phishing Detection","date":"2024-06-10","arxiv_id":"2406.06663","n_code_links":0,"syntology":null},{"paper":"/paper/domainrag-a-chinese-benchmark-for-evaluating","slug":"domainrag-a-chinese-benchmark-for-evaluating","title":"DomainRAG: A Chinese Benchmark for Evaluating Domain-specific Retrieval-Augmented Generation","date":"2024-06-09","arxiv_id":"2406.05654","n_code_links":2,"syntology":{"ran":12,"of":12,"n_ran_checked":11,"n_instrument":1,"unverified":0,"pointer_only":12,"phrase":"12 ran (of which 0 constructed an object rather than computing a result; 11 with no instrument failure: 0 honoured, 0 violated, 11 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["ShootingWong/DomainRAG"],"state":"official (archive's flag): 12 ran","n_ran":12,"n_constructed":0,"n_ran_no_instrument_failure":11,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"hidden-holes-topological-aspects-of-language","title":"Hidden Holes: topological aspects of language models","date":"2024-06-09","arxiv_id":"2406.05798","n_code_links":0,"syntology":null},{"paper":null,"slug":"medreqal-examining-medical-knowledge-recall","title":"MedREQAL: Examining Medical Knowledge Recall of Large Language Models via Question Answering","date":"2024-06-09","arxiv_id":"2406.05845","n_code_links":0,"syntology":null},{"paper":null,"slug":"text2vp-generative-ai-for-visual-programming","title":"Text2VP: Generative AI for Visual Programming and Parametric Modeling","date":"2024-06-09","arxiv_id":"2407.07732","n_code_links":0,"syntology":null},{"paper":null,"slug":"matablegpt-gpt-based-table-data-extractor","title":"MaTableGPT: GPT-based Table Data Extractor from Materials Science Literature","date":"2024-06-08","arxiv_id":"2406.05431","n_code_links":0,"syntology":null},{"paper":"/paper/berts-are-generative-in-context-learners","slug":"berts-are-generative-in-context-learners","title":"BERTs are Generative In-Context Learners","date":"2024-06-07","arxiv_id":"2406.04823","n_code_links":1,"syntology":{"ran":13,"of":26,"n_ran_checked":12,"n_instrument":1,"unverified":13,"pointer_only":0,"phrase":"13 ran (of which 0 constructed an object rather than computing a result; 12 with no instrument failure: 0 honoured, 1 violated, 11 with no contract checked; 1 where Syntology's instrument failed) · 13 unverified","official":{"repos":["ltgoslo/bert-in-context"],"state":"official (archive's flag): 13 ran","n_ran":13,"n_constructed":0,"n_ran_no_instrument_failure":12,"n_unverified":13,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"large-generative-graph-models","title":"Large Generative Graph Models","date":"2024-06-07","arxiv_id":"2406.05109","n_code_links":0,"syntology":null},{"paper":"/paper/on-subjective-uncertainty-quantification-and","slug":"on-subjective-uncertainty-quantification-and","title":"On Subjective Uncertainty Quantification and Calibration in Natural Language Generation","date":"2024-06-07","arxiv_id":"2406.05213","n_code_links":1,"syntology":{"ran":12,"of":14,"n_ran_checked":12,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"12 ran (of which 0 constructed an object rather than computing a result; 12 with no instrument failure: 0 honoured, 0 violated, 12 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["meta-inf/suq-nlg"],"state":"official (archive's flag): 12 ran","n_ran":12,"n_constructed":0,"n_ran_no_instrument_failure":12,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"a-survey-on-medical-large-language-models","title":"A Survey on Medical Large Language Models: Technology, Application, Trustworthiness, and Future Directions","date":"2024-06-06","arxiv_id":"2406.03712","n_code_links":0,"syntology":null},{"paper":"/paper/do-language-models-understand-morality","slug":"do-language-models-understand-morality","title":"Do Language Models Understand Morality? Towards a Robust Detection of Moral Content","date":"2024-06-06","arxiv_id":"2406.04143","n_code_links":1,"syntology":null},{"paper":"/paper/missci-reconstructing-fallacies-in","slug":"missci-reconstructing-fallacies-in","title":"Missci: Reconstructing Fallacies in Misrepresented Science","date":"2024-06-05","arxiv_id":"2406.03181","n_code_links":2,"syntology":null},{"paper":"/paper/checkembed-effective-verification-of-llm","slug":"checkembed-effective-verification-of-llm","title":"CheckEmbed: Effective Verification of LLM Solutions to Open-Ended Tasks","date":"2024-06-04","arxiv_id":"2406.02524","n_code_links":1,"syntology":null},{"paper":null,"slug":"occamllm-fast-and-exact-language-model","title":"OccamLLM: Fast and Exact Language Model Arithmetic in a Single Step","date":"2024-06-04","arxiv_id":"2406.06576","n_code_links":0,"syntology":null},{"paper":"/paper/randomized-geometric-algebra-methods-for","slug":"randomized-geometric-algebra-methods-for","title":"Randomized Geometric Algebra Methods for Convex Neural Networks","date":"2024-06-04","arxiv_id":"2406.02806","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["pilancilab/Randomized-Geometric-Algebra-Methods-for-Convex-Neural-Networks"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"annotation-guidelines-based-knowledge","title":"Annotation Guidelines-Based Knowledge Augmentation: Towards Enhancing Large Language Models for Educational Text Classification","date":"2024-06-03","arxiv_id":"2406.00954","n_code_links":0,"syntology":null},{"paper":"/paper/spatialrgpt-grounded-spatial-reasoning-in","slug":"spatialrgpt-grounded-spatial-reasoning-in","title":"SpatialRGPT: Grounded Spatial Reasoning in Vision Language Models","date":"2024-06-03","arxiv_id":"2406.01584","n_code_links":1,"syntology":null},{"paper":null,"slug":"focus-forging-originality-through-contrastive","title":"FOCUS: Forging Originality through Contrastive Use in Self-Plagiarism for Language Models","date":"2024-06-02","arxiv_id":"2406.00839","n_code_links":0,"syntology":null},{"paper":null,"slug":"pretrained-hybrids-with-mad-skills","title":"Pretrained Hybrids with MAD Skills","date":"2024-06-02","arxiv_id":"2406.00894","n_code_links":0,"syntology":null},{"paper":null,"slug":"2406-07572","title":"Domain-specific ReAct for physics-integrated iterative modeling: A case study of LLM agents for gas path analysis of gas turbines","date":"2024-06-01","arxiv_id":"2406.07572","n_code_links":0,"syntology":null},{"paper":"/paper/generative-ai-voting-fair-collective-choice","slug":"generative-ai-voting-fair-collective-choice","title":"Generative AI Voting: Fair Collective Choice is Resilient to LLM Biases and Inconsistencies","date":"2024-05-31","arxiv_id":"2406.11871","n_code_links":1,"syntology":null},{"paper":"/paper/hard-cases-detection-in-motion-prediction-by","slug":"hard-cases-detection-in-motion-prediction-by","title":"Hard Cases Detection in Motion Prediction by Vision-Language Foundation Models","date":"2024-05-31","arxiv_id":"2405.20991","n_code_links":1,"syntology":null},{"paper":"/paper/llamea-a-large-language-model-evolutionary","slug":"llamea-a-large-language-model-evolutionary","title":"LLaMEA: A Large Language Model Evolutionary Algorithm for Automatically Generating Metaheuristics","date":"2024-05-30","arxiv_id":"2405.20132","n_code_links":2,"syntology":{"ran":6,"of":6,"n_ran_checked":6,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["nikivanstein/LLaMEA"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["listed","official"]}}},{"paper":null,"slug":"significance-of-chain-of-thought-in-gender","title":"Significance of Chain of Thought in Gender Bias Mitigation for English-Dravidian Machine Translation","date":"2024-05-30","arxiv_id":"2405.19701","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-multi-source-retrieval-question-answering","title":"A Multi-Source Retrieval Question Answering Framework Based on RAG","date":"2024-05-29","arxiv_id":"2405.19207","n_code_links":0,"syntology":null},{"paper":null,"slug":"can-gpt-redefine-medical-understanding","title":"Can GPT Redefine Medical Understanding? Evaluating GPT on Biomedical Machine Reading Comprehension","date":"2024-05-29","arxiv_id":"2405.18682","n_code_links":0,"syntology":null},{"paper":"/paper/map-neo-highly-capable-and-transparent","slug":"map-neo-highly-capable-and-transparent","title":"MAP-Neo: Highly Capable and Transparent Bilingual Large Language Model Series","date":"2024-05-29","arxiv_id":"2405.19327","n_code_links":1,"syntology":{"ran":12,"of":14,"n_ran_checked":12,"n_instrument":0,"unverified":2,"pointer_only":14,"phrase":"12 ran (of which 0 constructed an object rather than computing a result; 12 with no instrument failure: 0 honoured, 0 violated, 12 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["multimodal-art-projection/map-neo"],"state":"official (archive's flag): 12 ran","n_ran":12,"n_constructed":0,"n_ran_no_instrument_failure":12,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/llms-and-memorization-on-quality-and","slug":"llms-and-memorization-on-quality-and","title":"LLMs and Memorization: On Quality and Specificity of Copyright Compliance","date":"2024-05-28","arxiv_id":"2405.18492","n_code_links":1,"syntology":null},{"paper":null,"slug":"performance-evaluation-of-reddit-comments","title":"Performance evaluation of Reddit Comments using Machine Learning and Natural Language Processing methods in Sentiment Analysis","date":"2024-05-27","arxiv_id":"2405.16810","n_code_links":0,"syntology":null},{"paper":"/paper/accelerating-transformers-with-spectrum-1","slug":"accelerating-transformers-with-spectrum-1","title":"Accelerating Transformers with Spectrum-Preserving Token Merging","date":"2024-05-25","arxiv_id":"2405.16148","n_code_links":1,"syntology":{"ran":6,"of":7,"n_ran_checked":2,"n_instrument":4,"unverified":1,"pointer_only":7,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 2 honoured, 0 violated, 0 with no contract checked; 4 where Syntology's instrument failed) · 1 unverified","official":{"repos":["hchautran/PiToMe"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/learning-the-language-of-protein-structure","slug":"learning-the-language-of-protein-structure","title":"Learning the Language of Protein Structure","date":"2024-05-24","arxiv_id":"2405.15840","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["instadeepai/protein-structure-tokenizer"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"the-impact-and-opportunities-of-generative-ai","title":"The Impact and Opportunities of Generative AI in Fact-Checking","date":"2024-05-24","arxiv_id":"2405.15985","n_code_links":0,"syntology":null},{"paper":"/paper/wise-rethinking-the-knowledge-memory-for","slug":"wise-rethinking-the-knowledge-memory-for","title":"WISE: Rethinking the Knowledge Memory for Lifelong Model Editing of Large Language Models","date":"2024-05-23","arxiv_id":"2405.14768","n_code_links":1,"syntology":{"ran":13,"of":16,"n_ran_checked":4,"n_instrument":9,"unverified":3,"pointer_only":0,"phrase":"13 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 1 honoured, 0 violated, 3 with no contract checked; 9 where Syntology's instrument failed) · 3 unverified","official":{"repos":["zjunlp/easyedit"],"state":"official (archive's flag): 13 ran","n_ran":13,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"ku-dmis-at-ehrsql-2024-generating-sql-query","title":"KU-DMIS at EHRSQL 2024:Generating SQL query via question templatization in EHR","date":"2024-05-22","arxiv_id":"2406.00014","n_code_links":0,"syntology":null},{"paper":null,"slug":"how-reliable-ai-chatbots-are-for-disease","title":"How Reliable AI Chatbots are for Disease Prediction from Patient Complaints?","date":"2024-05-21","arxiv_id":"2405.13219","n_code_links":0,"syntology":null},{"paper":null,"slug":"investigating-persuasion-techniques-in-arabic","title":"Investigating Persuasion Techniques in Arabic: An Empirical Study Leveraging Large Language Models","date":"2024-05-21","arxiv_id":"2405.12884","n_code_links":0,"syntology":null},{"paper":"/paper/evaluating-and-modeling-social-intelligence-a","slug":"evaluating-and-modeling-social-intelligence-a","title":"Evaluating and Modeling Social Intelligence: A Comparative Study of Human and AI Capabilities","date":"2024-05-20","arxiv_id":"2405.11841","n_code_links":1,"syntology":null},{"paper":"/paper/human-centered-llm-agent-user-interface-a","slug":"human-centered-llm-agent-user-interface-a","title":"Human-Centered LLM-Agent User Interface: A Position Paper","date":"2024-05-19","arxiv_id":"2405.13050","n_code_links":1,"syntology":null},{"paper":"/paper/your-transformer-is-secretly-linear","slug":"your-transformer-is-secretly-linear","title":"Your Transformer is Secretly Linear","date":"2024-05-19","arxiv_id":"2405.12250","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":3,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["AIRI-Institute/LLM-Microscope"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/language-models-can-exploit-cross-task-in","slug":"language-models-can-exploit-cross-task-in","title":"Language Models can Exploit Cross-Task In-context Learning for Data-Scarce Novel Tasks","date":"2024-05-17","arxiv_id":"2405.10548","n_code_links":1,"syntology":{"ran":7,"of":9,"n_ran_checked":7,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["c-anwoy/cross-task-icl"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":7,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"gpt-store-mining-and-analysis","title":"GPT Store Mining and Analysis","date":"2024-05-16","arxiv_id":"2405.10210","n_code_links":0,"syntology":null}],"record_sha256":"6b01ca24b7eee59fe60b56f380dea2d86e24acd16dde9c2c88ed9055fd09f0ab","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}