{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/task/chatbot/papers/2","list_of":"/task/chatbot","task":"Chatbot","archive":{"snapshot":"2025-07-28"},"key_notes":{"n_ran_checked":"legacy name, kept unchanged so existing readers do not break: it counts the samples that ran with no instrument failure (honoured, violated, and ran with no contract checked); it does not mean a contract was checked, and the pages print it as 'K with no instrument failure', not 'K checked'","n_constructed":"a sub-count of the samples that ran, never subtracted from them and never a failure: an executed sample whose run returned an instance of its own class (fixture_out_type equals the entry name): the run built an object and did not compute a result (Syntology's RAN record, counts.constructed)"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"repositories listed in the archive (most first), then date (newest first), then slug","page":2,"pages_in_order":10,"rows_per_page":100,"rows":[101,200],"of":971,"counts":{"archive_papers_tagged":971,"with_a_code_link":269,"where_syntology_ran_a_sample":56,"not_listed_spam_title":0,"listed":971,"listed_where_code_ran":56,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":45,"every_run_a_failure_of_syntologys_instrument":11,"listed_with_a_run_with_no_instrument_failure":45,"listed_every_run_a_failure_of_syntologys_instrument":11,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/task/chatbot","prev":"/task/chatbot","next":"/task/chatbot/papers/3","papers":[{"url":"/paper/conditioning-chat-gpt-for-information","slug":"conditioning-chat-gpt-for-information","title":"Unipa-GPT: Large Language Models for university-oriented QA in Italian","date":"2024-07-19","arxiv_id":"2407.14246","repositories_listed":1,"syntology":null},{"url":"/paper/trust-no-bot-discovering-personal-disclosures","slug":"trust-no-bot-discovering-personal-disclosures","title":"Trust No Bot: Discovering Personal Disclosures in Human-LLM Conversations in the Wild","date":"2024-07-16","arxiv_id":"2407.11438","repositories_listed":1,"syntology":null},{"url":"/paper/representing-rule-based-chatbots-with","slug":"representing-rule-based-chatbots-with","title":"Representing Rule-based Chatbots with Transformers","date":"2024-07-15","arxiv_id":"2407.10949","repositories_listed":1,"syntology":null},{"url":"/paper/a-chatbot-for-asylum-seeking-migrants-in","slug":"a-chatbot-for-asylum-seeking-migrants-in","title":"A Chatbot for Asylum-Seeking Migrants in Europe","date":"2024-07-12","arxiv_id":"2407.09197","repositories_listed":1,"syntology":null},{"url":"/paper/llm-roleplay-simulating-human-chatbot","slug":"llm-roleplay-simulating-human-chatbot","title":"LLM Roleplay: Simulating Human-Chatbot Interaction","date":"2024-07-04","arxiv_id":"2407.03974","repositories_listed":1,"syntology":null},{"url":"/paper/convocache-smart-re-use-of-chatbot-responses","slug":"convocache-smart-re-use-of-chatbot-responses","title":"ConvoCache: Smart Re-Use of Chatbot Responses","date":"2024-06-26","arxiv_id":"2406.18133","repositories_listed":1,"syntology":null},{"url":"/paper/eden-empathetic-dialogues-for-english","slug":"eden-empathetic-dialogues-for-english","title":"EDEN: Empathetic Dialogues for English learning","date":"2024-06-25","arxiv_id":"2406.17982","repositories_listed":1,"syntology":null},{"url":"/paper/designing-a-dashboard-for-transparency-and","slug":"designing-a-dashboard-for-transparency-and","title":"Designing a Dashboard for Transparency and Control of Conversational AI","date":"2024-06-12","arxiv_id":"2406.07882","repositories_listed":1,"syntology":{"n":25,"n_ran":21,"n_constructed":0,"n_ran_checked":21,"n_instrument":0,"n_unverified":4,"n_honours":0,"n_violates":0,"n_no_contract":21,"n_pointer_only":0,"phrase":"21 ran (of which 0 constructed an object rather than computing a result; 21 with no instrument failure: 0 honoured, 0 violated, 21 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","sample_list":"/paper/designing-a-dashboard-for-transparency-and#ran","syntology_url":"https://syntology.ai/paper/2406.07882","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2406.07882"}},"official":{"repos":["yc015/talktuner-chatbot-llm-dashboard"],"state":"official (archive's flag): 21 ran","n_ran":21,"n_constructed":0,"n_ran_no_instrument_failure":21,"n_unverified":4,"ran_from_kinds":["official"]}}},{"url":"/paper/language-model-council-benchmarking","slug":"language-model-council-benchmarking","title":"Language Model Council: Democratically Benchmarking Foundation Models on Highly Subjective Tasks","date":"2024-06-12","arxiv_id":"2406.08598","repositories_listed":1,"syntology":null},{"url":"/paper/tailoring-generative-ai-chatbots-for","slug":"tailoring-generative-ai-chatbots-for","title":"Tailoring Generative AI Chatbots for Multiethnic Communities in Disaster Preparedness Communication: Extending the CASA Paradigm","date":"2024-06-12","arxiv_id":"2406.08411","repositories_listed":1,"syntology":null},{"url":"/paper/wildbench-benchmarking-llms-with-challenging","slug":"wildbench-benchmarking-llms-with-challenging","title":"WildBench: Benchmarking LLMs with Challenging Tasks from Real Users in the Wild","date":"2024-06-07","arxiv_id":"2406.04770","repositories_listed":1,"syntology":{"n":9,"n_ran":5,"n_constructed":0,"n_ran_checked":5,"n_instrument":0,"n_unverified":4,"n_honours":0,"n_violates":0,"n_no_contract":5,"n_pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","sample_list":"/paper/wildbench-benchmarking-llms-with-challenging#ran","syntology_url":"https://syntology.ai/paper/2406.04770","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2406.04770"}},"official":{"repos":["allenai/wildbench"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":4,"ran_from_kinds":["official"]}}},{"url":"/paper/demo-soccer-information-retrieval-via-natural","slug":"demo-soccer-information-retrieval-via-natural","title":"Demo: Soccer Information Retrieval via Natural Queries using SoccerRAG","date":"2024-06-03","arxiv_id":"2406.01280","repositories_listed":1,"syntology":null},{"url":"/paper/inverse-constitutional-ai-compressing","slug":"inverse-constitutional-ai-compressing","title":"Inverse Constitutional AI: Compressing Preferences into Principles","date":"2024-06-02","arxiv_id":"2406.06560","repositories_listed":1,"syntology":{"n":10,"n_ran":6,"n_constructed":3,"n_ran_checked":3,"n_instrument":3,"n_unverified":4,"n_honours":0,"n_violates":0,"n_no_contract":3,"n_pointer_only":0,"phrase":"6 ran (of which 3 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 3 where Syntology's instrument failed) · 4 unverified","sample_list":"/paper/inverse-constitutional-ai-compressing#ran","syntology_url":"https://syntology.ai/paper/2406.06560","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2406.06560"}},"official":{"repos":["rdnfn/icai"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":3,"n_ran_no_instrument_failure":3,"n_unverified":4,"ran_from_kinds":["official"]}}},{"url":"/paper/auto-arena-of-llms-automating-llm-evaluations","slug":"auto-arena-of-llms-automating-llm-evaluations","title":"Auto-Arena: Automating LLM Evaluations with Agent Peer Battles and Committee Discussions","date":"2024-05-30","arxiv_id":"2405.20267","repositories_listed":1,"syntology":null},{"url":"/paper/designing-an-evaluation-framework-for-large","slug":"designing-an-evaluation-framework-for-large","title":"Designing an Evaluation Framework for Large Language Models in Astronomy Research","date":"2024-05-30","arxiv_id":"2405.20389","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":0,"n_instrument":1,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/designing-an-evaluation-framework-for-large#ran","syntology_url":"https://syntology.ai/paper/2405.20389","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2405.20389"}},"official":{"repos":["jsalt2024-evaluating-llms-for-astronomy/astro-arxiv-bot"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/duanzai-slang-enhanced-llm-with-prompt-for","slug":"duanzai-slang-enhanced-llm-with-prompt-for","title":"DuanzAI: Slang-Enhanced LLM with Prompt for Humor Understanding","date":"2024-05-23","arxiv_id":"2405.15818","repositories_listed":1,"syntology":null},{"url":"/paper/evaluating-large-language-models-with-human","slug":"evaluating-large-language-models-with-human","title":"Evaluating Large Language Models with Human Feedback: Establishing a Swedish Benchmark","date":"2024-05-22","arxiv_id":"2405.14006","repositories_listed":1,"syntology":null},{"url":"/paper/can-ai-relate-testing-large-language-model","slug":"can-ai-relate-testing-large-language-model","title":"Can AI Relate: Testing Large Language Model Response for Mental Health Support","date":"2024-05-20","arxiv_id":"2405.12021","repositories_listed":1,"syntology":{"n":10,"n_ran":10,"n_constructed":0,"n_ran_checked":10,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":10,"n_pointer_only":0,"phrase":"10 ran (of which 0 constructed an object rather than computing a result; 10 with no instrument failure: 0 honoured, 0 violated, 10 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/can-ai-relate-testing-large-language-model#ran","syntology_url":"https://syntology.ai/paper/2405.12021","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2405.12021"}},"official":{"repos":["skgabriel/mh-eval"],"state":"official (archive's flag): 10 ran","n_ran":10,"n_constructed":0,"n_ran_no_instrument_failure":10,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/tailoring-vaccine-messaging-with-common","slug":"tailoring-vaccine-messaging-with-common","title":"Tailoring Vaccine Messaging with Common-Ground Opinions","date":"2024-05-17","arxiv_id":"2405.10861","repositories_listed":1,"syntology":null},{"url":"/paper/conformity-confabulation-and-impersonation","slug":"conformity-confabulation-and-impersonation","title":"Persona Inconstancy in Multi-Agent LLM Collaboration: Conformity, Confabulation, and Impersonation","date":"2024-05-06","arxiv_id":"2405.03862","repositories_listed":1,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":3,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":3,"n_pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/conformity-confabulation-and-impersonation#ran","syntology_url":"https://syntology.ai/paper/2405.03862","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2405.03862"}},"official":{"repos":["baltaci-r/CulturedAgents"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/meddoc-bot-a-chat-tool-for-comparative","slug":"meddoc-bot-a-chat-tool-for-comparative","title":"MedDoc-Bot: A Chat Tool for Comparative Analysis of Large Language Models in the Context of the Pediatric Hypertension Guideline","date":"2024-05-06","arxiv_id":"2405.03359","repositories_listed":1,"syntology":null},{"url":"/paper/physics-event-classification-using-large","slug":"physics-event-classification-using-large","title":"Physics Event Classification Using Large Language Models","date":"2024-04-05","arxiv_id":"2404.05752","repositories_listed":1,"syntology":{"n":3,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":3,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/physics-event-classification-using-large#ran","syntology_url":"https://syntology.ai/paper/2404.05752","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2404.05752"}},"official":{"repos":["ai4eic/ai4eichackathon2023-streamlit"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/facilitating-pornographic-text-detection-for","slug":"facilitating-pornographic-text-detection-for","title":"Facilitating Pornographic Text Detection for Open-Domain Dialogue Systems via Knowledge Distillation of Large Language Models","date":"2024-03-20","arxiv_id":"2403.13250","repositories_listed":1,"syntology":null},{"url":"/paper/characteristic-ai-agents-via-large-language","slug":"characteristic-ai-agents-via-large-language","title":"Characteristic AI Agents via Large Language Models","date":"2024-03-19","arxiv_id":"2403.12368","repositories_listed":1,"syntology":null},{"url":"/paper/deepseek-vl-towards-real-world-vision","slug":"deepseek-vl-towards-real-world-vision","title":"DeepSeek-VL: Towards Real-World Vision-Language Understanding","date":"2024-03-08","arxiv_id":"2403.05525","repositories_listed":1,"syntology":{"n":11,"n_ran":11,"n_constructed":0,"n_ran_checked":7,"n_instrument":4,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":7,"n_pointer_only":1,"phrase":"11 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 4 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/deepseek-vl-towards-real-world-vision#ran","syntology_url":"https://syntology.ai/paper/2403.05525","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2403.05525"}},"official":{"repos":["deepseek-ai/deepseek-vl"],"state":"official (archive's flag): 11 ran","n_ran":11,"n_constructed":0,"n_ran_no_instrument_failure":7,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/yi-open-foundation-models-by-01-ai","slug":"yi-open-foundation-models-by-01-ai","title":"Yi: Open Foundation Models by 01.AI","date":"2024-03-07","arxiv_id":"2403.04652","repositories_listed":1,"syntology":{"n":8,"n_ran":2,"n_constructed":0,"n_ran_checked":0,"n_instrument":2,"n_unverified":6,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 6 unverified","sample_list":"/paper/yi-open-foundation-models-by-01-ai#ran","syntology_url":"https://syntology.ai/paper/2403.04652","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2403.04652"}},"official":{"repos":["01-ai/yi"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":6,"ran_from_kinds":["official"]}}},{"url":"/paper/making-them-ask-and-answer-jailbreaking-large","slug":"making-them-ask-and-answer-jailbreaking-large","title":"Making Them Ask and Answer: Jailbreaking Large Language Models in Few Queries via Disguise and Reconstruction","date":"2024-02-28","arxiv_id":"2402.18104","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":0,"n_instrument":1,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/making-them-ask-and-answer-jailbreaking-large#ran","syntology_url":"https://syntology.ai/paper/2402.18104","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.18104"}},"official":{"repos":["llm-dra/dra"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/kodialogbench-evaluating-conversational","slug":"kodialogbench-evaluating-conversational","title":"KoDialogBench: Evaluating Conversational Understanding of Language Models with Korean Dialogue Benchmark","date":"2024-02-27","arxiv_id":"2402.17377","repositories_listed":1,"syntology":null},{"url":"/paper/prediction-powered-ranking-of-large-language","slug":"prediction-powered-ranking-of-large-language","title":"Prediction-Powered Ranking of Large Language Models","date":"2024-02-27","arxiv_id":"2402.17826","repositories_listed":1,"syntology":{"n":8,"n_ran":6,"n_constructed":0,"n_ran_checked":5,"n_instrument":1,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":5,"n_pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/prediction-powered-ranking-of-large-language#ran","syntology_url":"https://syntology.ai/paper/2402.17826","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.17826"}},"official":{"repos":["networks-learning/prediction-powered-ranking"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/asem-enhancing-empathy-in-chatbot-through","slug":"asem-enhancing-empathy-in-chatbot-through","title":"ASEM: Enhancing Empathy in Chatbot through Attention-based Sentiment and Emotion Modeling","date":"2024-02-25","arxiv_id":"2402.16194","repositories_listed":1,"syntology":null},{"url":"/paper/citation-enhanced-generation-for-llm-based","slug":"citation-enhanced-generation-for-llm-based","title":"Citation-Enhanced Generation for LLM-based Chatbots","date":"2024-02-25","arxiv_id":"2402.16063","repositories_listed":1,"syntology":null},{"url":"/paper/hypotermqa-hypothetical-terms-dataset-for","slug":"hypotermqa-hypothetical-terms-dataset-for","title":"HypoTermQA: Hypothetical Terms Dataset for Benchmarking Hallucination Tendency of LLMs","date":"2024-02-25","arxiv_id":"2402.16211","repositories_listed":1,"syntology":null},{"url":"/paper/compress-to-impress-unleashing-the-potential","slug":"compress-to-impress-unleashing-the-potential","title":"Compress to Impress: Unleashing the Potential of Compressive Memory in Real-World Long-Term Conversations","date":"2024-02-19","arxiv_id":"2402.11975","repositories_listed":1,"syntology":null},{"url":"/paper/safedecoding-defending-against-jailbreak","slug":"safedecoding-defending-against-jailbreak","title":"SafeDecoding: Defending against Jailbreak Attacks via Safety-Aware Decoding","date":"2024-02-14","arxiv_id":"2402.08983","repositories_listed":1,"syntology":{"n":1,"n_ran":0,"n_constructed":0,"n_ran_checked":0,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"0 ran · 1 unverified","sample_list":"/paper/safedecoding-defending-against-jailbreak#ran","syntology_url":"https://syntology.ai/paper/2402.08983","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.08983"}},"official":{"repos":["uw-nsl/safedecoding"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":[]}}},{"url":"/paper/measuring-and-controlling-instruction-in","slug":"measuring-and-controlling-instruction-in","title":"Measuring and Controlling Instruction (In)Stability in Language Model Dialogs","date":"2024-02-13","arxiv_id":"2402.10962","repositories_listed":1,"syntology":{"n":5,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":3,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":5,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","sample_list":"/paper/measuring-and-controlling-instruction-in#ran","syntology_url":"https://syntology.ai/paper/2402.10962","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.10962"}},"official":{"repos":["likenneth/persona_drift"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":3,"ran_from_kinds":["official"]}}},{"url":"/paper/cataractbot-an-llm-powered-expert-in-the-loop","slug":"cataractbot-an-llm-powered-expert-in-the-loop","title":"CataractBot: An LLM-Powered Expert-in-the-Loop Chatbot for Cataract Patients","date":"2024-02-07","arxiv_id":"2402.04620","repositories_listed":1,"syntology":null},{"url":"/paper/hydragen-high-throughput-llm-inference-with","slug":"hydragen-high-throughput-llm-inference-with","title":"Hydragen: High-Throughput LLM Inference with Shared Prefixes","date":"2024-02-07","arxiv_id":"2402.05099","repositories_listed":1,"syntology":{"n":7,"n_ran":6,"n_constructed":0,"n_ran_checked":6,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":6,"n_pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/hydragen-high-throughput-llm-inference-with#ran","syntology_url":"https://syntology.ai/paper/2402.05099","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.05099"}},"official":{"repos":["jordan-benjamin/hydragen"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/from-rag-to-qa-rag-integrating-generative-ai","slug":"from-rag-to-qa-rag-integrating-generative-ai","title":"From RAG to QA-RAG: Integrating Generative AI for Pharmaceutical Regulatory Compliance Process","date":"2024-01-26","arxiv_id":"2402.01717","repositories_listed":1,"syntology":null},{"url":"/paper/walert-putting-conversational-search","slug":"walert-putting-conversational-search","title":"Walert: Putting Conversational Search Knowledge into Action by Building and Evaluating a Large Language Model-Powered Chatbot","date":"2024-01-14","arxiv_id":"2401.07216","repositories_listed":1,"syntology":null},{"url":"/paper/quokka-an-open-source-large-language-model","slug":"quokka-an-open-source-large-language-model","title":"Quokka: An Open-source Large Language Model ChatBot for Material Science","date":"2024-01-02","arxiv_id":"2401.01089","repositories_listed":1,"syntology":null},{"url":"/paper/llm4eda-emerging-progress-in-large-language","slug":"llm4eda-emerging-progress-in-large-language","title":"LLM4EDA: Emerging Progress in Large Language Models for Electronic Design Automation","date":"2023-12-28","arxiv_id":"2401.12224","repositories_listed":1,"syntology":null},{"url":"/paper/align-on-the-fly-adapting-chatbot-behavior-to","slug":"align-on-the-fly-adapting-chatbot-behavior-to","title":"Align on the Fly: Adapting Chatbot Behavior to Established Norms","date":"2023-12-26","arxiv_id":"2312.15907","repositories_listed":1,"syntology":{"n":2,"n_ran":2,"n_constructed":0,"n_ran_checked":0,"n_instrument":2,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/align-on-the-fly-adapting-chatbot-behavior-to#ran","syntology_url":"https://syntology.ai/paper/2312.15907","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2312.15907"}},"official":{"repos":["gair-nlp/opo"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/designing-guiding-principles-for-nlp-for","slug":"designing-guiding-principles-for-nlp-for","title":"NLP for Maternal Healthcare: Perspectives and Guiding Principles in the Age of LLMs","date":"2023-12-19","arxiv_id":"2312.11803","repositories_listed":1,"syntology":null},{"url":"/paper/faithful-persona-based-conversational-dataset","slug":"faithful-persona-based-conversational-dataset","title":"Faithful Persona-based Conversational Dataset Generation with Large Language Models","date":"2023-12-15","arxiv_id":"2312.10007","repositories_listed":1,"syntology":null},{"url":"/paper/dr-jekyll-and-mr-hyde-two-faces-of-llms","slug":"dr-jekyll-and-mr-hyde-two-faces-of-llms","title":"Dr. Jekyll and Mr. Hyde: Two Faces of LLMs","date":"2023-12-06","arxiv_id":"2312.03853","repositories_listed":1,"syntology":null},{"url":"/paper/language-model-alignment-with-elastic-reset-1","slug":"language-model-alignment-with-elastic-reset-1","title":"Language Model Alignment with Elastic Reset","date":"2023-12-06","arxiv_id":"2312.07551","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":0,"n_instrument":1,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/language-model-alignment-with-elastic-reset-1#ran","syntology_url":"https://syntology.ai/paper/2312.07551","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2312.07551"}},"official":{"repos":["mnoukhov/elastic-reset"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/mka-a-scalable-medical-knowledge-assisted","slug":"mka-a-scalable-medical-knowledge-assisted","title":"MKA: A Scalable Medical Knowledge Assisted Mechanism for Generative Models on Medical Conversation Tasks","date":"2023-12-05","arxiv_id":"2312.02496","repositories_listed":1,"syntology":null},{"url":"/paper/efficient-stitchable-task-adaptation","slug":"efficient-stitchable-task-adaptation","title":"Efficient Stitchable Task Adaptation","date":"2023-11-29","arxiv_id":"2311.17352","repositories_listed":1,"syntology":null},{"url":"/paper/learning-to-love-diligent-trolls-accounting","slug":"learning-to-love-diligent-trolls-accounting","title":"Learning to love diligent trolls: Accounting for rater effects in the dialogue safety task","date":"2023-10-30","arxiv_id":"2310.19271","repositories_listed":1,"syntology":null},{"url":"/paper/ina-an-integrative-approach-for-enhancing","slug":"ina-an-integrative-approach-for-enhancing","title":"INA: An Integrative Approach for Enhancing Negotiation Strategies with Reward-Based Dialogue System","date":"2023-10-27","arxiv_id":"2310.18207","repositories_listed":1,"syntology":null},{"url":"/paper/the-impact-of-using-an-ai-chatbot-to-respond","slug":"the-impact-of-using-an-ai-chatbot-to-respond","title":"The impact of responding to patient messages with large language model assistance","date":"2023-10-26","arxiv_id":"2310.17703","repositories_listed":1,"syntology":null},{"url":"/paper/a-multilingual-virtual-guide-for-self","slug":"a-multilingual-virtual-guide-for-self","title":"A Multilingual Virtual Guide for Self-Attachment Technique","date":"2023-10-25","arxiv_id":"2310.18366","repositories_listed":1,"syntology":null},{"url":"/paper/detection-of-news-written-by-the-chatgpt","slug":"detection-of-news-written-by-the-chatgpt","title":"Detection of news written by the ChatGPT through authorship attribution performed by a Bidirectional LSTM model","date":"2023-10-25","arxiv_id":"2310.16685","repositories_listed":1,"syntology":null},{"url":"/paper/can-you-follow-me-testing-situational","slug":"can-you-follow-me-testing-situational","title":"Can You Follow Me? Testing Situational Understanding in ChatGPT","date":"2023-10-24","arxiv_id":"2310.16135","repositories_listed":1,"syntology":null},{"url":"/paper/learning-from-free-text-human-feedback","slug":"learning-from-free-text-human-feedback","title":"Learning From Free-Text Human Feedback -- Collect New Datasets Or Extend Existing Ones?","date":"2023-10-24","arxiv_id":"2310.15758","repositories_listed":1,"syntology":{"n":2,"n_ran":1,"n_constructed":0,"n_ran_checked":1,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/learning-from-free-text-human-feedback#ran","syntology_url":"https://syntology.ai/paper/2310.15758","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2310.15758"}},"official":{"repos":["ukplab/emnlp2023-learning-from-free-text-human-feedback"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/bioimage-io-chatbot-a-personalized-assistant","slug":"bioimage-io-chatbot-a-personalized-assistant","title":"BioImage.IO Chatbot: A Community-Driven AI Assistant for Integrative Computational Bioimaging","date":"2023-10-23","arxiv_id":"2310.18351","repositories_listed":1,"syntology":null},{"url":"/paper/miracle-towards-personalized-dialogue","slug":"miracle-towards-personalized-dialogue","title":"MIRACLE: Towards Personalized Dialogue Generation with Latent-Space Multiple Personal Attribute Control","date":"2023-10-22","arxiv_id":"2310.18342","repositories_listed":1,"syntology":null},{"url":"/paper/building-persona-consistent-dialogue-agents","slug":"building-persona-consistent-dialogue-agents","title":"Building Persona Consistent Dialogue Agents with Offline Reinforcement Learning","date":"2023-10-16","arxiv_id":"2310.10735","repositories_listed":1,"syntology":null},{"url":"/paper/p5-plug-and-play-persona-prompting-for","slug":"p5-plug-and-play-persona-prompting-for","title":"P5: Plug-and-Play Persona Prompting for Personalized Response Selection","date":"2023-10-10","arxiv_id":"2310.06390","repositories_listed":1,"syntology":{"n":1,"n_ran":0,"n_constructed":0,"n_ran_checked":0,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":1,"phrase":"0 ran · 1 unverified","sample_list":"/paper/p5-plug-and-play-persona-prompting-for#ran","syntology_url":"https://syntology.ai/paper/2310.06390","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2310.06390"}},"official":{"repos":["rungjoo/plug-and-play-prompt-persona"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":[]}}},{"url":"/paper/boilerbot-a-reliable-task-oriented-chatbot","slug":"boilerbot-a-reliable-task-oriented-chatbot","title":"BoilerBot: A reliable task-oriented chatbot enhanced with large language models","date":"2023-10-03","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/mathvista-evaluating-mathematical-reasoning","slug":"mathvista-evaluating-mathematical-reasoning","title":"MathVista: Evaluating Mathematical Reasoning of Foundation Models in Visual Contexts","date":"2023-10-03","arxiv_id":"2310.02255","repositories_listed":1,"syntology":{"n":4,"n_ran":4,"n_constructed":0,"n_ran_checked":1,"n_instrument":3,"n_unverified":0,"n_honours":0,"n_violates":1,"n_no_contract":0,"n_pointer_only":4,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 1 violated, 0 with no contract checked; 3 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/mathvista-evaluating-mathematical-reasoning#ran","syntology_url":"https://syntology.ai/paper/2310.02255","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2310.02255"}},"official":null}},{"url":"/paper/an-ai-chatbot-for-explaining-deep","slug":"an-ai-chatbot-for-explaining-deep","title":"An AI Chatbot for Explaining Deep Reinforcement Learning Decisions of Service-oriented Systems","date":"2023-09-25","arxiv_id":"2309.14391","repositories_listed":1,"syntology":null},{"url":"/paper/how-robust-is-google-s-bard-to-adversarial","slug":"how-robust-is-google-s-bard-to-adversarial","title":"How Robust is Google's Bard to Adversarial Image Attacks?","date":"2023-09-21","arxiv_id":"2309.11751","repositories_listed":1,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":3,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":3,"n_pointer_only":3,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/how-robust-is-google-s-bard-to-adversarial#ran","syntology_url":"https://syntology.ai/paper/2309.11751","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2309.11751"}},"official":{"repos":["thu-ml/attack-bard"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/towards-joint-modeling-of-dialogue-response","slug":"towards-joint-modeling-of-dialogue-response","title":"Towards Joint Modeling of Dialogue Response and Speech Synthesis based on Large Language Model","date":"2023-09-20","arxiv_id":"2309.11000","repositories_listed":1,"syntology":null},{"url":"/paper/facilitating-nsfw-text-detection-in-open","slug":"facilitating-nsfw-text-detection-in-open","title":"Facilitating NSFW Text Detection in Open-Domain Dialogue Systems via Knowledge Distillation","date":"2023-09-18","arxiv_id":"2309.09749","repositories_listed":1,"syntology":null},{"url":"/paper/exploring-the-impact-of-human-evaluator-group","slug":"exploring-the-impact-of-human-evaluator-group","title":"Exploring the Impact of Human Evaluator Group on Chat-Oriented Dialogue Evaluation","date":"2023-09-14","arxiv_id":"2309.07998","repositories_listed":1,"syntology":null},{"url":"/paper/generative-social-choice","slug":"generative-social-choice","title":"Generative Social Choice","date":"2023-09-03","arxiv_id":"2309.01291","repositories_listed":1,"syntology":null},{"url":"/paper/diagnosing-infeasible-optimization-problems","slug":"diagnosing-infeasible-optimization-problems","title":"Diagnosing Infeasible Optimization Problems Using Large Language Models","date":"2023-08-23","arxiv_id":"2308.12923","repositories_listed":1,"syntology":null},{"url":"/paper/mulmarker-a-gpt-assisted-comprehensive","slug":"mulmarker-a-gpt-assisted-comprehensive","title":"MulMarker: a comprehensive framework for identifying multi-gene prognostic signatures","date":"2023-08-22","arxiv_id":"2308.11349","repositories_listed":1,"syntology":null},{"url":"/paper/llm4ts-two-stage-fine-tuning-for-time-series","slug":"llm4ts-two-stage-fine-tuning-for-time-series","title":"LLM4TS: Aligning Pre-Trained LLMs as Data-Efficient Time-Series Forecasters","date":"2023-08-16","arxiv_id":"2308.08469","repositories_listed":1,"syntology":{"n":8,"n_ran":6,"n_constructed":0,"n_ran_checked":6,"n_instrument":0,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":6,"n_pointer_only":2,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/llm4ts-two-stage-fine-tuning-for-time-series#ran","syntology_url":"https://syntology.ai/paper/2308.08469","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2308.08469"}},"official":{"repos":["blacksnail789521/LLM4TS"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/exploring-chatgpt-s-empathic-abilities","slug":"exploring-chatgpt-s-empathic-abilities","title":"Exploring ChatGPT's Empathic Abilities","date":"2023-08-07","arxiv_id":"2308.03527","repositories_listed":1,"syntology":null},{"url":"/paper/ontochatgpt-information-system-ontology","slug":"ontochatgpt-information-system-ontology","title":"OntoChatGPT Information System: Ontology-Driven Structured Prompts for ChatGPT Meta-Learning","date":"2023-07-11","arxiv_id":"2307.05082","repositories_listed":1,"syntology":null},{"url":"/paper/trac-trustworthy-retrieval-augmented-chatbot","slug":"trac-trustworthy-retrieval-augmented-chatbot","title":"TRAQ: Trustworthy Retrieval Augmented Question Answering via Conformal Prediction","date":"2023-07-07","arxiv_id":"2307.04642","repositories_listed":1,"syntology":null},{"url":"/paper/enhancing-dialogue-generation-via-dynamic","slug":"enhancing-dialogue-generation-via-dynamic","title":"Enhancing Dialogue Generation via Dynamic Graph Knowledge Aggregation","date":"2023-06-28","arxiv_id":"2306.16195","repositories_listed":1,"syntology":null},{"url":"/paper/bring-your-own-data-self-supervised","slug":"bring-your-own-data-self-supervised","title":"Bring Your Own Data! Self-Supervised Evaluation for Large Language Models","date":"2023-06-23","arxiv_id":"2306.13651","repositories_listed":1,"syntology":{"n":15,"n_ran":4,"n_constructed":0,"n_ran_checked":4,"n_instrument":0,"n_unverified":11,"n_honours":0,"n_violates":0,"n_no_contract":4,"n_pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 11 unverified","sample_list":"/paper/bring-your-own-data-self-supervised#ran","syntology_url":"https://syntology.ai/paper/2306.13651","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2306.13651"}},"official":{"repos":["neelsjain/byod"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":11,"ran_from_kinds":["official"]}}},{"url":"/paper/tracking-public-attitudes-toward-chatgpt-on","slug":"tracking-public-attitudes-toward-chatgpt-on","title":"Public Attitudes Toward ChatGPT on Twitter: Sentiments, Topics, and Occupations","date":"2023-06-22","arxiv_id":"2306.12951","repositories_listed":1,"syntology":null},{"url":"/paper/chatgpt-chemistry-assistant-for-text-mining","slug":"chatgpt-chemistry-assistant-for-text-mining","title":"ChatGPT Chemistry Assistant for Text Mining and Prediction of MOF Synthesis","date":"2023-06-20","arxiv_id":"2306.11296","repositories_listed":1,"syntology":null},{"url":"/paper/domain-specific-chatbots-for-science-using","slug":"domain-specific-chatbots-for-science-using","title":"Domain-specific ChatBots for Science using Embeddings","date":"2023-06-15","arxiv_id":"2306.10067","repositories_listed":1,"syntology":null},{"url":"/paper/datachat-prototyping-a-conversational-agent","slug":"datachat-prototyping-a-conversational-agent","title":"DataChat: Prototyping a Conversational Agent for Dataset Search and Visualization","date":"2023-05-26","arxiv_id":"2305.18358","repositories_listed":1,"syntology":null},{"url":"/paper/cheap-and-quick-efficient-vision-language","slug":"cheap-and-quick-efficient-vision-language","title":"Cheap and Quick: Efficient Vision-Language Instruction Tuning for Large Language Models","date":"2023-05-24","arxiv_id":"2305.15023","repositories_listed":1,"syntology":null},{"url":"/paper/wikichat-a-few-shot-llm-based-chatbot","slug":"wikichat-a-few-shot-llm-based-chatbot","title":"WikiChat: Stopping the Hallucination of Large Language Model Chatbots by Few-Shot Grounding on Wikipedia","date":"2023-05-23","arxiv_id":"2305.14292","repositories_listed":1,"syntology":{"n":7,"n_ran":5,"n_constructed":0,"n_ran_checked":5,"n_instrument":0,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":5,"n_pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/wikichat-a-few-shot-llm-based-chatbot#ran","syntology_url":"https://syntology.ai/paper/2305.14292","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2305.14292"}},"official":{"repos":["stanford-oval/wikichat"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/class-meet-spock-an-education-tutoring","slug":"class-meet-spock-an-education-tutoring","title":"CLASS: A Design Framework for building Intelligent Tutoring Systems based on Learning Science principles","date":"2023-05-22","arxiv_id":"2305.13272","repositories_listed":1,"syntology":{"n":5,"n_ran":4,"n_constructed":0,"n_ran_checked":3,"n_instrument":1,"n_unverified":1,"n_honours":1,"n_violates":2,"n_no_contract":0,"n_pointer_only":5,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 1 honoured, 2 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/class-meet-spock-an-education-tutoring#ran","syntology_url":"https://syntology.ai/paper/2305.13272","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2305.13272"}},"official":{"repos":["luffycodes/tutorbot-spock"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/boosting-distress-support-dialogue-responses","slug":"boosting-distress-support-dialogue-responses","title":"Boosting Distress Support Dialogue Responses with Motivational Interviewing Strategy","date":"2023-05-17","arxiv_id":"2305.10195","repositories_listed":1,"syntology":null},{"url":"/paper/memorybank-enhancing-large-language-models","slug":"memorybank-enhancing-large-language-models","title":"MemoryBank: Enhancing Large Language Models with Long-Term Memory","date":"2023-05-17","arxiv_id":"2305.10250","repositories_listed":1,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":0,"n_instrument":3,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 3 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/memorybank-enhancing-large-language-models#ran","syntology_url":"https://syntology.ai/paper/2305.10250","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2305.10250"}},"official":{"repos":["zhongwanjun/memorybank-siliconfriend"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/prompted-llms-as-chatbot-modules-for-long","slug":"prompted-llms-as-chatbot-modules-for-long","title":"Prompted LLMs as Chatbot Modules for Long Open-domain Conversation","date":"2023-05-08","arxiv_id":"2305.04533","repositories_listed":1,"syntology":{"n":11,"n_ran":1,"n_constructed":0,"n_ran_checked":1,"n_instrument":0,"n_unverified":10,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 10 unverified","sample_list":"/paper/prompted-llms-as-chatbot-modules-for-long#ran","syntology_url":"https://syntology.ai/paper/2305.04533","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2305.04533"}},"official":{"repos":["krafton-ai/mpc"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":10,"ran_from_kinds":["official"]}}},{"url":"/paper/semantic-space-grounded-weighted-decoding-for","slug":"semantic-space-grounded-weighted-decoding-for","title":"Semantic Space Grounded Weighted Decoding for Multi-Attribute Controllable Dialogue Generation","date":"2023-05-04","arxiv_id":"2305.02820","repositories_listed":1,"syntology":null},{"url":"/paper/smile-single-turn-to-multi-turn-inclusive","slug":"smile-single-turn-to-multi-turn-inclusive","title":"SMILE: Single-turn to Multi-turn Inclusive Language Expansion via ChatGPT for Mental Health Support","date":"2023-04-30","arxiv_id":"2305.00450","repositories_listed":1,"syntology":null},{"url":"/paper/boosting-big-brother-attacking-search-engines","slug":"boosting-big-brother-attacking-search-engines","title":"Boosting Big Brother: Attacking Search Engines with Encodings","date":"2023-04-27","arxiv_id":"2304.14031","repositories_listed":1,"syntology":null},{"url":"/paper/building-multimodal-ai-chatbots","slug":"building-multimodal-ai-chatbots","title":"Building Multimodal AI Chatbots","date":"2023-04-21","arxiv_id":"2305.03512","repositories_listed":1,"syntology":null},{"url":"/paper/v3det-vast-vocabulary-visual-detection","slug":"v3det-vast-vocabulary-visual-detection","title":"V3Det: Vast Vocabulary Visual Detection Dataset","date":"2023-04-07","arxiv_id":"2304.03752","repositories_listed":1,"syntology":{"n":11,"n_ran":8,"n_constructed":0,"n_ran_checked":6,"n_instrument":2,"n_unverified":3,"n_honours":0,"n_violates":0,"n_no_contract":6,"n_pointer_only":4,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 2 where Syntology's instrument failed) · 3 unverified","sample_list":"/paper/v3det-vast-vocabulary-visual-detection#ran","syntology_url":"https://syntology.ai/paper/2304.03752","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2304.03752"}},"official":null}},{"url":"/paper/ten-quick-tips-for-harnessing-the-power-of","slug":"ten-quick-tips-for-harnessing-the-power-of","title":"Ten Quick Tips for Harnessing the Power of ChatGPT/GPT-4 in Computational Biology","date":"2023-03-29","arxiv_id":"2303.16429","repositories_listed":1,"syntology":null},{"url":"/paper/on-the-robustness-of-chatgpt-an-adversarial","slug":"on-the-robustness-of-chatgpt-an-adversarial","title":"On the Robustness of ChatGPT: An Adversarial and Out-of-distribution Perspective","date":"2023-02-22","arxiv_id":"2302.12095","repositories_listed":1,"syntology":null},{"url":"/paper/chatgpt-jack-of-all-trades-master-of-none","slug":"chatgpt-jack-of-all-trades-master-of-none","title":"ChatGPT: Jack of all trades, master of none","date":"2023-02-21","arxiv_id":"2302.10724","repositories_listed":1,"syntology":{"n":2,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/chatgpt-jack-of-all-trades-master-of-none#ran","syntology_url":"https://syntology.ai/paper/2302.10724","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2302.10724"}},"official":{"repos":["clarin-pl/chatgpt-evaluation-01-2023"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/search-engine-augmented-dialogue-response","slug":"search-engine-augmented-dialogue-response","title":"Search-Engine-augmented Dialogue Response Generation with Cheaply Supervised Query Production","date":"2023-02-16","arxiv_id":"2302.09300","repositories_listed":1,"syntology":null},{"url":"/paper/prescriptive-process-monitoring-in","slug":"prescriptive-process-monitoring-in","title":"Prescriptive Process Monitoring in Intelligent Process Automation with Chatbot Orchestration","date":"2022-12-13","arxiv_id":"2212.06564","repositories_listed":1,"syntology":null},{"url":"/paper/dagfinn-a-conversational-conference-assistant","slug":"dagfinn-a-conversational-conference-assistant","title":"DAGFiNN: A Conversational Conference Assistant","date":"2022-11-29","arxiv_id":"2211.16281","repositories_listed":1,"syntology":null},{"url":"/paper/unified-multimodal-model-with-unlikelihood","slug":"unified-multimodal-model-with-unlikelihood","title":"Unified Multimodal Model with Unlikelihood Training for Visual Dialog","date":"2022-11-23","arxiv_id":"2211.13235","repositories_listed":1,"syntology":null},{"url":"/paper/causal-inference-for-chatting-handoff","slug":"causal-inference-for-chatting-handoff","title":"Causal Inference for Chatting Handoff","date":"2022-10-06","arxiv_id":"2210.02862","repositories_listed":1,"syntology":null},{"url":"/paper/ieval-interactive-evaluation-framework-for","slug":"ieval-interactive-evaluation-framework-for","title":"iEval: Interactive Evaluation Framework for Open-Domain Empathetic Chatbots","date":"2022-09-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/towards-boosting-the-open-domain-chatbot-with","slug":"towards-boosting-the-open-domain-chatbot-with","title":"Towards Boosting the Open-Domain Chatbot with Human Feedback","date":"2022-08-30","arxiv_id":"2208.14165","repositories_listed":1,"syntology":null}],"record_sha256":"d8b17688e92d95a4bf2a07ad7ac2293f97bcf11cb9e35012a6c4e08cdf251e4c","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}