{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/task/language-modelling/papers/37","list_of":"/task/language-modelling","task":"Language Modelling","archive":{"snapshot":"2025-07-28"},"key_notes":{"n_ran_checked":"legacy name, kept unchanged so existing readers do not break: it counts the samples that ran with no instrument failure (honoured, violated, and ran with no contract checked); it does not mean a contract was checked, and the pages print it as 'K with no instrument failure', not 'K checked'","n_constructed":"a sub-count of the samples that ran, never subtracted from them and never a failure: an executed sample whose run returned an instance of its own class (fixture_out_type equals the entry name): the run built an object and did not compute a result (Syntology's RAN record, counts.constructed)"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"repositories listed in the archive (most first), then date (newest first), then slug","page":37,"pages_in_order":177,"rows_per_page":100,"rows":[3601,3700],"of":17610,"counts":{"archive_papers_tagged":17610,"with_a_code_link":7012,"where_syntology_ran_a_sample":2428,"not_listed_spam_title":0,"listed":17610,"listed_where_code_ran":2428,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":2027,"every_run_a_failure_of_syntologys_instrument":401,"listed_with_a_run_with_no_instrument_failure":2027,"listed_every_run_a_failure_of_syntologys_instrument":401,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/task/language-modelling","prev":"/task/language-modelling/papers/36","next":"/task/language-modelling/papers/38","papers":[{"url":"/paper/preact-predicting-future-in-react-enhances","slug":"preact-predicting-future-in-react-enhances","title":"PreAct: Prediction Enhances Agent's Planning Ability","date":"2024-02-18","arxiv_id":"2402.11534","repositories_listed":1,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":2,"n_instrument":1,"n_unverified":0,"n_honours":1,"n_violates":0,"n_no_contract":1,"n_pointer_only":3,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 1 honoured, 0 violated, 1 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/preact-predicting-future-in-react-enhances#ran","syntology_url":"https://syntology.ai/paper/2402.11534","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.11534"}},"official":{"repos":["fu-dayuan/preact"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/stealthy-attack-on-large-language-model-based","slug":"stealthy-attack-on-large-language-model-based","title":"Stealthy Attack on Large Language Model based Recommendation","date":"2024-02-18","arxiv_id":"2402.14836","repositories_listed":1,"syntology":null},{"url":"/paper/controlled-text-generation-for-large-language","slug":"controlled-text-generation-for-large-language","title":"Controlled Text Generation for Large Language Model with Dynamic Attribute Graphs","date":"2024-02-17","arxiv_id":"2402.11218","repositories_listed":1,"syntology":null},{"url":"/paper/dissecting-human-and-llm-preferences","slug":"dissecting-human-and-llm-preferences","title":"Dissecting Human and LLM Preferences","date":"2024-02-17","arxiv_id":"2402.11296","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":0,"n_instrument":1,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/dissecting-human-and-llm-preferences#ran","syntology_url":"https://syntology.ai/paper/2402.11296","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.11296"}},"official":{"repos":["gair-nlp/preference-dissection"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/grasping-the-essentials-tailoring-large","slug":"grasping-the-essentials-tailoring-large","title":"Grasping the Essentials: Tailoring Large Language Models for Zero-Shot Relation Extraction","date":"2024-02-17","arxiv_id":"2402.11142","repositories_listed":1,"syntology":null},{"url":"/paper/i-learn-better-if-you-speak-my-language","slug":"i-learn-better-if-you-speak-my-language","title":"I Learn Better If You Speak My Language: Understanding the Superior Performance of Fine-Tuning Large Language Models with LLM-Generated Responses","date":"2024-02-17","arxiv_id":"2402.11192","repositories_listed":1,"syntology":{"n":9,"n_ran":5,"n_constructed":0,"n_ran_checked":5,"n_instrument":0,"n_unverified":4,"n_honours":0,"n_violates":0,"n_no_contract":5,"n_pointer_only":9,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","sample_list":"/paper/i-learn-better-if-you-speak-my-language#ran","syntology_url":"https://syntology.ai/paper/2402.11192","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.11192"}},"official":{"repos":["xuanren4470/i-learn-better-if-you-speak-my-language"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":4,"ran_from_kinds":["official"]}}},{"url":"/paper/training-language-model-agents-without","slug":"training-language-model-agents-without","title":"Offline Training of Language Model Agents with Functions as Learnable Weights","date":"2024-02-17","arxiv_id":"2402.11359","repositories_listed":1,"syntology":null},{"url":"/paper/zerog-investigating-cross-dataset-zero-shot","slug":"zerog-investigating-cross-dataset-zero-shot","title":"ZeroG: Investigating Cross-dataset Zero-shot Transferability in Graphs","date":"2024-02-17","arxiv_id":"2402.11235","repositories_listed":1,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":0,"n_instrument":3,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 3 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/zerog-investigating-cross-dataset-zero-shot#ran","syntology_url":"https://syntology.ai/paper/2402.11235","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.11235"}},"official":{"repos":["nineabyss/zerog"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/an-empirical-study-on-cross-lingual","slug":"an-empirical-study-on-cross-lingual","title":"An Empirical Study on Cross-lingual Vocabulary Adaptation for Efficient Language Model Inference","date":"2024-02-16","arxiv_id":"2402.10712","repositories_listed":1,"syntology":{"n":7,"n_ran":6,"n_constructed":0,"n_ran_checked":6,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":6,"n_pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/an-empirical-study-on-cross-lingual#ran","syntology_url":"https://syntology.ai/paper/2402.10712","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.10712"}},"official":{"repos":["gucci-j/llm-cva"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/question-instructed-visual-descriptions-for","slug":"question-instructed-visual-descriptions-for","title":"Question-Instructed Visual Descriptions for Zero-Shot Video Question Answering","date":"2024-02-16","arxiv_id":"2402.10698","repositories_listed":1,"syntology":null},{"url":"/paper/rag-driver-generalisable-driving-explanations","slug":"rag-driver-generalisable-driving-explanations","title":"RAG-Driver: Generalisable Driving Explanations with Retrieval-Augmented In-Context Learning in Multi-Modal Large Language Model","date":"2024-02-16","arxiv_id":"2402.10828","repositories_listed":1,"syntology":{"n":2,"n_ran":1,"n_constructed":0,"n_ran_checked":0,"n_instrument":1,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/rag-driver-generalisable-driving-explanations#ran","syntology_url":"https://syntology.ai/paper/2402.10828","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.10828"}},"official":null}},{"url":"/paper/unsupervised-llm-adaptation-for-question","slug":"unsupervised-llm-adaptation-for-question","title":"Where is the answer? Investigating Positional Bias in Language Model Knowledge Extraction","date":"2024-02-16","arxiv_id":"2402.12170","repositories_listed":1,"syntology":null},{"url":"/paper/both-matter-enhancing-the-emotional","slug":"both-matter-enhancing-the-emotional","title":"Both Matter: Enhancing the Emotional Intelligence of Large Language Models without Compromising the General Intelligence","date":"2024-02-15","arxiv_id":"2402.10073","repositories_listed":1,"syntology":null},{"url":"/paper/chemreasoner-heuristic-search-over-a-large","slug":"chemreasoner-heuristic-search-over-a-large","title":"ChemReasoner: Heuristic Search over a Large Language Model's Knowledge Space using Quantum-Chemical Feedback","date":"2024-02-15","arxiv_id":"2402.10980","repositories_listed":1,"syntology":{"n":7,"n_ran":5,"n_constructed":0,"n_ran_checked":0,"n_instrument":5,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":7,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 5 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/chemreasoner-heuristic-search-over-a-large#ran","syntology_url":"https://syntology.ai/paper/2402.10980","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.10980"}},"official":{"repos":["pnnl/chemreasoner"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/fast-vocabulary-transfer-for-language-model","slug":"fast-vocabulary-transfer-for-language-model","title":"Fast Vocabulary Transfer for Language Model Compression","date":"2024-02-15","arxiv_id":"2402.09977","repositories_listed":1,"syntology":null},{"url":"/paper/llm-based-federated-recommendation","slug":"llm-based-federated-recommendation","title":"A Federated Framework for LLM-based Recommendation","date":"2024-02-15","arxiv_id":"2402.09959","repositories_listed":1,"syntology":null},{"url":"/paper/on-the-safety-concerns-of-deploying-llms-vlms","slug":"on-the-safety-concerns-of-deploying-llms-vlms","title":"On the Vulnerability of LLM/VLM-Controlled Robotics","date":"2024-02-15","arxiv_id":"2402.10340","repositories_listed":1,"syntology":null},{"url":"/paper/optimus-scalable-optimization-modeling-with","slug":"optimus-scalable-optimization-modeling-with","title":"OptiMUS: Scalable Optimization Modeling with (MI)LP Solvers and Large Language Models","date":"2024-02-15","arxiv_id":"2402.10172","repositories_listed":1,"syntology":null},{"url":"/paper/visually-dehallucinative-instruction-1","slug":"visually-dehallucinative-instruction-1","title":"Visually Dehallucinative Instruction Generation: Know What You Don't Know","date":"2024-02-15","arxiv_id":"2402.09717","repositories_listed":1,"syntology":null},{"url":"/paper/chinese-mentalbert-domain-adaptive-pre","slug":"chinese-mentalbert-domain-adaptive-pre","title":"Chinese MentalBERT: Domain-Adaptive Pre-training on Social Media for Chinese Mental Health Text Analysis","date":"2024-02-14","arxiv_id":"2402.09151","repositories_listed":1,"syntology":null},{"url":"/paper/fgeo-drl-deductive-reasoning-for-geometric","slug":"fgeo-drl-deductive-reasoning-for-geometric","title":"FGeo-DRL: Deductive Reasoning for Geometric Problems through Deep Reinforcement Learning","date":"2024-02-14","arxiv_id":"2402.09051","repositories_listed":1,"syntology":null},{"url":"/paper/massively-multi-cultural-knowledge","slug":"massively-multi-cultural-knowledge","title":"Massively Multi-Cultural Knowledge Acquisition & LM Benchmarking","date":"2024-02-14","arxiv_id":"2402.09369","repositories_listed":1,"syntology":{"n":14,"n_ran":7,"n_constructed":0,"n_ran_checked":7,"n_instrument":0,"n_unverified":7,"n_honours":0,"n_violates":0,"n_no_contract":7,"n_pointer_only":0,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 0 where Syntology's instrument failed) · 7 unverified","sample_list":"/paper/massively-multi-cultural-knowledge#ran","syntology_url":"https://syntology.ai/paper/2402.09369","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.09369"}},"official":{"repos":["yrf1/llm-massivemulticulturenormsknowledge-nclb"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":7,"n_unverified":7,"ran_from_kinds":["official"]}}},{"url":"/paper/mustard-mastering-uniform-synthesis-of","slug":"mustard-mastering-uniform-synthesis-of","title":"MUSTARD: Mastering Uniform Synthesis of Theorem and Proof Data","date":"2024-02-14","arxiv_id":"2402.08957","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":0,"n_instrument":1,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/mustard-mastering-uniform-synthesis-of#ran","syntology_url":"https://syntology.ai/paper/2402.08957","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.08957"}},"official":{"repos":["eleanor-h/mustard"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/pretraining-vision-language-model-for","slug":"pretraining-vision-language-model-for","title":"Pretraining Vision-Language Model for Difference Visual Question Answering in Longitudinal Chest X-rays","date":"2024-02-14","arxiv_id":"2402.08966","repositories_listed":1,"syntology":null},{"url":"/paper/rapid-adoption-hidden-risks-the-dual-impact","slug":"rapid-adoption-hidden-risks-the-dual-impact","title":"Instruction Backdoor Attacks Against Customized LLMs","date":"2024-02-14","arxiv_id":"2402.09179","repositories_listed":1,"syntology":{"n":2,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/rapid-adoption-hidden-risks-the-dual-impact#ran","syntology_url":"https://syntology.ai/paper/2402.09179","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.09179"}},"official":{"repos":["zhangrui4041/instruction_backdoor_attack"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/scaling-the-authoring-of-autotutors-with","slug":"scaling-the-authoring-of-autotutors-with","title":"AutoTutor meets Large Language Models: A Language Model Tutor with Rich Pedagogy and Guardrails","date":"2024-02-14","arxiv_id":"2402.09216","repositories_listed":1,"syntology":{"n":10,"n_ran":4,"n_constructed":0,"n_ran_checked":4,"n_instrument":0,"n_unverified":6,"n_honours":0,"n_violates":0,"n_no_contract":4,"n_pointer_only":10,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 6 unverified","sample_list":"/paper/scaling-the-authoring-of-autotutors-with#ran","syntology_url":"https://syntology.ai/paper/2402.09216","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.09216"}},"official":{"repos":["eth-lre/mwptutor"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":6,"ran_from_kinds":["official"]}}},{"url":"/paper/tell-me-more-towards-implicit-user-intention","slug":"tell-me-more-towards-implicit-user-intention","title":"Tell Me More! Towards Implicit User Intention Understanding of Language Model Driven Agents","date":"2024-02-14","arxiv_id":"2402.09205","repositories_listed":1,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":0,"n_instrument":3,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 3 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/tell-me-more-towards-implicit-user-intention#ran","syntology_url":"https://syntology.ai/paper/2402.09205","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.09205"}},"official":{"repos":["hbx-hbx/mistral-interact"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/agent-smith-a-single-image-can-jailbreak-one","slug":"agent-smith-a-single-image-can-jailbreak-one","title":"Agent Smith: A Single Image Can Jailbreak One Million Multimodal LLM Agents Exponentially Fast","date":"2024-02-13","arxiv_id":"2402.08567","repositories_listed":1,"syntology":{"n":2,"n_ran":1,"n_constructed":0,"n_ran_checked":0,"n_instrument":1,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/agent-smith-a-single-image-can-jailbreak-one#ran","syntology_url":"https://syntology.ai/paper/2402.08567","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.08567"}},"official":{"repos":["sail-sg/agent-smith"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/evaluating-the-data-model-robustness-of-text","slug":"evaluating-the-data-model-robustness-of-text","title":"Evaluating the Data Model Robustness of Text-to-SQL Systems Based on Real User Queries","date":"2024-02-13","arxiv_id":"2402.08349","repositories_listed":1,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":3,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":3,"n_pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/evaluating-the-data-model-robustness-of-text#ran","syntology_url":"https://syntology.ai/paper/2402.08349","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.08349"}},"official":{"repos":["jf87/footballdb"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/measuring-and-controlling-instruction-in","slug":"measuring-and-controlling-instruction-in","title":"Measuring and Controlling Instruction (In)Stability in Language Model Dialogs","date":"2024-02-13","arxiv_id":"2402.10962","repositories_listed":1,"syntology":{"n":5,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":3,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":5,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","sample_list":"/paper/measuring-and-controlling-instruction-in#ran","syntology_url":"https://syntology.ai/paper/2402.10962","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.10962"}},"official":{"repos":["likenneth/persona_drift"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":3,"ran_from_kinds":["official"]}}},{"url":"/paper/prompt-optimization-in-multi-step-tasks","slug":"prompt-optimization-in-multi-step-tasks","title":"PRompt Optimization in Multi-Step Tasks (PROMST): Integrating Human Feedback and Heuristic-based Sampling","date":"2024-02-13","arxiv_id":"2402.08702","repositories_listed":1,"syntology":{"n":12,"n_ran":9,"n_constructed":0,"n_ran_checked":9,"n_instrument":0,"n_unverified":3,"n_honours":0,"n_violates":0,"n_no_contract":9,"n_pointer_only":0,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 9 with no instrument failure: 0 honoured, 0 violated, 9 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","sample_list":"/paper/prompt-optimization-in-multi-step-tasks#ran","syntology_url":"https://syntology.ai/paper/2402.08702","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.08702"}},"official":{"repos":["yongchao98/promst"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":9,"n_unverified":3,"ran_from_kinds":["official"]}}},{"url":"/paper/punctuation-restoration-improves-structure","slug":"punctuation-restoration-improves-structure","title":"Punctuation Restoration Improves Structure Understanding Without Supervision","date":"2024-02-13","arxiv_id":"2402.08382","repositories_listed":1,"syntology":null},{"url":"/paper/verified-multi-step-synthesis-using-large","slug":"verified-multi-step-synthesis-using-large","title":"VerMCTS: Synthesizing Multi-Step Programs using a Verifier, a Large Language Model, and Tree Search","date":"2024-02-13","arxiv_id":"2402.08147","repositories_listed":1,"syntology":{"n":8,"n_ran":8,"n_constructed":0,"n_ran_checked":8,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":8,"n_pointer_only":0,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 0 violated, 8 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/verified-multi-step-synthesis-using-large#ran","syntology_url":"https://syntology.ai/paper/2402.08147","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.08147"}},"official":{"repos":["namin/llm-verified-with-monte-carlo-tree-search"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/visually-dehallucinative-instruction","slug":"visually-dehallucinative-instruction","title":"Visually Dehallucinative Instruction Generation","date":"2024-02-13","arxiv_id":"2402.08348","repositories_listed":1,"syntology":null},{"url":"/paper/breakgpt-a-large-language-model-with-multi","slug":"breakgpt-a-large-language-model-with-multi","title":"BreakGPT: A Large Language Model with Multi-stage Structure for Financial Breakout Detection","date":"2024-02-12","arxiv_id":"2402.07536","repositories_listed":1,"syntology":null},{"url":"/paper/careless-whisper-speech-to-text-hallucination","slug":"careless-whisper-speech-to-text-hallucination","title":"Careless Whisper: Speech-to-Text Hallucination Harms","date":"2024-02-12","arxiv_id":"2402.08021","repositories_listed":1,"syntology":null},{"url":"/paper/detecting-the-clinical-features-of-difficult","slug":"detecting-the-clinical-features-of-difficult","title":"Detecting the Clinical Features of Difficult-to-Treat Depression using Synthetic Data from Large Language Models","date":"2024-02-12","arxiv_id":"2402.07645","repositories_listed":1,"syntology":null},{"url":"/paper/diffusion-of-thoughts-chain-of-thought","slug":"diffusion-of-thoughts-chain-of-thought","title":"Diffusion of Thoughts: Chain-of-Thought Reasoning in Diffusion Language Models","date":"2024-02-12","arxiv_id":"2402.07754","repositories_listed":1,"syntology":{"n":17,"n_ran":16,"n_constructed":0,"n_ran_checked":14,"n_instrument":2,"n_unverified":1,"n_honours":1,"n_violates":0,"n_no_contract":13,"n_pointer_only":17,"phrase":"16 ran (of which 0 constructed an object rather than computing a result; 14 with no instrument failure: 1 honoured, 0 violated, 13 with no contract checked; 2 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/diffusion-of-thoughts-chain-of-thought#ran","syntology_url":"https://syntology.ai/paper/2402.07754","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.07754"}},"official":{"repos":["hkunlp/diffusion-of-thoughts"],"state":"official (archive's flag): 16 ran","n_ran":16,"n_constructed":0,"n_ran_no_instrument_failure":14,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/wildfiregpt-tailored-large-language-model-for","slug":"wildfiregpt-tailored-large-language-model-for","title":"A RAG-Based Multi-Agent LLM System for Natural Hazard Resilience and Adaptation","date":"2024-02-12","arxiv_id":"2402.07877","repositories_listed":1,"syntology":null},{"url":"/paper/graphtranslator-aligning-graph-model-to-large","slug":"graphtranslator-aligning-graph-model-to-large","title":"GraphTranslator: Aligning Graph Model to Large Language Model for Open-ended Tasks","date":"2024-02-11","arxiv_id":"2402.07197","repositories_listed":1,"syntology":{"n":9,"n_ran":8,"n_constructed":0,"n_ran_checked":6,"n_instrument":2,"n_unverified":1,"n_honours":0,"n_violates":1,"n_no_contract":5,"n_pointer_only":1,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 1 violated, 5 with no contract checked; 2 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/graphtranslator-aligning-graph-model-to-large#ran","syntology_url":"https://syntology.ai/paper/2402.07197","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.07197"}},"official":{"repos":["alibaba/graphtranslator"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/hyperbert-mixing-hypergraph-aware-layers-with","slug":"hyperbert-mixing-hypergraph-aware-layers-with","title":"HyperBERT: Mixing Hypergraph-Aware Layers with Language Models for Node Classification on Text-Attributed Hypergraphs","date":"2024-02-11","arxiv_id":"2402.07309","repositories_listed":1,"syntology":null},{"url":"/paper/using-large-language-models-for-student-code","slug":"using-large-language-models-for-student-code","title":"Using Large Language Models for Student-Code Guided Test Case Generation in Computer Science Education","date":"2024-02-11","arxiv_id":"2402.07081","repositories_listed":1,"syntology":null},{"url":"/paper/chemllm-a-chemical-large-language-model","slug":"chemllm-a-chemical-large-language-model","title":"ChemLLM: A Chemical Large Language Model","date":"2024-02-10","arxiv_id":"2402.06852","repositories_listed":1,"syntology":null},{"url":"/paper/urbankgent-a-unified-large-language-model","slug":"urbankgent-a-unified-large-language-model","title":"UrbanKGent: A Unified Large Language Model Agent Framework for Urban Knowledge Graph Construction","date":"2024-02-10","arxiv_id":"2402.06861","repositories_listed":1,"syntology":{"n":5,"n_ran":3,"n_constructed":1,"n_ran_checked":1,"n_instrument":2,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":5,"phrase":"3 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 2 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/urbankgent-a-unified-large-language-model#ran","syntology_url":"https://syntology.ai/paper/2402.06861","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.06861"}},"official":{"repos":["usail-hkust/urbankgent"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/aya-dataset-an-open-access-collection-for","slug":"aya-dataset-an-open-access-collection-for","title":"Aya Dataset: An Open-Access Collection for Multilingual Instruction Tuning","date":"2024-02-09","arxiv_id":"2402.06619","repositories_listed":1,"syntology":null},{"url":"/paper/entropy-regularized-token-level-policy","slug":"entropy-regularized-token-level-policy","title":"Entropy-Regularized Token-Level Policy Optimization for Language Agent Reinforcement","date":"2024-02-09","arxiv_id":"2402.06700","repositories_listed":1,"syntology":{"n":9,"n_ran":6,"n_constructed":0,"n_ran_checked":3,"n_instrument":3,"n_unverified":3,"n_honours":0,"n_violates":1,"n_no_contract":2,"n_pointer_only":9,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 1 violated, 2 with no contract checked; 3 where Syntology's instrument failed) · 3 unverified","sample_list":"/paper/entropy-regularized-token-level-policy#ran","syntology_url":"https://syntology.ai/paper/2402.06700","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.06700"}},"official":{"repos":["morning9393/etpo"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":3,"ran_from_kinds":["official"]}}},{"url":"/paper/g-sciedbert-a-contextualized-llm-for-science","slug":"g-sciedbert-a-contextualized-llm-for-science","title":"G-SciEdBERT: A Contextualized LLM for Science Assessment Tasks in German","date":"2024-02-09","arxiv_id":"2402.06584","repositories_listed":1,"syntology":null},{"url":"/paper/language-model-sentence-completion-with-a","slug":"language-model-sentence-completion-with-a","title":"Language Model Sentence Completion with a Parser-Driven Rhetorical Control Method","date":"2024-02-09","arxiv_id":"2402.06125","repositories_listed":1,"syntology":{"n":5,"n_ran":3,"n_constructed":0,"n_ran_checked":1,"n_instrument":2,"n_unverified":2,"n_honours":1,"n_violates":0,"n_no_contract":0,"n_pointer_only":5,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/language-model-sentence-completion-with-a#ran","syntology_url":"https://syntology.ai/paper/2402.06125","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.06125"}},"official":{"repos":["joshua-zingale/plug-and-play-rst-ctg"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/model-editing-with-canonical-examples","slug":"model-editing-with-canonical-examples","title":"Model Editing with Canonical Examples","date":"2024-02-09","arxiv_id":"2402.06155","repositories_listed":1,"syntology":{"n":4,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":4,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/model-editing-with-canonical-examples#ran","syntology_url":"https://syntology.ai/paper/2402.06155","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.06155"}},"official":{"repos":["john-hewitt/model-editing-canonical-examples"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/on-the-efficacy-of-eviction-policy-for-key","slug":"on-the-efficacy-of-eviction-policy-for-key","title":"On the Efficacy of Eviction Policy for Key-Value Constrained Generative Language Model Inference","date":"2024-02-09","arxiv_id":"2402.06262","repositories_listed":1,"syntology":{"n":7,"n_ran":7,"n_constructed":0,"n_ran_checked":5,"n_instrument":2,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":5,"n_pointer_only":7,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/on-the-efficacy-of-eviction-policy-for-key#ran","syntology_url":"https://syntology.ai/paper/2402.06262","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.06262"}},"official":{"repos":["drsy/easykv"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/resumeflow-an-llm-facilitated-pipeline-for","slug":"resumeflow-an-llm-facilitated-pipeline-for","title":"ResumeFlow: An LLM-facilitated Pipeline for Personalized Resume Generation and Refinement","date":"2024-02-09","arxiv_id":"2402.06221","repositories_listed":1,"syntology":{"n":9,"n_ran":5,"n_constructed":0,"n_ran_checked":5,"n_instrument":0,"n_unverified":4,"n_honours":0,"n_violates":0,"n_no_contract":5,"n_pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","sample_list":"/paper/resumeflow-an-llm-facilitated-pipeline-for#ran","syntology_url":"https://syntology.ai/paper/2402.06221","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.06221"}},"official":{"repos":["Ztrimus/job-llm"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":4,"ran_from_kinds":["official"]}}},{"url":"/paper/screenagent-a-vision-language-model-driven","slug":"screenagent-a-vision-language-model-driven","title":"ScreenAgent: A Vision Language Model-driven Computer Control Agent","date":"2024-02-09","arxiv_id":"2402.07945","repositories_listed":1,"syntology":{"n":9,"n_ran":6,"n_constructed":0,"n_ran_checked":6,"n_instrument":0,"n_unverified":3,"n_honours":0,"n_violates":0,"n_no_contract":6,"n_pointer_only":9,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","sample_list":"/paper/screenagent-a-vision-language-model-driven#ran","syntology_url":"https://syntology.ai/paper/2402.07945","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.07945"}},"official":{"repos":["niuzaisheng/screenagent"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":3,"ran_from_kinds":["official"]}}},{"url":"/paper/the-quantified-boolean-bayesian-network","slug":"the-quantified-boolean-bayesian-network","title":"The Quantified Boolean Bayesian Network: Theory and Experiments with a Logical Graphical Model","date":"2024-02-09","arxiv_id":"2402.06557","repositories_listed":1,"syntology":null},{"url":"/paper/understanding-the-weakness-of-large-language","slug":"understanding-the-weakness-of-large-language","title":"Understanding the Weakness of Large Language Model Agents within a Complex Android Environment","date":"2024-02-09","arxiv_id":"2402.06596","repositories_listed":1,"syntology":{"n":4,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":4,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/understanding-the-weakness-of-large-language#ran","syntology_url":"https://syntology.ai/paper/2402.06596","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.06596"}},"official":{"repos":["androidarenaagent/androidarena"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/anfinsen-goes-neural-a-graphical-model-for","slug":"anfinsen-goes-neural-a-graphical-model-for","title":"Decoupled Sequence and Structure Generation for Realistic Antibody Design","date":"2024-02-08","arxiv_id":"2402.05982","repositories_listed":1,"syntology":null},{"url":"/paper/editable-scene-simulation-for-autonomous","slug":"editable-scene-simulation-for-autonomous","title":"Editable Scene Simulation for Autonomous Driving via Collaborative LLM-Agents","date":"2024-02-08","arxiv_id":"2402.05746","repositories_listed":1,"syntology":{"n":4,"n_ran":4,"n_constructed":0,"n_ran_checked":4,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":4,"n_pointer_only":4,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/editable-scene-simulation-for-autonomous#ran","syntology_url":"https://syntology.ai/paper/2402.05746","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.05746"}},"official":{"repos":["yifanlu0227/chatsim"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/sphinx-x-scaling-data-and-parameters-for-a","slug":"sphinx-x-scaling-data-and-parameters-for-a","title":"SPHINX-X: Scaling Data and Parameters for a Family of Multi-modal Large Language Models","date":"2024-02-08","arxiv_id":"2402.05935","repositories_listed":1,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":1,"n_instrument":2,"n_unverified":0,"n_honours":0,"n_violates":1,"n_no_contract":0,"n_pointer_only":3,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 1 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/sphinx-x-scaling-data-and-parameters-for-a#ran","syntology_url":"https://syntology.ai/paper/2402.05935","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.05935"}},"official":{"repos":["alpha-vllm/llama2-accessory"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/spirit-lm-interleaved-spoken-and-written","slug":"spirit-lm-interleaved-spoken-and-written","title":"Spirit LM: Interleaved Spoken and Written Language Model","date":"2024-02-08","arxiv_id":"2402.05755","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":0,"n_instrument":1,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/spirit-lm-interleaved-spoken-and-written#ran","syntology_url":"https://syntology.ai/paper/2402.05755","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.05755"}},"official":{"repos":["facebookresearch/spiritlm"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/apiq-finetuning-of-2-bit-quantized-large","slug":"apiq-finetuning-of-2-bit-quantized-large","title":"ApiQ: Finetuning of 2-Bit Quantized Large Language Model","date":"2024-02-07","arxiv_id":"2402.05147","repositories_listed":1,"syntology":{"n":2,"n_ran":1,"n_constructed":0,"n_ran_checked":1,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/apiq-finetuning-of-2-bit-quantized-large#ran","syntology_url":"https://syntology.ai/paper/2402.05147","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.05147"}},"official":{"repos":["baohaoliao/apiq"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/can-large-language-model-agents-simulate","slug":"can-large-language-model-agents-simulate","title":"Can Large Language Model Agents Simulate Human Trust Behavior?","date":"2024-02-07","arxiv_id":"2402.04559","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":0,"n_instrument":1,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/can-large-language-model-agents-simulate#ran","syntology_url":"https://syntology.ai/paper/2402.04559","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.04559"}},"official":{"repos":["camel-ai/agent-trust"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/codeit-self-improving-language-models-with","slug":"codeit-self-improving-language-models-with","title":"CodeIt: Self-Improving Language Models with Prioritized Hindsight Replay","date":"2024-02-07","arxiv_id":"2402.04858","repositories_listed":1,"syntology":{"n":41,"n_ran":32,"n_constructed":4,"n_ran_checked":10,"n_instrument":22,"n_unverified":9,"n_honours":4,"n_violates":2,"n_no_contract":4,"n_pointer_only":41,"phrase":"32 ran (of which 4 constructed an object rather than computing a result; 10 with no instrument failure: 4 honoured, 2 violated, 4 with no contract checked; 22 where Syntology's instrument failed) · 9 unverified","sample_list":"/paper/codeit-self-improving-language-models-with#ran","syntology_url":"https://syntology.ai/paper/2402.04858","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.04858"}},"official":{"repos":["Qualcomm-AI-research/codeit"],"state":"official (archive's flag): 32 ran","n_ran":32,"n_constructed":4,"n_ran_no_instrument_failure":10,"n_unverified":9,"ran_from_kinds":["official"]}}},{"url":"/paper/language-based-augmentation-to-address","slug":"language-based-augmentation-to-address","title":"Language-Based Augmentation to Address Shortcut Learning in Object Goal Navigation","date":"2024-02-07","arxiv_id":"2402.05090","repositories_listed":1,"syntology":null},{"url":"/paper/leveraging-llms-for-unsupervised-dense","slug":"leveraging-llms-for-unsupervised-dense","title":"Leveraging LLMs for Unsupervised Dense Retriever Ranking","date":"2024-02-07","arxiv_id":"2402.04853","repositories_listed":1,"syntology":null},{"url":"/paper/structure-informed-protein-language-model","slug":"structure-informed-protein-language-model","title":"Structure-Informed Protein Language Model","date":"2024-02-07","arxiv_id":"2402.05856","repositories_listed":1,"syntology":null},{"url":"/paper/sumrec-a-framework-for-recommendation-using","slug":"sumrec-a-framework-for-recommendation-using","title":"SumRec: A Framework for Recommendation using Open-Domain Dialogue","date":"2024-02-07","arxiv_id":"2402.04523","repositories_listed":1,"syntology":null},{"url":"/paper/2402-03766","slug":"2402-03766","title":"MobileVLM V2: Faster and Stronger Baseline for Vision Language Model","date":"2024-02-06","arxiv_id":"2402.03766","repositories_listed":1,"syntology":{"n":8,"n_ran":8,"n_constructed":0,"n_ran_checked":4,"n_instrument":4,"n_unverified":0,"n_honours":1,"n_violates":1,"n_no_contract":2,"n_pointer_only":4,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 1 honoured, 1 violated, 2 with no contract checked; 4 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/2402-03766#ran","syntology_url":"https://syntology.ai/paper/2402.03766","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.03766"}},"official":{"repos":["meituan-automl/mobilevlm"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/anytool-self-reflective-hierarchical-agents","slug":"anytool-self-reflective-hierarchical-agents","title":"AnyTool: Self-Reflective, Hierarchical Agents for Large-Scale API Calls","date":"2024-02-06","arxiv_id":"2402.04253","repositories_listed":1,"syntology":{"n":6,"n_ran":4,"n_constructed":0,"n_ran_checked":4,"n_instrument":0,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":4,"n_pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/anytool-self-reflective-hierarchical-agents#ran","syntology_url":"https://syntology.ai/paper/2402.04253","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.04253"}},"official":{"repos":["dyabel/anytool"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/identifying-reasons-for-contraceptive","slug":"identifying-reasons-for-contraceptive","title":"Identifying Reasons for Contraceptive Switching from Real-World Data Using Large Language Models","date":"2024-02-06","arxiv_id":"2402.03597","repositories_listed":1,"syntology":null},{"url":"/paper/personalized-language-modeling-from","slug":"personalized-language-modeling-from","title":"Personalized Language Modeling from Personalized Human Feedback","date":"2024-02-06","arxiv_id":"2402.05133","repositories_listed":1,"syntology":{"n":9,"n_ran":8,"n_constructed":0,"n_ran_checked":8,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":8,"n_pointer_only":0,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 0 violated, 8 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/personalized-language-modeling-from#ran","syntology_url":"https://syntology.ai/paper/2402.05133","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.05133"}},"official":{"repos":["humainlab/personalized_rlhf"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/positive-concave-deep-equilibrium-models","slug":"positive-concave-deep-equilibrium-models","title":"Positive concave deep equilibrium models","date":"2024-02-06","arxiv_id":"2402.04029","repositories_listed":1,"syntology":null},{"url":"/paper/retrieve-to-explain-evidence-driven","slug":"retrieve-to-explain-evidence-driven","title":"Retrieve to Explain: Evidence-driven Predictions with Language Models","date":"2024-02-06","arxiv_id":"2402.04068","repositories_listed":1,"syntology":null},{"url":"/paper/scientific-language-modeling-a-quantitative","slug":"scientific-language-modeling-a-quantitative","title":"A quantitative analysis of knowledge-learning preferences in large language models in molecular science","date":"2024-02-06","arxiv_id":"2402.04119","repositories_listed":1,"syntology":null},{"url":"/paper/sentiment-enhanced-graph-based-sarcasm","slug":"sentiment-enhanced-graph-based-sarcasm","title":"Sentiment-enhanced Graph-based Sarcasm Explanation in Dialogue","date":"2024-02-06","arxiv_id":"2402.03658","repositories_listed":1,"syntology":null},{"url":"/paper/arabic-synonym-bert-based-adversarial","slug":"arabic-synonym-bert-based-adversarial","title":"Arabic Synonym BERT-based Adversarial Examples for Text Classification","date":"2024-02-05","arxiv_id":"2402.03477","repositories_listed":1,"syntology":null},{"url":"/paper/large-language-model-distilling-medication","slug":"large-language-model-distilling-medication","title":"Large Language Model Distilling Medication Recommendation Model","date":"2024-02-05","arxiv_id":"2402.02803","repositories_listed":1,"syntology":null},{"url":"/paper/learning-from-teaching-regularization","slug":"learning-from-teaching-regularization","title":"Learning from Teaching Regularization: Generalizable Correlations Should be Easy to Imitate","date":"2024-02-05","arxiv_id":"2402.02769","repositories_listed":1,"syntology":null},{"url":"/paper/racer-an-llm-powered-methodology-for-scalable","slug":"racer-an-llm-powered-methodology-for-scalable","title":"RACER: An LLM-powered Methodology for Scalable Analysis of Semi-structured Mental Health Interviews","date":"2024-02-05","arxiv_id":"2402.02656","repositories_listed":1,"syntology":null},{"url":"/paper/rethinking-optimization-and-architecture-for","slug":"rethinking-optimization-and-architecture-for","title":"Rethinking Optimization and Architecture for Tiny Language Models","date":"2024-02-05","arxiv_id":"2402.02791","repositories_listed":1,"syntology":{"n":2,"n_ran":1,"n_constructed":0,"n_ran_checked":1,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":2,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/rethinking-optimization-and-architecture-for#ran","syntology_url":"https://syntology.ai/paper/2402.02791","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.02791"}},"official":{"repos":["yuchuantian/rethinktinylm"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/skill-set-optimization-reinforcing-language","slug":"skill-set-optimization-reinforcing-language","title":"Skill Set Optimization: Reinforcing Language Model Behavior via Transferable Skills","date":"2024-02-05","arxiv_id":"2402.03244","repositories_listed":1,"syntology":{"n":3,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/skill-set-optimization-reinforcing-language#ran","syntology_url":"https://syntology.ai/paper/2402.03244","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.03244"}},"official":{"repos":["allenai/sso"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/texshape-information-theoretic-sentence","slug":"texshape-information-theoretic-sentence","title":"TexShape: Information Theoretic Sentence Embedding for Language Models","date":"2024-02-05","arxiv_id":"2402.05132","repositories_listed":1,"syntology":null},{"url":"/paper/unsupervised-semantic-segmentation-of-high","slug":"unsupervised-semantic-segmentation-of-high","title":"Applying Unsupervised Semantic Segmentation to High-Resolution UAV Imagery for Enhanced Road Scene Parsing","date":"2024-02-05","arxiv_id":"2402.02985","repositories_listed":1,"syntology":null},{"url":"/paper/autotimes-autoregressive-time-series","slug":"autotimes-autoregressive-time-series","title":"AutoTimes: Autoregressive Time Series Forecasters via Large Language Models","date":"2024-02-04","arxiv_id":"2402.02370","repositories_listed":1,"syntology":null},{"url":"/paper/can-large-language-models-learn-independent","slug":"can-large-language-models-learn-independent","title":"Can Large Language Models Learn Independent Causal Mechanisms?","date":"2024-02-04","arxiv_id":"2402.02636","repositories_listed":1,"syntology":{"n":8,"n_ran":6,"n_constructed":0,"n_ran_checked":6,"n_instrument":0,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":6,"n_pointer_only":1,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/can-large-language-models-learn-independent#ran","syntology_url":"https://syntology.ai/paper/2402.02636","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.02636"}},"official":{"repos":["strong-ai-lab/modular-lm"],"state":"official: no sample here; runs from other or unrecorded repositories","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["found_in_text"]}}},{"url":"/paper/gerea-question-aware-prompt-captions-for","slug":"gerea-question-aware-prompt-captions-for","title":"GeReA: Question-Aware Prompt Captions for Knowledge-based Visual Question Answering","date":"2024-02-04","arxiv_id":"2402.02503","repositories_listed":1,"syntology":{"n":18,"n_ran":13,"n_constructed":0,"n_ran_checked":13,"n_instrument":0,"n_unverified":5,"n_honours":0,"n_violates":0,"n_no_contract":13,"n_pointer_only":18,"phrase":"13 ran (of which 0 constructed an object rather than computing a result; 13 with no instrument failure: 0 honoured, 0 violated, 13 with no contract checked; 0 where Syntology's instrument failed) · 5 unverified","sample_list":"/paper/gerea-question-aware-prompt-captions-for#ran","syntology_url":"https://syntology.ai/paper/2402.02503","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.02503"}},"official":{"repos":["upper9527/gerea"],"state":"official (archive's flag): 13 ran","n_ran":13,"n_constructed":0,"n_ran_no_instrument_failure":13,"n_unverified":5,"ran_from_kinds":["official"]}}},{"url":"/paper/girt-model-automated-generation-of-issue","slug":"girt-model-automated-generation-of-issue","title":"GIRT-Model: Automated Generation of Issue Report Templates","date":"2024-02-04","arxiv_id":"2402.02632","repositories_listed":1,"syntology":null},{"url":"/paper/glape-gold-label-agnostic-prompt-evaluation","slug":"glape-gold-label-agnostic-prompt-evaluation","title":"GLaPE: Gold Label-agnostic Prompt Evaluation and Optimization for Large Language Model","date":"2024-02-04","arxiv_id":"2402.02408","repositories_listed":1,"syntology":null},{"url":"/paper/kicgpt-large-language-model-with-knowledge-in","slug":"kicgpt-large-language-model-with-knowledge-in","title":"KICGPT: Large Language Model with Knowledge in Context for Knowledge Graph Completion","date":"2024-02-04","arxiv_id":"2402.02389","repositories_listed":1,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":1,"n_instrument":2,"n_unverified":0,"n_honours":1,"n_violates":0,"n_no_contract":0,"n_pointer_only":1,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/kicgpt-large-language-model-with-knowledge-in#ran","syntology_url":"https://syntology.ai/paper/2402.02389","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.02389"}},"official":{"repos":["weiyanbin1999/kicgpt"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official","unlocated"]}}},{"url":"/paper/lhrs-bot-empowering-remote-sensing-with-vgi","slug":"lhrs-bot-empowering-remote-sensing-with-vgi","title":"LHRS-Bot: Empowering Remote Sensing with VGI-Enhanced Large Multimodal Language Model","date":"2024-02-04","arxiv_id":"2402.02544","repositories_listed":1,"syntology":{"n":5,"n_ran":5,"n_constructed":0,"n_ran_checked":2,"n_instrument":3,"n_unverified":0,"n_honours":2,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 2 honoured, 0 violated, 0 with no contract checked; 3 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/lhrs-bot-empowering-remote-sensing-with-vgi#ran","syntology_url":"https://syntology.ai/paper/2402.02544","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.02544"}},"official":{"repos":["NJU-LHRS/LHRS-Bot"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/selecting-large-language-model-to-fine-tune","slug":"selecting-large-language-model-to-fine-tune","title":"Selecting Large Language Model to Fine-tune via Rectified Scaling Law","date":"2024-02-04","arxiv_id":"2402.02314","repositories_listed":1,"syntology":{"n":6,"n_ran":6,"n_constructed":0,"n_ran_checked":2,"n_instrument":4,"n_unverified":0,"n_honours":0,"n_violates":2,"n_no_contract":0,"n_pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 2 violated, 0 with no contract checked; 4 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/selecting-large-language-model-to-fine-tune#ran","syntology_url":"https://syntology.ai/paper/2402.02314","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.02314"}},"official":null}},{"url":"/paper/anthroscore-a-computational-linguistic","slug":"anthroscore-a-computational-linguistic","title":"AnthroScore: A Computational Linguistic Measure of Anthropomorphism","date":"2024-02-03","arxiv_id":"2402.02056","repositories_listed":1,"syntology":null},{"url":"/paper/frequency-explains-the-inverse-correlation-of","slug":"frequency-explains-the-inverse-correlation-of","title":"Frequency Explains the Inverse Correlation of Large Language Models' Size, Training Data Amount, and Surprisal's Fit to Reading Times","date":"2024-02-03","arxiv_id":"2402.02255","repositories_listed":1,"syntology":null},{"url":"/paper/apiserve-efficient-api-support-for-large","slug":"apiserve-efficient-api-support-for-large","title":"InferCept: Efficient Intercept Support for Augmented Large Language Model Inference","date":"2024-02-02","arxiv_id":"2402.01869","repositories_listed":1,"syntology":{"n":4,"n_ran":3,"n_constructed":0,"n_ran_checked":3,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":3,"n_pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/apiserve-efficient-api-support-for-large#ran","syntology_url":"https://syntology.ai/paper/2402.01869","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.01869"}},"official":{"repos":["wuklab/infercept"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/audio-flamingo-a-novel-audio-language-model","slug":"audio-flamingo-a-novel-audio-language-model","title":"Audio Flamingo: A Novel Audio Language Model with Few-Shot Learning and Dialogue Abilities","date":"2024-02-02","arxiv_id":"2402.01831","repositories_listed":1,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":3,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":3,"n_no_contract":0,"n_pointer_only":3,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 3 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/audio-flamingo-a-novel-audio-language-model#ran","syntology_url":"https://syntology.ai/paper/2402.01831","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.01831"}},"official":{"repos":["NVIDIA/audio-flamingo"],"state":"official: no sample here; runs from other or unrecorded repositories","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["unlocated"]}}},{"url":"/paper/decoding-speculative-decoding","slug":"decoding-speculative-decoding","title":"Decoding Speculative Decoding","date":"2024-02-02","arxiv_id":"2402.01528","repositories_listed":1,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":0,"n_instrument":3,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":3,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 3 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/decoding-speculative-decoding#ran","syntology_url":"https://syntology.ai/paper/2402.01528","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.01528"}},"official":{"repos":["uw-mad-dash/decoding-speculative-decoding"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/interpretation-of-intracardiac-electrograms","slug":"interpretation-of-intracardiac-electrograms","title":"Interpretation of Intracardiac Electrograms Through Textual Representations","date":"2024-02-02","arxiv_id":"2402.01115","repositories_listed":1,"syntology":null},{"url":"/paper/magdi-structured-distillation-of-multi-agent","slug":"magdi-structured-distillation-of-multi-agent","title":"MAGDi: Structured Distillation of Multi-Agent Interaction Graphs Improves Reasoning in Smaller Language Models","date":"2024-02-02","arxiv_id":"2402.01620","repositories_listed":1,"syntology":{"n":6,"n_ran":6,"n_constructed":0,"n_ran_checked":6,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":6,"n_pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/magdi-structured-distillation-of-multi-agent#ran","syntology_url":"https://syntology.ai/paper/2402.01620","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.01620"}},"official":{"repos":["dinobby/magdi"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/rethinking-the-role-of-proxy-rewards-in","slug":"rethinking-the-role-of-proxy-rewards-in","title":"Rethinking the Role of Proxy Rewards in Language Model Alignment","date":"2024-02-02","arxiv_id":"2402.03469","repositories_listed":1,"syntology":null},{"url":"/paper/style-vectors-for-steering-generative-large","slug":"style-vectors-for-steering-generative-large","title":"Style Vectors for Steering Generative Large Language Model","date":"2024-02-02","arxiv_id":"2402.01618","repositories_listed":1,"syntology":{"n":12,"n_ran":8,"n_constructed":0,"n_ran_checked":7,"n_instrument":1,"n_unverified":4,"n_honours":0,"n_violates":0,"n_no_contract":7,"n_pointer_only":1,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 1 where Syntology's instrument failed) · 4 unverified","sample_list":"/paper/style-vectors-for-steering-generative-large#ran","syntology_url":"https://syntology.ai/paper/2402.01618","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.01618"}},"official":{"repos":["dlr-sc/style-vectors-for-steering-llms"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":7,"n_unverified":4,"ran_from_kinds":["official"]}}},{"url":"/paper/vaccine-perturbation-aware-alignment-for","slug":"vaccine-perturbation-aware-alignment-for","title":"Vaccine: Perturbation-aware Alignment for Large Language Models against Harmful Fine-tuning Attack","date":"2024-02-02","arxiv_id":"2402.01109","repositories_listed":1,"syntology":null},{"url":"/paper/blackmamba-mixture-of-experts-for-state-space","slug":"blackmamba-mixture-of-experts-for-state-space","title":"BlackMamba: Mixture of Experts for State-Space Models","date":"2024-02-01","arxiv_id":"2402.01771","repositories_listed":1,"syntology":null}],"record_sha256":"6461181a6d1dd48c85ab375a5cf3b5d4e69e7beae9c3e3b9512f749ec353d490","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}