{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/task/text-generation/papers/8","list_of":"/task/text-generation","task":"Text Generation","archive":{"snapshot":"2025-07-28"},"key_notes":{"n_ran_checked":"legacy name, kept unchanged so existing readers do not break: it counts the samples that ran with no instrument failure (honoured, violated, and ran with no contract checked); it does not mean a contract was checked, and the pages print it as 'K with no instrument failure', not 'K checked'","n_constructed":"a sub-count of the samples that ran, never subtracted from them and never a failure: an executed sample whose run returned an instance of its own class (fixture_out_type equals the entry name): the run built an object and did not compute a result (Syntology's RAN record, counts.constructed)"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"repositories listed in the archive (most first), then date (newest first), then slug","page":8,"pages_in_order":54,"rows_per_page":100,"rows":[701,800],"of":5335,"counts":{"archive_papers_tagged":5335,"with_a_code_link":2047,"where_syntology_ran_a_sample":610,"not_listed_spam_title":0,"listed":5335,"listed_where_code_ran":610,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":503,"every_run_a_failure_of_syntologys_instrument":107,"listed_with_a_run_with_no_instrument_failure":503,"listed_every_run_a_failure_of_syntologys_instrument":107,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/task/text-generation","prev":"/task/text-generation/papers/7","next":"/task/text-generation/papers/9","papers":[{"url":"/paper/lamda-large-model-fine-tuning-via-spectrally","slug":"lamda-large-model-fine-tuning-via-spectrally","title":"LaMDA: Large Model Fine-Tuning via Spectrally Decomposed Low-Dimensional Adaptation","date":"2024-06-18","arxiv_id":"2406.12832","repositories_listed":1,"syntology":{"n":1,"n_ran":0,"n_constructed":0,"n_ran_checked":0,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"0 ran · 1 unverified","sample_list":"/paper/lamda-large-model-fine-tuning-via-spectrally#ran","syntology_url":"https://syntology.ai/paper/2406.12832","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2406.12832"}},"official":{"repos":["arminazizi98/lamda"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":[]}}},{"url":"/paper/shield-evaluation-and-defense-strategies-for","slug":"shield-evaluation-and-defense-strategies-for","title":"SHIELD: Evaluation and Defense Strategies for Copyright Compliance in LLM Text Generation","date":"2024-06-18","arxiv_id":"2406.12975","repositories_listed":1,"syntology":{"n":2,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/shield-evaluation-and-defense-strategies-for#ran","syntology_url":"https://syntology.ai/paper/2406.12975","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2406.12975"}},"official":{"repos":["xz-liu/shield"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/dtgb-a-comprehensive-benchmark-for-dynamic","slug":"dtgb-a-comprehensive-benchmark-for-dynamic","title":"DTGB: A Comprehensive Benchmark for Dynamic Text-Attributed Graphs","date":"2024-06-17","arxiv_id":"2406.12072","repositories_listed":1,"syntology":{"n":1,"n_ran":0,"n_constructed":0,"n_ran_checked":0,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":1,"phrase":"0 ran · 1 unverified","sample_list":"/paper/dtgb-a-comprehensive-benchmark-for-dynamic#ran","syntology_url":"https://syntology.ai/paper/2406.12072","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2406.12072"}},"official":{"repos":["zjs123/DTGB"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":[]}}},{"url":"/paper/extrinsic-evaluation-of-cultural-competence","slug":"extrinsic-evaluation-of-cultural-competence","title":"Extrinsic Evaluation of Cultural Competence in Large Language Models","date":"2024-06-17","arxiv_id":"2406.11565","repositories_listed":1,"syntology":null},{"url":"/paper/incentivizing-quality-text-generation-via","slug":"incentivizing-quality-text-generation-via","title":"Incentivizing Quality Text Generation via Statistical Contracts","date":"2024-06-17","arxiv_id":"2406.11118","repositories_listed":1,"syntology":{"n":1,"n_ran":0,"n_constructed":0,"n_ran_checked":0,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":1,"phrase":"0 ran · 1 unverified","sample_list":"/paper/incentivizing-quality-text-generation-via#ran","syntology_url":"https://syntology.ai/paper/2406.11118","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2406.11118"}},"official":{"repos":["edensaig/llm-contracts"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":[]}}},{"url":"/paper/analyzing-key-neurons-in-large-language","slug":"analyzing-key-neurons-in-large-language","title":"Identifying Query-Relevant Neurons in Large Language Models for Long-Form Texts","date":"2024-06-16","arxiv_id":"2406.10868","repositories_listed":1,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":2,"n_instrument":1,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":1,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/analyzing-key-neurons-in-large-language#ran","syntology_url":"https://syntology.ai/paper/2406.10868","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2406.10868"}},"official":{"repos":["tigerchen52/qrneuron"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/post-hoc-utterance-refining-method-by-entity","slug":"post-hoc-utterance-refining-method-by-entity","title":"Post-hoc Utterance Refining Method by Entity Mining for Faithful Knowledge Grounded Conversations","date":"2024-06-16","arxiv_id":"2406.10809","repositories_listed":1,"syntology":null},{"url":"/paper/comm-a-coherent-interleaved-image-text","slug":"comm-a-coherent-interleaved-image-text","title":"CoMM: A Coherent Interleaved Image-Text Dataset for Multimodal Understanding and Generation","date":"2024-06-15","arxiv_id":"2406.10462","repositories_listed":1,"syntology":null},{"url":"/paper/freectrl-constructing-control-centers-with","slug":"freectrl-constructing-control-centers-with","title":"FreeCtrl: Constructing Control Centers with Feedforward Layers for Learning-Free Controllable Text Generation","date":"2024-06-14","arxiv_id":"2406.09688","repositories_listed":1,"syntology":{"n":1,"n_ran":0,"n_constructed":0,"n_ran_checked":0,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":1,"phrase":"0 ran · 1 unverified","sample_list":"/paper/freectrl-constructing-control-centers-with#ran","syntology_url":"https://syntology.ai/paper/2406.09688","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2406.09688"}},"official":{"repos":["zijian678/FreeCtrl"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":[]}}},{"url":"/paper/mmrel-a-relation-understanding-dataset-and","slug":"mmrel-a-relation-understanding-dataset-and","title":"MMRel: A Relation Understanding Benchmark in the MLLM Era","date":"2024-06-13","arxiv_id":"2406.09121","repositories_listed":1,"syntology":null},{"url":"/paper/analyzing-constrained-llm-through-pdfa","slug":"analyzing-constrained-llm-through-pdfa","title":"Analyzing constrained LLM through PDFA-learning","date":"2024-06-12","arxiv_id":"2406.08269","repositories_listed":1,"syntology":null},{"url":"/paper/conme-rethinking-evaluation-of-compositional","slug":"conme-rethinking-evaluation-of-compositional","title":"ConMe: Rethinking Evaluation of Compositional Reasoning for Modern VLMs","date":"2024-06-12","arxiv_id":"2406.08164","repositories_listed":1,"syntology":{"n":12,"n_ran":11,"n_constructed":0,"n_ran_checked":11,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":11,"n_pointer_only":12,"phrase":"11 ran (of which 0 constructed an object rather than computing a result; 11 with no instrument failure: 0 honoured, 0 violated, 11 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/conme-rethinking-evaluation-of-compositional#ran","syntology_url":"https://syntology.ai/paper/2406.08164","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2406.08164"}},"official":{"repos":["jmiemirza/conme"],"state":"official (archive's flag): 11 ran","n_ran":11,"n_constructed":0,"n_ran_no_instrument_failure":11,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/defining-and-detecting-vulnerability-in-human","slug":"defining-and-detecting-vulnerability-in-human","title":"Defining and Detecting Vulnerability in Human Evaluation Guidelines: A Preliminary Study Towards Reliable NLG Evaluation","date":"2024-06-12","arxiv_id":"2406.07935","repositories_listed":1,"syntology":null},{"url":"/paper/ceret-cost-effective-extrinsic-refinement-for","slug":"ceret-cost-effective-extrinsic-refinement-for","title":"CERET: Cost-Effective Extrinsic Refinement for Text Generation","date":"2024-06-08","arxiv_id":"2406.05588","repositories_listed":1,"syntology":null},{"url":"/paper/annotating-framenet-via-structure-conditioned","slug":"annotating-framenet-via-structure-conditioned","title":"Annotating FrameNet via Structure-Conditioned Language Generation","date":"2024-06-07","arxiv_id":"2406.04834","repositories_listed":1,"syntology":null},{"url":"/paper/extroversion-or-introversion-controlling-the","slug":"extroversion-or-introversion-controlling-the","title":"Extroversion or Introversion? Controlling The Personality of Your Large Language Models","date":"2024-06-07","arxiv_id":"2406.04583","repositories_listed":1,"syntology":null},{"url":"/paper/improving-logits-based-detector-without","slug":"improving-logits-based-detector-without","title":"DALD: Improving Logits-based Detector without Logits from Black-box LLMs","date":"2024-06-07","arxiv_id":"2406.05232","repositories_listed":1,"syntology":null},{"url":"/paper/on-subjective-uncertainty-quantification-and","slug":"on-subjective-uncertainty-quantification-and","title":"On Subjective Uncertainty Quantification and Calibration in Natural Language Generation","date":"2024-06-07","arxiv_id":"2406.05213","repositories_listed":1,"syntology":{"n":14,"n_ran":12,"n_constructed":0,"n_ran_checked":12,"n_instrument":0,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":12,"n_pointer_only":0,"phrase":"12 ran (of which 0 constructed an object rather than computing a result; 12 with no instrument failure: 0 honoured, 0 violated, 12 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/on-subjective-uncertainty-quantification-and#ran","syntology_url":"https://syntology.ai/paper/2406.05213","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2406.05213"}},"official":{"repos":["meta-inf/suq-nlg"],"state":"official (archive's flag): 12 ran","n_ran":12,"n_constructed":0,"n_ran_no_instrument_failure":12,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/evaluating-the-smooth-control-of-attribute","slug":"evaluating-the-smooth-control-of-attribute","title":"Evaluating the Smooth Control of Attribute Intensity in Text Generation with LLMs","date":"2024-06-06","arxiv_id":"2406.04460","repositories_listed":1,"syntology":null},{"url":"/paper/maira-2-grounded-radiology-report-generation","slug":"maira-2-grounded-radiology-report-generation","title":"MAIRA-2: Grounded Radiology Report Generation","date":"2024-06-06","arxiv_id":"2406.04449","repositories_listed":1,"syntology":{"n":4,"n_ran":4,"n_constructed":0,"n_ran_checked":4,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":4,"n_pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/maira-2-grounded-radiology-report-generation#ran","syntology_url":"https://syntology.ai/paper/2406.04449","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2406.04449"}},"official":{"repos":["microsoft/RadFact"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/semantically-diverse-language-generation-for","slug":"semantically-diverse-language-generation-for","title":"Semantically Diverse Language Generation for Uncertainty Estimation in Language Models","date":"2024-06-06","arxiv_id":"2406.04306","repositories_listed":1,"syntology":{"n":11,"n_ran":5,"n_constructed":0,"n_ran_checked":0,"n_instrument":5,"n_unverified":6,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":11,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 5 where Syntology's instrument failed) · 6 unverified","sample_list":"/paper/semantically-diverse-language-generation-for#ran","syntology_url":"https://syntology.ai/paper/2406.04306","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2406.04306"}},"official":{"repos":["ml-jku/SDLG"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":6,"ran_from_kinds":["official"]}}},{"url":"/paper/uncovering-limitations-of-large-language","slug":"uncovering-limitations-of-large-language","title":"Uncovering Limitations of Large Language Models in Information Seeking from Tables","date":"2024-06-06","arxiv_id":"2406.04113","repositories_listed":1,"syntology":null},{"url":"/paper/css-contrastive-semantic-similarity-for","slug":"css-contrastive-semantic-similarity-for","title":"CSS: Contrastive Semantic Similarity for Uncertainty Quantification of LLMs","date":"2024-06-05","arxiv_id":"2406.03158","repositories_listed":1,"syntology":{"n":11,"n_ran":9,"n_constructed":0,"n_ran_checked":8,"n_instrument":1,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":8,"n_pointer_only":11,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 0 violated, 8 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/css-contrastive-semantic-similarity-for#ran","syntology_url":"https://syntology.ai/paper/2406.03158","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2406.03158"}},"official":{"repos":["aoshuang92/css_uq_llms"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/drivlme-enhancing-llm-based-autonomous","slug":"drivlme-enhancing-llm-based-autonomous","title":"DriVLMe: Enhancing LLM-based Autonomous Driving Agents with Embodied and Social Experiences","date":"2024-06-05","arxiv_id":"2406.03008","repositories_listed":1,"syntology":{"n":2,"n_ran":2,"n_constructed":0,"n_ran_checked":0,"n_instrument":2,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/drivlme-enhancing-llm-based-autonomous#ran","syntology_url":"https://syntology.ai/paper/2406.03008","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2406.03008"}},"official":{"repos":["sled-group/driVLMe"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/efficient-minimum-bayes-risk-decoding-using","slug":"efficient-minimum-bayes-risk-decoding-using","title":"Efficient Minimum Bayes Risk Decoding using Low-Rank Matrix Completion Algorithms","date":"2024-06-05","arxiv_id":"2406.02832","repositories_listed":1,"syntology":null},{"url":"/paper/fusionbench-a-comprehensive-benchmark-of-deep","slug":"fusionbench-a-comprehensive-benchmark-of-deep","title":"FusionBench: A Comprehensive Benchmark of Deep Model Fusion","date":"2024-06-05","arxiv_id":"2406.03280","repositories_listed":1,"syntology":{"n":4,"n_ran":4,"n_constructed":0,"n_ran_checked":4,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":4,"n_pointer_only":4,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/fusionbench-a-comprehensive-benchmark-of-deep#ran","syntology_url":"https://syntology.ai/paper/2406.03280","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2406.03280"}},"official":{"repos":["tanganke/fusion_bench"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/open-grounded-planning-challenges-and","slug":"open-grounded-planning-challenges-and","title":"Open Grounded Planning: Challenges and Benchmark Construction","date":"2024-06-05","arxiv_id":"2406.02903","repositories_listed":1,"syntology":null},{"url":"/paper/patenteval-understanding-errors-in-patent","slug":"patenteval-understanding-errors-in-patent","title":"PatentEval: Understanding Errors in Patent Generation","date":"2024-06-05","arxiv_id":"2406.06589","repositories_listed":1,"syntology":{"n":5,"n_ran":3,"n_constructed":0,"n_ran_checked":3,"n_instrument":0,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":3,"n_pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/patenteval-understanding-errors-in-patent#ran","syntology_url":"https://syntology.ai/paper/2406.06589","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2406.06589"}},"official":{"repos":["zoeyou/patenteval"],"state":"official: no sample here; runs from other or unrecorded repositories","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["found_in_text"]}}},{"url":"/paper/fedmkt-federated-mutual-knowledge-transfer","slug":"fedmkt-federated-mutual-knowledge-transfer","title":"FedMKT: Federated Mutual Knowledge Transfer for Large and Small Language Models","date":"2024-06-04","arxiv_id":"2406.02224","repositories_listed":1,"syntology":null},{"url":"/paper/set-based-prompting-provably-solving-the","slug":"set-based-prompting-provably-solving-the","title":"Order-Independence Without Fine Tuning","date":"2024-06-04","arxiv_id":"2406.06581","repositories_listed":1,"syntology":{"n":15,"n_ran":8,"n_constructed":0,"n_ran_checked":1,"n_instrument":7,"n_unverified":7,"n_honours":1,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 7 where Syntology's instrument failed) · 7 unverified","sample_list":"/paper/set-based-prompting-provably-solving-the#ran","syntology_url":"https://syntology.ai/paper/2406.06581","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2406.06581"}},"official":{"repos":["reidmcy/set-based-prompting"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":7,"ran_from_kinds":["official"]}}},{"url":"/paper/specexec-massively-parallel-speculative","slug":"specexec-massively-parallel-speculative","title":"SpecExec: Massively Parallel Speculative Decoding for Interactive LLM Inference on Consumer Devices","date":"2024-06-04","arxiv_id":"2406.02532","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":1,"n_instrument":0,"n_unverified":0,"n_honours":1,"n_violates":0,"n_no_contract":0,"n_pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/specexec-massively-parallel-speculative#ran","syntology_url":"https://syntology.ai/paper/2406.02532","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2406.02532"}},"official":{"repos":["yandex-research/specexec"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/contextualized-sequence-likelihood-enhanced","slug":"contextualized-sequence-likelihood-enhanced","title":"Contextualized Sequence Likelihood: Enhanced Confidence Scores for Natural Language Generation","date":"2024-06-03","arxiv_id":"2406.01806","repositories_listed":1,"syntology":{"n":14,"n_ran":8,"n_constructed":0,"n_ran_checked":7,"n_instrument":1,"n_unverified":6,"n_honours":0,"n_violates":0,"n_no_contract":7,"n_pointer_only":2,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 1 where Syntology's instrument failed) · 6 unverified","sample_list":"/paper/contextualized-sequence-likelihood-enhanced#ran","syntology_url":"https://syntology.ai/paper/2406.01806","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2406.01806"}},"official":{"repos":["zlin7/contextsl"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":7,"n_unverified":6,"ran_from_kinds":["official"]}}},{"url":"/paper/mad-multi-alignment-meg-to-text-decoding","slug":"mad-multi-alignment-meg-to-text-decoding","title":"MAD: Multi-Alignment MEG-to-Text Decoding","date":"2024-06-03","arxiv_id":"2406.01512","repositories_listed":1,"syntology":{"n":13,"n_ran":12,"n_constructed":0,"n_ran_checked":12,"n_instrument":0,"n_unverified":1,"n_honours":1,"n_violates":2,"n_no_contract":9,"n_pointer_only":0,"phrase":"12 ran (of which 0 constructed an object rather than computing a result; 12 with no instrument failure: 1 honoured, 2 violated, 9 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/mad-multi-alignment-meg-to-text-decoding#ran","syntology_url":"https://syntology.ai/paper/2406.01512","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2406.01512"}},"official":{"repos":["neuspeech/mad-meg2text"],"state":"official (archive's flag): 12 ran","n_ran":12,"n_constructed":0,"n_ran_no_instrument_failure":12,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/tcmbench-a-comprehensive-benchmark-for","slug":"tcmbench-a-comprehensive-benchmark-for","title":"TCMBench: A Comprehensive Benchmark for Evaluating Large Language Models in Traditional Chinese Medicine","date":"2024-06-03","arxiv_id":"2406.01126","repositories_listed":1,"syntology":null},{"url":"/paper/the-power-of-summary-source-alignments","slug":"the-power-of-summary-source-alignments","title":"The Power of Summary-Source Alignments","date":"2024-06-02","arxiv_id":"2406.00842","repositories_listed":1,"syntology":null},{"url":"/paper/evaluating-large-language-model-biases-in","slug":"evaluating-large-language-model-biases-in","title":"Evaluating Large Language Model Biases in Persona-Steered Generation","date":"2024-05-30","arxiv_id":"2405.20253","repositories_listed":1,"syntology":{"n":10,"n_ran":8,"n_constructed":0,"n_ran_checked":8,"n_instrument":0,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":8,"n_pointer_only":10,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 0 violated, 8 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/evaluating-large-language-model-biases-in#ran","syntology_url":"https://syntology.ai/paper/2405.20253","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2405.20253"}},"official":{"repos":["andyjliu/persona-steered-generation-bias"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/multi-aspect-controllable-text-generation","slug":"multi-aspect-controllable-text-generation","title":"Multi-Aspect Controllable Text Generation with Disentangled Counterfactual Augmentation","date":"2024-05-30","arxiv_id":"2405.19958","repositories_listed":1,"syntology":{"n":4,"n_ran":4,"n_constructed":3,"n_ran_checked":3,"n_instrument":1,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":3,"n_pointer_only":4,"phrase":"4 ran (of which 3 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/multi-aspect-controllable-text-generation#ran","syntology_url":"https://syntology.ai/paper/2405.19958","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2405.19958"}},"official":{"repos":["nju-websoft/magic"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":3,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/rtgen-generating-region-text-pairs-for-open","slug":"rtgen-generating-region-text-pairs-for-open","title":"RTGen: Generating Region-Text Pairs for Open-Vocabulary Object Detection","date":"2024-05-30","arxiv_id":"2405.19854","repositories_listed":1,"syntology":null},{"url":"/paper/language-generation-with-strictly-proper","slug":"language-generation-with-strictly-proper","title":"Language Generation with Strictly Proper Scoring Rules","date":"2024-05-29","arxiv_id":"2405.18906","repositories_listed":1,"syntology":{"n":6,"n_ran":5,"n_constructed":0,"n_ran_checked":3,"n_instrument":2,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":3,"n_pointer_only":2,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 2 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/language-generation-with-strictly-proper#ran","syntology_url":"https://syntology.ai/paper/2405.18906","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2405.18906"}},"official":{"repos":["shaochenze/scoringruleslm"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":1,"ran_from_kinds":["official","unlocated"]}}},{"url":"/paper/wrdscore-new-metric-for-evaluation-of-natural","slug":"wrdscore-new-metric-for-evaluation-of-natural","title":"WRDScore: New Metric for Evaluation of Natural Language Generation Models","date":"2024-05-29","arxiv_id":"2405.19220","repositories_listed":1,"syntology":null},{"url":"/paper/hardware-aware-parallel-prompt-decoding-for","slug":"hardware-aware-parallel-prompt-decoding-for","title":"Hardware-Aware Parallel Prompt Decoding for Memory-Efficient Acceleration of LLM Inference","date":"2024-05-28","arxiv_id":"2405.18628","repositories_listed":1,"syntology":null},{"url":"/paper/on-the-noise-robustness-of-in-context","slug":"on-the-noise-robustness-of-in-context","title":"On the Noise Robustness of In-Context Learning for Text Generation","date":"2024-05-27","arxiv_id":"2405.17264","repositories_listed":1,"syntology":{"n":11,"n_ran":5,"n_constructed":0,"n_ran_checked":5,"n_instrument":0,"n_unverified":6,"n_honours":0,"n_violates":0,"n_no_contract":5,"n_pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 6 unverified","sample_list":"/paper/on-the-noise-robustness-of-in-context#ran","syntology_url":"https://syntology.ai/paper/2405.17264","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2405.17264"}},"official":{"repos":["ml-stat-sustech/local-perplexity-ranking"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":6,"ran_from_kinds":["official"]}}},{"url":"/paper/automatic-jailbreaking-of-the-text-to-image","slug":"automatic-jailbreaking-of-the-text-to-image","title":"Automatic Jailbreaking of the Text-to-Image Generative AI Systems","date":"2024-05-26","arxiv_id":"2405.16567","repositories_listed":1,"syntology":{"n":10,"n_ran":8,"n_constructed":0,"n_ran_checked":7,"n_instrument":1,"n_unverified":2,"n_honours":1,"n_violates":0,"n_no_contract":6,"n_pointer_only":10,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 1 honoured, 0 violated, 6 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/automatic-jailbreaking-of-the-text-to-image#ran","syntology_url":"https://syntology.ai/paper/2405.16567","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2405.16567"}},"official":{"repos":["Kim-Minseon/APGP"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":7,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/on-the-algorithmic-bias-of-aligning-large","slug":"on-the-algorithmic-bias-of-aligning-large","title":"On the Algorithmic Bias of Aligning Large Language Models with RLHF: Preference Collapse and Matching Regularization","date":"2024-05-26","arxiv_id":"2405.16455","repositories_listed":1,"syntology":{"n":14,"n_ran":11,"n_constructed":0,"n_ran_checked":10,"n_instrument":1,"n_unverified":3,"n_honours":0,"n_violates":0,"n_no_contract":10,"n_pointer_only":14,"phrase":"11 ran (of which 0 constructed an object rather than computing a result; 10 with no instrument failure: 0 honoured, 0 violated, 10 with no contract checked; 1 where Syntology's instrument failed) · 3 unverified","sample_list":"/paper/on-the-algorithmic-bias-of-aligning-large#ran","syntology_url":"https://syntology.ai/paper/2405.16455","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2405.16455"}},"official":{"repos":["JiancongXiao/PM_RLHF"],"state":"official (archive's flag): 11 ran","n_ran":11,"n_constructed":0,"n_ran_no_instrument_failure":10,"n_unverified":3,"ran_from_kinds":["official"]}}},{"url":"/paper/a-misleading-gallery-of-fluid-motion-by","slug":"a-misleading-gallery-of-fluid-motion-by","title":"A Misleading Gallery of Fluid Motion by Generative Artificial Intelligence","date":"2024-05-24","arxiv_id":"2405.15406","repositories_listed":1,"syntology":null},{"url":"/paper/certifiably-robust-rag-against-retrieval","slug":"certifiably-robust-rag-against-retrieval","title":"Certifiably Robust RAG against Retrieval Corruption","date":"2024-05-24","arxiv_id":"2405.15556","repositories_listed":1,"syntology":null},{"url":"/paper/text-generation-a-systematic-literature","slug":"text-generation-a-systematic-literature","title":"Text Generation: A Systematic Literature Review of Tasks, Evaluation, and Challenges","date":"2024-05-24","arxiv_id":"2405.15604","repositories_listed":1,"syntology":null},{"url":"/paper/vb-lora-extreme-parameter-efficient-fine","slug":"vb-lora-extreme-parameter-efficient-fine","title":"VB-LoRA: Extreme Parameter Efficient Fine-Tuning with Vector Banks","date":"2024-05-24","arxiv_id":"2405.15179","repositories_listed":1,"syntology":{"n":17,"n_ran":17,"n_constructed":0,"n_ran_checked":11,"n_instrument":6,"n_unverified":0,"n_honours":0,"n_violates":1,"n_no_contract":10,"n_pointer_only":17,"phrase":"17 ran (of which 0 constructed an object rather than computing a result; 11 with no instrument failure: 0 honoured, 1 violated, 10 with no contract checked; 6 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/vb-lora-extreme-parameter-efficient-fine#ran","syntology_url":"https://syntology.ai/paper/2405.15179","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2405.15179"}},"official":{"repos":["leo-yangli/vb-lora"],"state":"official (archive's flag): 17 ran","n_ran":17,"n_constructed":0,"n_ran_no_instrument_failure":11,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/structural-entities-extraction-and-patient","slug":"structural-entities-extraction-and-patient","title":"Structural Entities Extraction and Patient Indications Incorporation for Chest X-ray Report Generation","date":"2024-05-23","arxiv_id":"2405.14905","repositories_listed":1,"syntology":null},{"url":"/paper/unveiling-the-achilles-heel-of-nlg-evaluators","slug":"unveiling-the-achilles-heel-of-nlg-evaluators","title":"Unveiling the Achilles' Heel of NLG Evaluators: A Unified Adversarial Framework Driven by Large Language Models","date":"2024-05-23","arxiv_id":"2405.14646","repositories_listed":1,"syntology":null},{"url":"/paper/vikhr-the-family-of-open-source-instruction","slug":"vikhr-the-family-of-open-source-instruction","title":"Vikhr: Constructing a State-of-the-art Bilingual Open-Source Instruction-Following Large Language Model for Russian","date":"2024-05-22","arxiv_id":"2405.13929","repositories_listed":1,"syntology":null},{"url":"/paper/olaph-improving-factuality-in-biomedical-long","slug":"olaph-improving-factuality-in-biomedical-long","title":"OLAPH: Improving Factuality in Biomedical Long-form Question Answering","date":"2024-05-21","arxiv_id":"2405.12701","repositories_listed":1,"syntology":null},{"url":"/paper/prott3-protein-to-text-generation-for-text","slug":"prott3-protein-to-text-generation-for-text","title":"ProtT3: Protein-to-Text Generation for Text-based Protein Understanding","date":"2024-05-21","arxiv_id":"2405.12564","repositories_listed":1,"syntology":{"n":13,"n_ran":8,"n_constructed":0,"n_ran_checked":8,"n_instrument":0,"n_unverified":5,"n_honours":0,"n_violates":1,"n_no_contract":7,"n_pointer_only":13,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 1 violated, 7 with no contract checked; 0 where Syntology's instrument failed) · 5 unverified","sample_list":"/paper/prott3-protein-to-text-generation-for-text#ran","syntology_url":"https://syntology.ai/paper/2405.12564","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2405.12564"}},"official":{"repos":["acharkq/prott3"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":5,"ran_from_kinds":["official"]}}},{"url":"/paper/sirllm-streaming-infinite-retentive-llm","slug":"sirllm-streaming-infinite-retentive-llm","title":"SirLLM: Streaming Infinite Retentive LLM","date":"2024-05-21","arxiv_id":"2405.12528","repositories_listed":1,"syntology":{"n":9,"n_ran":7,"n_constructed":1,"n_ran_checked":3,"n_instrument":4,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":3,"n_pointer_only":9,"phrase":"7 ran (of which 1 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 4 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/sirllm-streaming-infinite-retentive-llm#ran","syntology_url":"https://syntology.ai/paper/2405.12528","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2405.12528"}},"official":{"repos":["zoeyyao27/sirllm"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":1,"n_ran_no_instrument_failure":3,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/unveiling-and-manipulating-prompt-influence","slug":"unveiling-and-manipulating-prompt-influence","title":"Unveiling and Manipulating Prompt Influence in Large Language Models","date":"2024-05-20","arxiv_id":"2405.11891","repositories_listed":1,"syntology":{"n":5,"n_ran":4,"n_constructed":0,"n_ran_checked":0,"n_instrument":4,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":5,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 4 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/unveiling-and-manipulating-prompt-influence#ran","syntology_url":"https://syntology.ai/paper/2405.11891","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2405.11891"}},"official":{"repos":["zijian678/tdd"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/language-models-can-evaluate-themselves-via","slug":"language-models-can-evaluate-themselves-via","title":"Language Models can Evaluate Themselves via Probability Discrepancy","date":"2024-05-17","arxiv_id":"2405.10516","repositories_listed":1,"syntology":null},{"url":"/paper/spor-a-comprehensive-and-practical-evaluation","slug":"spor-a-comprehensive-and-practical-evaluation","title":"SPOR: A Comprehensive and Practical Evaluation Method for Compositional Generalization in Data-to-Text Generation","date":"2024-05-17","arxiv_id":"2405.10650","repositories_listed":1,"syntology":{"n":9,"n_ran":6,"n_constructed":1,"n_ran_checked":3,"n_instrument":3,"n_unverified":3,"n_honours":1,"n_violates":0,"n_no_contract":2,"n_pointer_only":0,"phrase":"6 ran (of which 1 constructed an object rather than computing a result; 3 with no instrument failure: 1 honoured, 0 violated, 2 with no contract checked; 3 where Syntology's instrument failed) · 3 unverified","sample_list":"/paper/spor-a-comprehensive-and-practical-evaluation#ran","syntology_url":"https://syntology.ai/paper/2405.10650","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2405.10650"}},"official":{"repos":["xzy-xzy/spor"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":1,"n_ran_no_instrument_failure":3,"n_unverified":3,"ran_from_kinds":["official"]}}},{"url":"/paper/debate-devil-s-advocate-based-assessment-and","slug":"debate-devil-s-advocate-based-assessment-and","title":"DEBATE: Devil's Advocate-Based Assessment and Text Evaluation","date":"2024-05-16","arxiv_id":"2405.09935","repositories_listed":1,"syntology":null},{"url":"/paper/a-comprehensive-survey-of-accelerated","slug":"a-comprehensive-survey-of-accelerated","title":"A Comprehensive Survey of Accelerated Generation Techniques in Large Language Models","date":"2024-05-15","arxiv_id":"2405.13019","repositories_listed":1,"syntology":null},{"url":"/paper/factual-serialization-enhancement-a-key","slug":"factual-serialization-enhancement-a-key","title":"Factual Serialization Enhancement: A Key Innovation for Chest X-ray Report Generation","date":"2024-05-15","arxiv_id":"2405.09586","repositories_listed":1,"syntology":null},{"url":"/paper/quite-good-but-not-enough-nationality-bias-in","slug":"quite-good-but-not-enough-nationality-bias-in","title":"Quite Good, but Not Enough: Nationality Bias in Large Language Models -- A Case Study of ChatGPT","date":"2024-05-11","arxiv_id":"2405.06996","repositories_listed":1,"syntology":null},{"url":"/paper/pseudo-prompt-generating-in-pre-trained","slug":"pseudo-prompt-generating-in-pre-trained","title":"Pseudo-Prompt Generating in Pre-trained Vision-Language Models for Multi-Label Medical Image Classification","date":"2024-05-10","arxiv_id":"2405.06468","repositories_listed":1,"syntology":null},{"url":"/paper/muting-whisper-a-universal-acoustic","slug":"muting-whisper-a-universal-acoustic","title":"Muting Whisper: A Universal Acoustic Adversarial Attack on Speech Foundation Models","date":"2024-05-09","arxiv_id":"2405.06134","repositories_listed":1,"syntology":null},{"url":"/paper/explanation-as-a-watermark-towards-harmless","slug":"explanation-as-a-watermark-towards-harmless","title":"Explanation as a Watermark: Towards Harmless and Multi-bit Model Ownership Verification via Watermarking Feature Attribution","date":"2024-05-08","arxiv_id":"2405.04825","repositories_listed":1,"syntology":{"n":11,"n_ran":9,"n_constructed":0,"n_ran_checked":9,"n_instrument":0,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":9,"n_pointer_only":0,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 9 with no instrument failure: 0 honoured, 0 violated, 9 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/explanation-as-a-watermark-towards-harmless#ran","syntology_url":"https://syntology.ai/paper/2405.04825","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2405.04825"}},"official":{"repos":["shaoshuo-ss/eaaw"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":9,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/meta-learning-for-efficient-fine-tuning-of","slug":"meta-learning-for-efficient-fine-tuning-of","title":"Meta-Learning for Efficient Fine-Tuning of Large Language Models","date":"2024-05-08","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/parameter-efficient-fine-tuning-with-discrete","slug":"parameter-efficient-fine-tuning-with-discrete","title":"Parameter-Efficient Fine-Tuning with Discrete Fourier Transform","date":"2024-05-05","arxiv_id":"2405.03003","repositories_listed":1,"syntology":null},{"url":"/paper/enhancing-contextual-understanding-in-large","slug":"enhancing-contextual-understanding-in-large","title":"Enhancing Contextual Understanding in Large Language Models through Contrastive Decoding","date":"2024-05-04","arxiv_id":"2405.02750","repositories_listed":1,"syntology":{"n":1,"n_ran":0,"n_constructed":0,"n_ran_checked":0,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":1,"phrase":"0 ran · 1 unverified","sample_list":"/paper/enhancing-contextual-understanding-in-large#ran","syntology_url":"https://syntology.ai/paper/2405.02750","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2405.02750"}},"official":{"repos":["amazon-science/contextualunderstanding-contrastivedecoding"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":[]}}},{"url":"/paper/investigating-wit-creativity-and","slug":"investigating-wit-creativity-and","title":"Investigating Wit, Creativity, and Detectability of Large Language Models in Domain-Specific Writing Style Adaptation of Reddit's Showerthoughts","date":"2024-05-02","arxiv_id":"2405.01660","repositories_listed":1,"syntology":null},{"url":"/paper/integrating-a-i-in-higher-education-protocol","slug":"integrating-a-i-in-higher-education-protocol","title":"Integrating A.I. in Higher Education: Protocol for a Pilot Study with 'SAMCares: An Adaptive Learning Hub'","date":"2024-05-01","arxiv_id":"2405.00330","repositories_listed":1,"syntology":null},{"url":"/paper/countering-reward-over-optimization-in-llm","slug":"countering-reward-over-optimization-in-llm","title":"Countering Reward Over-optimization in LLM with Demonstration-Guided Reinforcement Learning","date":"2024-04-30","arxiv_id":"2404.19409","repositories_listed":1,"syntology":null},{"url":"/paper/pecc-problem-extraction-and-coding-challenges","slug":"pecc-problem-extraction-and-coding-challenges","title":"PECC: Problem Extraction and Coding Challenges","date":"2024-04-29","arxiv_id":"2404.18766","repositories_listed":1,"syntology":null},{"url":"/paper/ceval-a-benchmark-for-evaluating","slug":"ceval-a-benchmark-for-evaluating","title":"CEval: A Benchmark for Evaluating Counterfactual Text Generation","date":"2024-04-26","arxiv_id":"2404.17475","repositories_listed":1,"syntology":null},{"url":"/paper/when-to-trust-llms-aligning-confidence-with","slug":"when-to-trust-llms-aligning-confidence-with","title":"When to Trust LLMs: Aligning Confidence with Response Quality","date":"2024-04-26","arxiv_id":"2404.17287","repositories_listed":1,"syntology":{"n":12,"n_ran":11,"n_constructed":0,"n_ran_checked":10,"n_instrument":1,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":10,"n_pointer_only":12,"phrase":"11 ran (of which 0 constructed an object rather than computing a result; 10 with no instrument failure: 0 honoured, 0 violated, 10 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/when-to-trust-llms-aligning-confidence-with#ran","syntology_url":"https://syntology.ai/paper/2404.17287","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2404.17287"}},"official":{"repos":["taoshuchang/conqord"],"state":"official (archive's flag): 11 ran","n_ran":11,"n_constructed":0,"n_ran_no_instrument_failure":10,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/large-language-models-in-healthcare-a","slug":"large-language-models-in-healthcare-a","title":"Large Language Models in the Clinic: A Comprehensive Benchmark","date":"2024-04-25","arxiv_id":"2405.00716","repositories_listed":1,"syntology":null},{"url":"/paper/simulating-task-oriented-dialogues-with-state","slug":"simulating-task-oriented-dialogues-with-state","title":"Simulating Task-Oriented Dialogues with State Transition Graphs and Large Language Models","date":"2024-04-23","arxiv_id":"2404.14772","repositories_listed":1,"syntology":null},{"url":"/paper/towards-smallers-faster-decoder-only","slug":"towards-smallers-faster-decoder-only","title":"Towards smaller, faster decoder-only transformers: Architectural variants and their implications","date":"2024-04-22","arxiv_id":"2404.14462","repositories_listed":1,"syntology":null},{"url":"/paper/ladic-are-diffusion-models-really-inferior-to","slug":"ladic-are-diffusion-models-really-inferior-to","title":"LaDiC: Are Diffusion Models Really Inferior to Autoregressive Counterparts for Image-to-Text Generation?","date":"2024-04-16","arxiv_id":"2404.10763","repositories_listed":1,"syntology":null},{"url":"/paper/modeling-low-resource-health-coaching","slug":"modeling-low-resource-health-coaching","title":"Modeling Low-Resource Health Coaching Dialogues via Neuro-Symbolic Goal Summarization and Text-Units-Text Generation","date":"2024-04-16","arxiv_id":"2404.10268","repositories_listed":1,"syntology":null},{"url":"/paper/bridging-the-gap-between-different","slug":"bridging-the-gap-between-different","title":"Bridging the Gap between Different Vocabularies for LLM Ensemble","date":"2024-04-15","arxiv_id":"2404.09492","repositories_listed":1,"syntology":{"n":6,"n_ran":4,"n_constructed":0,"n_ran_checked":0,"n_instrument":4,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":6,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 4 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/bridging-the-gap-between-different#ran","syntology_url":"https://syntology.ai/paper/2404.09492","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2404.09492"}},"official":{"repos":["xydaytoy/eva"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/wikisplit-easy-data-refinement-for-split-and","slug":"wikisplit-easy-data-refinement-for-split-and","title":"WikiSplit++: Easy Data Refinement for Split and Rephrase","date":"2024-04-13","arxiv_id":"2404.09002","repositories_listed":1,"syntology":null},{"url":"/paper/continuous-language-model-interpolation-for","slug":"continuous-language-model-interpolation-for","title":"Continuous Language Model Interpolation for Dynamic and Controllable Text Generation","date":"2024-04-10","arxiv_id":"2404.07117","repositories_listed":1,"syntology":{"n":8,"n_ran":5,"n_constructed":0,"n_ran_checked":5,"n_instrument":0,"n_unverified":3,"n_honours":0,"n_violates":0,"n_no_contract":5,"n_pointer_only":8,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","sample_list":"/paper/continuous-language-model-interpolation-for#ran","syntology_url":"https://syntology.ai/paper/2404.07117","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2404.07117"}},"official":{"repos":["skangasl/continuous-lm-interpolation"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":3,"ran_from_kinds":["official"]}}},{"url":"/paper/control-dag-constrained-decoding-for-non","slug":"control-dag-constrained-decoding-for-non","title":"Control-DAG: Constrained Decoding for Non-Autoregressive Directed Acyclic T5 using Weighted Finite State Automata","date":"2024-04-10","arxiv_id":"2404.06854","repositories_listed":1,"syntology":null},{"url":"/paper/take-a-look-at-it-rethinking-how-to-evaluate","slug":"take-a-look-at-it-rethinking-how-to-evaluate","title":"Rethinking How to Evaluate Language Model Jailbreak","date":"2024-04-09","arxiv_id":"2404.06407","repositories_listed":1,"syntology":null},{"url":"/paper/benchmarking-and-improving-compositional","slug":"benchmarking-and-improving-compositional","title":"Benchmarking and Improving Compositional Generalization of Multi-aspect Controllable Text Generation","date":"2024-04-05","arxiv_id":"2404.04232","repositories_listed":1,"syntology":{"n":8,"n_ran":8,"n_constructed":0,"n_ran_checked":8,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":8,"n_pointer_only":0,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 0 violated, 8 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/benchmarking-and-improving-compositional#ran","syntology_url":"https://syntology.ai/paper/2404.04232","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2404.04232"}},"official":{"repos":["tqzhong/cg4mctg"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/uncertainty-in-language-models-assessment","slug":"uncertainty-in-language-models-assessment","title":"Uncertainty in Language Models: Assessment through Rank-Calibration","date":"2024-04-04","arxiv_id":"2404.03163","repositories_listed":1,"syntology":{"n":2,"n_ran":2,"n_constructed":0,"n_ran_checked":0,"n_instrument":2,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/uncertainty-in-language-models-assessment#ran","syntology_url":"https://syntology.ai/paper/2404.03163","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2404.03163"}},"official":{"repos":["shuoli90/rank-calibration"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/on-the-multilingual-ability-of-decoder-based","slug":"on-the-multilingual-ability-of-decoder-based","title":"On the Multilingual Ability of Decoder-based Pre-trained Language Models: Finding and Controlling Language-Specific Neurons","date":"2024-04-03","arxiv_id":"2404.02431","repositories_listed":1,"syntology":null},{"url":"/paper/prompting-for-numerical-sequences-a-case","slug":"prompting-for-numerical-sequences-a-case","title":"Prompting for Numerical Sequences: A Case Study on Market Comment Generation","date":"2024-04-03","arxiv_id":"2404.02466","repositories_listed":1,"syntology":null},{"url":"/paper/developing-safe-and-responsible-large","slug":"developing-safe-and-responsible-large","title":"Developing Safe and Responsible Large Language Model : Can We Balance Bias Reduction and Language Understanding in Large Language Models?","date":"2024-04-01","arxiv_id":"2404.01399","repositories_listed":1,"syntology":{"n":2,"n_ran":1,"n_constructed":0,"n_ran_checked":1,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":2,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/developing-safe-and-responsible-large#ran","syntology_url":"https://syntology.ai/paper/2404.01399","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2404.01399"}},"official":{"repos":["shainarazavi/safe-responsible-llm"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/from-pixels-to-graphs-open-vocabulary-scene","slug":"from-pixels-to-graphs-open-vocabulary-scene","title":"From Pixels to Graphs: Open-Vocabulary Scene Graph Generation with Vision-Language Models","date":"2024-04-01","arxiv_id":"2404.00906","repositories_listed":1,"syntology":{"n":9,"n_ran":7,"n_constructed":0,"n_ran_checked":6,"n_instrument":1,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":6,"n_pointer_only":1,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/from-pixels-to-graphs-open-vocabulary-scene#ran","syntology_url":"https://syntology.ai/paper/2404.00906","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2404.00906"}},"official":{"repos":["shtuplus/pix2grp_cvpr2024"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/learning-by-correction-efficient-tuning-task","slug":"learning-by-correction-efficient-tuning-task","title":"Learning by Correction: Efficient Tuning Task for Zero-Shot Generative Vision-Language Reasoning","date":"2024-04-01","arxiv_id":"2404.00909","repositories_listed":1,"syntology":{"n":10,"n_ran":8,"n_constructed":1,"n_ran_checked":4,"n_instrument":4,"n_unverified":2,"n_honours":1,"n_violates":2,"n_no_contract":1,"n_pointer_only":0,"phrase":"8 ran (of which 1 constructed an object rather than computing a result; 4 with no instrument failure: 1 honoured, 2 violated, 1 with no contract checked; 4 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/learning-by-correction-efficient-tuning-task#ran","syntology_url":"https://syntology.ai/paper/2404.00909","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2404.00909"}},"official":{"repos":["shtuplus/iccc_cvpr2024"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":1,"n_ran_no_instrument_failure":4,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/set-aligning-framework-for-auto-regressive","slug":"set-aligning-framework-for-auto-regressive","title":"Set-Aligning Framework for Auto-Regressive Event Temporal Graph Generation","date":"2024-04-01","arxiv_id":"2404.01532","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":1,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":1,"n_no_contract":0,"n_pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 1 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/set-aligning-framework-for-auto-regressive#ran","syntology_url":"https://syntology.ai/paper/2404.01532","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2404.01532"}},"official":{"repos":["xingwei-warwick/set-aligning-event-temporal-graph-generation"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/humane-speech-synthesis-through-zero-shot","slug":"humane-speech-synthesis-through-zero-shot","title":"Humane Speech Synthesis through Zero-Shot Emotion and Disfluency Generation","date":"2024-03-31","arxiv_id":"2404.01339","repositories_listed":1,"syntology":null},{"url":"/paper/on-the-true-distribution-approximation-of","slug":"on-the-true-distribution-approximation-of","title":"On the True Distribution Approximation of Minimum Bayes-Risk Decoding","date":"2024-03-31","arxiv_id":"2404.00752","repositories_listed":1,"syntology":null},{"url":"/paper/luq-long-text-uncertainty-quantification-for","slug":"luq-long-text-uncertainty-quantification-for","title":"LUQ: Long-text Uncertainty Quantification for LLMs","date":"2024-03-29","arxiv_id":"2403.20279","repositories_listed":1,"syntology":null},{"url":"/paper/the-impact-of-prompts-on-zero-shot-detection","slug":"the-impact-of-prompts-on-zero-shot-detection","title":"The Impact of Prompts on Zero-Shot Detection of AI-Generated Text","date":"2024-03-29","arxiv_id":"2403.20127","repositories_listed":1,"syntology":{"n":9,"n_ran":8,"n_constructed":0,"n_ran_checked":8,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":8,"n_pointer_only":9,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 0 violated, 8 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/the-impact-of-prompts-on-zero-shot-detection#ran","syntology_url":"https://syntology.ai/paper/2403.20127","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2403.20127"}},"official":{"repos":["kaito25atugich/detector"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/img2loc-revisiting-image-geolocalization","slug":"img2loc-revisiting-image-geolocalization","title":"Img2Loc: Revisiting Image Geolocalization using Multi-modality Foundation Models and Image-based Retrieval-Augmented Generation","date":"2024-03-28","arxiv_id":"2403.19584","repositories_listed":1,"syntology":null},{"url":"/paper/omniparser-a-unified-framework-for-text","slug":"omniparser-a-unified-framework-for-text","title":"OmniParser: A Unified Framework for Text Spotting, Key Information Extraction and Table Recognition","date":"2024-03-28","arxiv_id":"2403.19128","repositories_listed":1,"syntology":{"n":4,"n_ran":3,"n_constructed":1,"n_ran_checked":2,"n_instrument":1,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":0,"phrase":"3 ran (of which 1 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/omniparser-a-unified-framework-for-text#ran","syntology_url":"https://syntology.ai/paper/2403.19128","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2403.19128"}},"official":{"repos":["alibabaresearch/advancedliteratemachinery"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":1,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/enhanced-generative-recommendation-via","slug":"enhanced-generative-recommendation-via","title":"Content-Based Collaborative Generation for Recommender Systems","date":"2024-03-27","arxiv_id":"2403.18480","repositories_listed":1,"syntology":{"n":9,"n_ran":8,"n_constructed":0,"n_ran_checked":8,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":8,"n_pointer_only":9,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 0 violated, 8 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/enhanced-generative-recommendation-via#ran","syntology_url":"https://syntology.ai/paper/2403.18480","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2403.18480"}},"official":{"repos":["junewang0614/colarec"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/lita-language-instructed-temporal","slug":"lita-language-instructed-temporal","title":"LITA: Language Instructed Temporal-Localization Assistant","date":"2024-03-27","arxiv_id":"2403.19046","repositories_listed":1,"syntology":{"n":2,"n_ran":2,"n_constructed":0,"n_ran_checked":0,"n_instrument":2,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/lita-language-instructed-temporal#ran","syntology_url":"https://syntology.ai/paper/2403.19046","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2403.19046"}},"official":{"repos":["nvlabs/lita"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/scaling-laws-for-dense-retrieval","slug":"scaling-laws-for-dense-retrieval","title":"Scaling Laws For Dense Retrieval","date":"2024-03-27","arxiv_id":"2403.18684","repositories_listed":1,"syntology":{"n":8,"n_ran":7,"n_constructed":0,"n_ran_checked":7,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":7,"n_pointer_only":0,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/scaling-laws-for-dense-retrieval#ran","syntology_url":"https://syntology.ai/paper/2403.18684","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2403.18684"}},"official":{"repos":["jingtaozhan/drscale"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":7,"n_unverified":1,"ran_from_kinds":["official"]}}}],"record_sha256":"76efd5abb9b10ca469755d527d6036ddad795615d9bcfed055f7af5fdaee1e5f","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}