{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/task/question-answering/papers/ran/8","list_of":"/task/question-answering","task":"Question Answering","archive":{"snapshot":"2025-07-28"},"key_notes":{"n_ran_checked":"legacy name, kept unchanged so existing readers do not break: it counts the samples that ran with no instrument failure (honoured, violated, and ran with no contract checked); it does not mean a contract was checked, and the pages print it as 'K with no instrument failure', not 'K checked'","n_constructed":"a sub-count of the samples that ran, never subtracted from them and never a failure: an executed sample whose run returned an instance of its own class (fixture_out_type equals the entry name): the run built an object and did not compute a result (Syntology's RAN record, counts.constructed)"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"ran","order_definition":"only papers where Syntology ran at least one harvested sample; date (newest first), ties by arXiv id","caption":"We ran code from the paper's repository; we did not run it on this task or check it against the task's benchmarks.","absence":"A paper missing from this list is not a recorded non-run: it may have no arXiv id, no harvested code, or only samples that have not run yet.","page":8,"pages_in_order":13,"rows_per_page":100,"rows":[701,800],"of":1274,"counts":{"archive_papers_tagged":10817,"with_a_code_link":4171,"where_syntology_ran_a_sample":1274,"not_listed_spam_title":0,"listed":10817,"listed_where_code_ran":1274,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":1073,"every_run_a_failure_of_syntologys_instrument":201,"listed_with_a_run_with_no_instrument_failure":1073,"listed_every_run_a_failure_of_syntologys_instrument":201,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/task/question-answering/papers/ran/1","prev":"/task/question-answering/papers/ran/7","next":"/task/question-answering/papers/ran/9","papers":[{"url":"/paper/pengi-an-audio-language-model-for-audio-tasks-1","slug":"pengi-an-audio-language-model-for-audio-tasks-1","title":"Pengi: An Audio Language Model for Audio Tasks","date":"2023-05-19","arxiv_id":"2305.11834","repositories_listed":1,"syntology":{"n":4,"n_ran":3,"n_constructed":0,"n_ran_checked":2,"n_instrument":1,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":1,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/pengi-an-audio-language-model-for-audio-tasks-1#ran","syntology_url":"https://syntology.ai/paper/2305.11834","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2305.11834"}},"official":{"repos":["microsoft/pengi"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/one-peace-exploring-one-general","slug":"one-peace-exploring-one-general","title":"ONE-PEACE: Exploring One General Representation Model Toward Unlimited Modalities","date":"2023-05-18","arxiv_id":"2305.11172","repositories_listed":2,"syntology":{"n":7,"n_ran":5,"n_constructed":2,"n_ran_checked":4,"n_instrument":1,"n_unverified":2,"n_honours":1,"n_violates":0,"n_no_contract":3,"n_pointer_only":1,"phrase":"5 ran (of which 2 constructed an object rather than computing a result; 4 with no instrument failure: 1 honoured, 0 violated, 3 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/one-peace-exploring-one-general#ran","syntology_url":"https://syntology.ai/paper/2305.11172","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2305.11172"}},"official":{"repos":["OFA-Sys/ONE-PEACE"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":2,"n_ran_no_instrument_failure":4,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/can-language-models-solve-graph-problems-in","slug":"can-language-models-solve-graph-problems-in","title":"Can Language Models Solve Graph Problems in Natural Language?","date":"2023-05-17","arxiv_id":"2305.10037","repositories_listed":2,"syntology":{"n":22,"n_ran":11,"n_constructed":0,"n_ran_checked":9,"n_instrument":2,"n_unverified":11,"n_honours":0,"n_violates":0,"n_no_contract":9,"n_pointer_only":0,"phrase":"11 ran (of which 0 constructed an object rather than computing a result; 9 with no instrument failure: 0 honoured, 0 violated, 9 with no contract checked; 2 where Syntology's instrument failed) · 11 unverified","sample_list":"/paper/can-language-models-solve-graph-problems-in#ran","syntology_url":"https://syntology.ai/paper/2305.10037","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2305.10037"}},"official":{"repos":["arthur-heng/nlgraph"],"state":"official (archive's flag): 11 ran","n_ran":11,"n_constructed":0,"n_ran_no_instrument_failure":9,"n_unverified":11,"ran_from_kinds":["official"]}}},{"url":"/paper/pmc-vqa-visual-instruction-tuning-for-medical","slug":"pmc-vqa-visual-instruction-tuning-for-medical","title":"PMC-VQA: Visual Instruction Tuning for Medical Visual Question Answering","date":"2023-05-17","arxiv_id":"2305.10415","repositories_listed":2,"syntology":{"n":5,"n_ran":4,"n_constructed":0,"n_ran_checked":4,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":1,"n_no_contract":3,"n_pointer_only":2,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 1 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/pmc-vqa-visual-instruction-tuning-for-medical#ran","syntology_url":"https://syntology.ai/paper/2305.10415","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2305.10415"}},"official":{"repos":["xiaoman-zhang/PMC-VQA"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/imad-image-augmented-multi-modal-dialogue","slug":"imad-image-augmented-multi-modal-dialogue","title":"IMAD: IMage-Augmented multi-modal Dialogue","date":"2023-05-17","arxiv_id":"2305.10512","repositories_listed":1,"syntology":{"n":2,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/imad-image-augmented-multi-modal-dialogue#ran","syntology_url":"https://syntology.ai/paper/2305.10512","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2305.10512"}},"official":{"repos":["vityavitalich/imad"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/tree-of-thoughts-deliberate-problem-solving-1","slug":"tree-of-thoughts-deliberate-problem-solving-1","title":"Tree of Thoughts: Deliberate Problem Solving with Large Language Models","date":"2023-05-17","arxiv_id":"2305.10601","repositories_listed":6,"syntology":{"n":24,"n_ran":19,"n_constructed":6,"n_ran_checked":18,"n_instrument":1,"n_unverified":5,"n_honours":0,"n_violates":0,"n_no_contract":18,"n_pointer_only":1,"phrase":"19 ran (of which 6 constructed an object rather than computing a result; 18 with no instrument failure: 0 honoured, 0 violated, 18 with no contract checked; 1 where Syntology's instrument failed) · 5 unverified","sample_list":"/paper/tree-of-thoughts-deliberate-problem-solving-1#ran","syntology_url":"https://syntology.ai/paper/2305.10601","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2305.10601"}},"official":{"repos":["princeton-nlp/tree-of-thought-llm","ysymyth/tree-of-thought-llm"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["community","listed","official"]}}},{"url":"/paper/structgpt-a-general-framework-for-large","slug":"structgpt-a-general-framework-for-large","title":"StructGPT: A General Framework for Large Language Model to Reason over Structured Data","date":"2023-05-16","arxiv_id":"2305.09645","repositories_listed":1,"syntology":{"n":10,"n_ran":8,"n_constructed":2,"n_ran_checked":8,"n_instrument":0,"n_unverified":2,"n_honours":0,"n_violates":3,"n_no_contract":5,"n_pointer_only":1,"phrase":"8 ran (of which 2 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 3 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/structgpt-a-general-framework-for-large#ran","syntology_url":"https://syntology.ai/paper/2305.09645","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2305.09645"}},"official":{"repos":["rucaibox/structgpt"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":2,"n_ran_no_instrument_failure":8,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/the-interpreter-understands-your-meaning-end","slug":"the-interpreter-understands-your-meaning-end","title":"The Interpreter Understands Your Meaning: End-to-end Spoken Language Understanding Aided by Speech Translation","date":"2023-05-16","arxiv_id":"2305.09652","repositories_listed":1,"syntology":{"n":11,"n_ran":10,"n_constructed":0,"n_ran_checked":10,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":10,"n_pointer_only":0,"phrase":"10 ran (of which 0 constructed an object rather than computing a result; 10 with no instrument failure: 0 honoured, 0 violated, 10 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/the-interpreter-understands-your-meaning-end#ran","syntology_url":"https://syntology.ai/paper/2305.09652","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2305.09652"}},"official":{"repos":["idiap/translation-aided-slu"],"state":"official (archive's flag): 10 ran","n_ran":10,"n_constructed":0,"n_ran_no_instrument_failure":10,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/matsci-nlp-evaluating-scientific-language","slug":"matsci-nlp-evaluating-scientific-language","title":"MatSci-NLP: Evaluating Scientific Language Models on Materials Science Language Tasks Using Text-to-Schema Modeling","date":"2023-05-14","arxiv_id":"2305.08264","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":1,"n_ran_checked":1,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":1,"phrase":"1 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified; the one sample that ran constructed an object rather than computing a result","sample_list":"/paper/matsci-nlp-evaluating-scientific-language#ran","syntology_url":"https://syntology.ai/paper/2305.08264","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2305.08264"}},"official":{"repos":["banglab-udem-mila/nlp4matsci-acl23"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/scene-self-labeled-counterfactuals-for","slug":"scene-self-labeled-counterfactuals-for","title":"SCENE: Self-Labeled Counterfactuals for Extrapolating to Negative Examples","date":"2023-05-13","arxiv_id":"2305.07984","repositories_listed":1,"syntology":{"n":3,"n_ran":2,"n_constructed":0,"n_ran_checked":1,"n_instrument":1,"n_unverified":1,"n_honours":0,"n_violates":1,"n_no_contract":0,"n_pointer_only":3,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 1 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/scene-self-labeled-counterfactuals-for#ran","syntology_url":"https://syntology.ai/paper/2305.07984","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2305.07984"}},"official":{"repos":["deqingfu/scene"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/webcpm-interactive-web-search-for-chinese","slug":"webcpm-interactive-web-search-for-chinese","title":"WebCPM: Interactive Web Search for Chinese Long-form Question Answering","date":"2023-05-11","arxiv_id":"2305.06849","repositories_listed":1,"syntology":{"n":4,"n_ran":3,"n_constructed":0,"n_ran_checked":3,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":3,"n_pointer_only":1,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/webcpm-interactive-web-search-for-chinese#ran","syntology_url":"https://syntology.ai/paper/2305.06849","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2305.06849"}},"official":{"repos":["thunlp/webcpm"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/evaluating-open-domain-question-answering-in","slug":"evaluating-open-domain-question-answering-in","title":"Evaluating Open-Domain Question Answering in the Era of Large Language Models","date":"2023-05-11","arxiv_id":"2305.06984","repositories_listed":1,"syntology":{"n":6,"n_ran":4,"n_constructed":0,"n_ran_checked":3,"n_instrument":1,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":3,"n_pointer_only":1,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/evaluating-open-domain-question-answering-in#ran","syntology_url":"https://syntology.ai/paper/2305.06984","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2305.06984"}},"official":{"repos":["ehsk/openqa-eval"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/self-chained-image-language-model-for-video-1","slug":"self-chained-image-language-model-for-video-1","title":"Self-Chained Image-Language Model for Video Localization and Question Answering","date":"2023-05-11","arxiv_id":"2305.06988","repositories_listed":1,"syntology":{"n":7,"n_ran":6,"n_constructed":0,"n_ran_checked":6,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":6,"n_pointer_only":7,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/self-chained-image-language-model-for-video-1#ran","syntology_url":"https://syntology.ai/paper/2305.06988","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2305.06988"}},"official":{"repos":["yui010206/sevila"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/say-what-you-mean-large-language-models-speak","slug":"say-what-you-mean-large-language-models-speak","title":"Say What You Mean! Large Language Models Speak Too Positively about Negative Commonsense Knowledge","date":"2023-05-10","arxiv_id":"2305.05976","repositories_listed":1,"syntology":{"n":4,"n_ran":4,"n_constructed":0,"n_ran_checked":2,"n_instrument":2,"n_unverified":0,"n_honours":1,"n_violates":1,"n_no_contract":0,"n_pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 1 honoured, 1 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/say-what-you-mean-large-language-models-speak#ran","syntology_url":"https://syntology.ai/paper/2305.05976","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2305.05976"}},"official":{"repos":["jiangjiechen/uncommongen"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/automatic-evaluation-of-attribution-by-large","slug":"automatic-evaluation-of-attribution-by-large","title":"Automatic Evaluation of Attribution by Large Language Models","date":"2023-05-10","arxiv_id":"2305.06311","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":1,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/automatic-evaluation-of-attribution-by-large#ran","syntology_url":"https://syntology.ai/paper/2305.06311","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2305.06311"}},"official":{"repos":["osu-nlp-group/attrscore"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/anetqa-a-large-scale-benchmark-for-fine","slug":"anetqa-a-large-scale-benchmark-for-fine","title":"ANetQA: A Large-scale Benchmark for Fine-grained Compositional Reasoning over Untrimmed Videos","date":"2023-05-04","arxiv_id":"2305.02519","repositories_listed":1,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":2,"n_instrument":1,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/anetqa-a-large-scale-benchmark-for-fine#ran","syntology_url":"https://syntology.ai/paper/2305.02519","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2305.02519"}},"official":{"repos":["MILVLG/anetqa-code"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/faithful-question-answering-with-monte-carlo","slug":"faithful-question-answering-with-monte-carlo","title":"Faithful Question Answering with Monte-Carlo Planning","date":"2023-05-04","arxiv_id":"2305.02556","repositories_listed":1,"syntology":{"n":4,"n_ran":3,"n_constructed":3,"n_ran_checked":3,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":3,"n_pointer_only":0,"phrase":"3 ran (of which 3 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified; every one of the 3 samples that ran constructed an object rather than computing a result","sample_list":"/paper/faithful-question-answering-with-monte-carlo#ran","syntology_url":"https://syntology.ai/paper/2305.02556","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2305.02556"}},"official":{"repos":["Raising-hrx/FAME"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":3,"n_ran_no_instrument_failure":3,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/retromae-2-duplex-masked-auto-encoder-for-pre","slug":"retromae-2-duplex-masked-auto-encoder-for-pre","title":"RetroMAE-2: Duplex Masked Auto-Encoder For Pre-Training Retrieval-Oriented Language Models","date":"2023-05-04","arxiv_id":"2305.02564","repositories_listed":1,"syntology":{"n":6,"n_ran":6,"n_constructed":0,"n_ran_checked":6,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":6,"n_pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/retromae-2-duplex-masked-auto-encoder-for-pre#ran","syntology_url":"https://syntology.ai/paper/2305.02564","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2305.02564"}},"official":{"repos":["staoxiao/retromae"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/few-shot-in-context-learning-for-knowledge","slug":"few-shot-in-context-learning-for-knowledge","title":"Few-shot In-context Learning for Knowledge Base Question Answering","date":"2023-05-02","arxiv_id":"2305.01750","repositories_listed":1,"syntology":{"n":16,"n_ran":9,"n_constructed":0,"n_ran_checked":9,"n_instrument":0,"n_unverified":7,"n_honours":0,"n_violates":0,"n_no_contract":9,"n_pointer_only":0,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 9 with no instrument failure: 0 honoured, 0 violated, 9 with no contract checked; 0 where Syntology's instrument failed) · 7 unverified","sample_list":"/paper/few-shot-in-context-learning-for-knowledge#ran","syntology_url":"https://syntology.ai/paper/2305.01750","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2305.01750"}},"official":{"repos":["ltl3a87/kb-binder"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":9,"n_unverified":7,"ran_from_kinds":["official"]}}},{"url":"/paper/pmc-llama-further-finetuning-llama-on-medical","slug":"pmc-llama-further-finetuning-llama-on-medical","title":"PMC-LLaMA: Towards Building Open-source Language Models for Medicine","date":"2023-04-27","arxiv_id":"2304.14454","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":0,"n_instrument":1,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/pmc-llama-further-finetuning-llama-on-medical#ran","syntology_url":"https://syntology.ai/paper/2304.14454","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2304.14454"}},"official":{"repos":["chaoyi-wu/pmc-llama"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/answering-questions-by-meta-reasoning-over","slug":"answering-questions-by-meta-reasoning-over","title":"Answering Questions by Meta-Reasoning over Multiple Chains of Thought","date":"2023-04-25","arxiv_id":"2304.13007","repositories_listed":1,"syntology":{"n":10,"n_ran":9,"n_constructed":0,"n_ran_checked":8,"n_instrument":1,"n_unverified":1,"n_honours":1,"n_violates":0,"n_no_contract":7,"n_pointer_only":0,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 1 honoured, 0 violated, 7 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/answering-questions-by-meta-reasoning-over#ran","syntology_url":"https://syntology.ai/paper/2304.13007","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2304.13007"}},"official":{"repos":["oriyor/reasoning-on-cots"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/mpmqa-multimodal-question-answering-on","slug":"mpmqa-multimodal-question-answering-on","title":"MPMQA: Multimodal Question Answering on Product Manuals","date":"2023-04-19","arxiv_id":"2304.09660","repositories_listed":1,"syntology":{"n":11,"n_ran":8,"n_constructed":0,"n_ran_checked":8,"n_instrument":0,"n_unverified":3,"n_honours":0,"n_violates":0,"n_no_contract":8,"n_pointer_only":0,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 0 violated, 8 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","sample_list":"/paper/mpmqa-multimodal-question-answering-on#ran","syntology_url":"https://syntology.ai/paper/2304.09660","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2304.09660"}},"official":{"repos":["aim3-ruc/mpmqa"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":3,"ran_from_kinds":["official"]}}},{"url":"/paper/surgicalgpt-end-to-end-language-vision-gpt","slug":"surgicalgpt-end-to-end-language-vision-gpt","title":"SurgicalGPT: End-to-End Language-Vision GPT for Visual Question Answering in Surgery","date":"2023-04-19","arxiv_id":"2304.09974","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":1,"n_instrument":0,"n_unverified":0,"n_honours":1,"n_violates":0,"n_no_contract":0,"n_pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/surgicalgpt-end-to-end-language-vision-gpt#ran","syntology_url":"https://syntology.ai/paper/2304.09974","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2304.09974"}},"official":{"repos":["lalithjets/surgicalgpt"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/learning-situation-hyper-graphs-for-video","slug":"learning-situation-hyper-graphs-for-video","title":"Learning Situation Hyper-Graphs for Video Question Answering","date":"2023-04-18","arxiv_id":"2304.08682","repositories_listed":1,"syntology":{"n":15,"n_ran":8,"n_constructed":5,"n_ran_checked":8,"n_instrument":0,"n_unverified":7,"n_honours":0,"n_violates":0,"n_no_contract":8,"n_pointer_only":15,"phrase":"8 ran (of which 5 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 0 violated, 8 with no contract checked; 0 where Syntology's instrument failed) · 7 unverified","sample_list":"/paper/learning-situation-hyper-graphs-for-video#ran","syntology_url":"https://syntology.ai/paper/2304.08682","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2304.08682"}},"official":{"repos":["aurooj/shg-vqa"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":5,"n_ran_no_instrument_failure":8,"n_unverified":7,"ran_from_kinds":["official"]}}},{"url":"/paper/verbs-in-action-improving-verb-understanding","slug":"verbs-in-action-improving-verb-understanding","title":"Verbs in Action: Improving verb understanding in video-language models","date":"2023-04-13","arxiv_id":"2304.06708","repositories_listed":1,"syntology":{"n":3,"n_ran":2,"n_constructed":0,"n_ran_checked":1,"n_instrument":1,"n_unverified":1,"n_honours":1,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/verbs-in-action-improving-verb-understanding#ran","syntology_url":"https://syntology.ai/paper/2304.06708","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2304.06708"}},"official":{"repos":["google-research/scenic"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/mphayaner-named-entity-recognition-for","slug":"mphayaner-named-entity-recognition-for","title":"MphayaNER: Named Entity Recognition for Tshivenda","date":"2023-04-08","arxiv_id":"2304.03952","repositories_listed":1,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":2,"n_instrument":1,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":1,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/mphayaner-named-entity-recognition-for#ran","syntology_url":"https://syntology.ai/paper/2304.03952","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2304.03952"}},"official":{"repos":["rendanim/mphayaner"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/mammut-a-simple-architecture-for-joint","slug":"mammut-a-simple-architecture-for-joint","title":"MaMMUT: A Simple Architecture for Joint Learning for MultiModal Tasks","date":"2023-03-29","arxiv_id":"2303.16839","repositories_listed":1,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":3,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":3,"n_no_contract":0,"n_pointer_only":2,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 3 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/mammut-a-simple-architecture-for-joint#ran","syntology_url":"https://syntology.ai/paper/2303.16839","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2303.16839"}},"official":null}},{"url":"/paper/explicit-planning-helps-language-models-in","slug":"explicit-planning-helps-language-models-in","title":"Explicit Planning Helps Language Models in Logical Reasoning","date":"2023-03-28","arxiv_id":"2303.15714","repositories_listed":2,"syntology":{"n":1,"n_ran":1,"n_constructed":1,"n_ran_checked":1,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":0,"phrase":"1 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified; the one sample that ran constructed an object rather than computing a result","sample_list":"/paper/explicit-planning-helps-language-models-in#ran","syntology_url":"https://syntology.ai/paper/2303.15714","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2303.15714"}},"official":{"repos":["cindermond/explicit-planning-for-reasoning","cindermond/leap"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/mgtbench-benchmarking-machine-generated-text","slug":"mgtbench-benchmarking-machine-generated-text","title":"MGTBench: Benchmarking Machine-Generated Text Detection","date":"2023-03-26","arxiv_id":"2303.14822","repositories_listed":4,"syntology":{"n":26,"n_ran":19,"n_constructed":0,"n_ran_checked":18,"n_instrument":1,"n_unverified":7,"n_honours":0,"n_violates":0,"n_no_contract":18,"n_pointer_only":5,"phrase":"19 ran (of which 0 constructed an object rather than computing a result; 18 with no instrument failure: 0 honoured, 0 violated, 18 with no contract checked; 1 where Syntology's instrument failed) · 7 unverified","sample_list":"/paper/mgtbench-benchmarking-machine-generated-text#ran","syntology_url":"https://syntology.ai/paper/2303.14822","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2303.14822"}},"official":{"repos":["xinleihe/mgtbench"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"url":"/paper/video-text-as-game-players-hierarchical","slug":"video-text-as-game-players-hierarchical","title":"Video-Text as Game Players: Hierarchical Banzhaf Interaction for Cross-Modal Representation Learning","date":"2023-03-25","arxiv_id":"2303.14369","repositories_listed":4,"syntology":{"n":16,"n_ran":12,"n_constructed":7,"n_ran_checked":11,"n_instrument":1,"n_unverified":4,"n_honours":0,"n_violates":0,"n_no_contract":11,"n_pointer_only":0,"phrase":"12 ran (of which 7 constructed an object rather than computing a result; 11 with no instrument failure: 0 honoured, 0 violated, 11 with no contract checked; 1 where Syntology's instrument failed) · 4 unverified","sample_list":"/paper/video-text-as-game-players-hierarchical#ran","syntology_url":"https://syntology.ai/paper/2303.14369","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2303.14369"}},"official":{"repos":["jpthu17/HBI"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["listed","official"]}}},{"url":"/paper/meltr-meta-loss-transformer-for-learning-to","slug":"meltr-meta-loss-transformer-for-learning-to","title":"MELTR: Meta Loss Transformer for Learning to Fine-tune Video Foundation Models","date":"2023-03-23","arxiv_id":"2303.13009","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":1,"n_ran_checked":1,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":0,"phrase":"1 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified; the one sample that ran constructed an object rather than computing a result","sample_list":"/paper/meltr-meta-loss-transformer-for-learning-to#ran","syntology_url":"https://syntology.ai/paper/2303.13009","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2303.13009"}},"official":{"repos":["mlvlab/MELTR"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/tifa-accurate-and-interpretable-text-to-image","slug":"tifa-accurate-and-interpretable-text-to-image","title":"TIFA: Accurate and Interpretable Text-to-Image Faithfulness Evaluation with Question Answering","date":"2023-03-21","arxiv_id":"2303.11897","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":0,"n_instrument":1,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/tifa-accurate-and-interpretable-text-to-image#ran","syntology_url":"https://syntology.ai/paper/2303.11897","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2303.11897"}},"official":{"repos":["Yushi-Hu/tifa"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/adaptive-budget-allocation-for-parameter","slug":"adaptive-budget-allocation-for-parameter","title":"AdaLoRA: Adaptive Budget Allocation for Parameter-Efficient Fine-Tuning","date":"2023-03-18","arxiv_id":"2303.10512","repositories_listed":2,"syntology":{"n":2,"n_ran":2,"n_constructed":1,"n_ran_checked":2,"n_instrument":0,"n_unverified":0,"n_honours":1,"n_violates":0,"n_no_contract":1,"n_pointer_only":0,"phrase":"2 ran (of which 1 constructed an object rather than computing a result; 2 with no instrument failure: 1 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/adaptive-budget-allocation-for-parameter#ran","syntology_url":"https://syntology.ai/paper/2303.10512","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2303.10512"}},"official":{"repos":["qingruzhang/adalora"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":1,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/gpt-4-technical-report-1","slug":"gpt-4-technical-report-1","title":"GPT-4 Technical Report","date":"2023-03-15","arxiv_id":"2303.08774","repositories_listed":11,"syntology":{"n":5,"n_ran":5,"n_constructed":0,"n_ran_checked":5,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":2,"n_no_contract":3,"n_pointer_only":1,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 2 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/gpt-4-technical-report-1#ran","syntology_url":"https://syntology.ai/paper/2303.08774","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2303.08774"}},"official":{"repos":["openai/evals"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"url":"/paper/chatgpt-asks-blip-2-answers-automatic","slug":"chatgpt-asks-blip-2-answers-automatic","title":"ChatGPT Asks, BLIP-2 Answers: Automatic Questioning Towards Enriched Visual Descriptions","date":"2023-03-12","arxiv_id":"2303.06594","repositories_listed":1,"syntology":{"n":8,"n_ran":6,"n_constructed":0,"n_ran_checked":6,"n_instrument":0,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":6,"n_pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/chatgpt-asks-blip-2-answers-automatic#ran","syntology_url":"https://syntology.ai/paper/2303.06594","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2303.06594"}},"official":{"repos":["vision-cair/chatcaptioner"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/open-ended-medical-visual-question-answering","slug":"open-ended-medical-visual-question-answering","title":"Open-Ended Medical Visual Question Answering Through Prefix Tuning of Language Models","date":"2023-03-10","arxiv_id":"2303.05977","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":1,"n_instrument":0,"n_unverified":0,"n_honours":1,"n_violates":0,"n_no_contract":0,"n_pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/open-ended-medical-visual-question-answering#ran","syntology_url":"https://syntology.ai/paper/2303.05977","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2303.05977"}},"official":{"repos":["tjvsonsbeek/open-ended-medical-vqa"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/prompting-large-language-models-with-answer","slug":"prompting-large-language-models-with-answer","title":"Prophet: Prompting Large Language Models with Complementary Answer Heuristics for Knowledge-based Visual Question Answering","date":"2023-03-03","arxiv_id":"2303.01903","repositories_listed":1,"syntology":{"n":8,"n_ran":6,"n_constructed":0,"n_ran_checked":6,"n_instrument":0,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":6,"n_pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/prompting-large-language-models-with-answer#ran","syntology_url":"https://syntology.ai/paper/2303.01903","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2303.01903"}},"official":{"repos":["milvlg/prophet"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/wice-real-world-entailment-for-claims-in","slug":"wice-real-world-entailment-for-claims-in","title":"WiCE: Real-World Entailment for Claims in Wikipedia","date":"2023-03-02","arxiv_id":"2303.01432","repositories_listed":2,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":2,"n_instrument":1,"n_unverified":0,"n_honours":1,"n_violates":1,"n_no_contract":0,"n_pointer_only":3,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 1 honoured, 1 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/wice-real-world-entailment-for-claims-in#ran","syntology_url":"https://syntology.ai/paper/2303.01432","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2303.01432"}},"official":{"repos":["ryokamoi/wice"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/contrastive-video-question-answering-via","slug":"contrastive-video-question-answering-via","title":"Contrastive Video Question Answering via Video Graph Transformer","date":"2023-02-27","arxiv_id":"2302.13668","repositories_listed":1,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":3,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":3,"n_pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/contrastive-video-question-answering-via#ran","syntology_url":"https://syntology.ai/paper/2302.13668","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2302.13668"}},"official":{"repos":["doc-doc/covgt"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/llama-open-and-efficient-foundation-language-1","slug":"llama-open-and-efficient-foundation-language-1","title":"LLaMA: Open and Efficient Foundation Language Models","date":"2023-02-27","arxiv_id":"2302.13971","repositories_listed":57,"syntology":{"n":58,"n_ran":37,"n_constructed":9,"n_ran_checked":25,"n_instrument":12,"n_unverified":21,"n_honours":3,"n_violates":0,"n_no_contract":22,"n_pointer_only":4,"phrase":"37 ran (of which 9 constructed an object rather than computing a result; 25 with no instrument failure: 3 honoured, 0 violated, 22 with no contract checked; 12 where Syntology's instrument failed) · 21 unverified","sample_list":"/paper/llama-open-and-efficient-foundation-language-1#ran","syntology_url":"https://syntology.ai/paper/2302.13971","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2302.13971"}},"official":{"repos":["facebookresearch/llama"],"state":"official: no sample here; runs from other or unrecorded repositories","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed","unlocated"]}}},{"url":"/paper/can-pre-trained-vision-and-language-models","slug":"can-pre-trained-vision-and-language-models","title":"Can Pre-trained Vision and Language Models Answer Visual Information-Seeking Questions?","date":"2023-02-23","arxiv_id":"2302.11713","repositories_listed":2,"syntology":{"n":6,"n_ran":3,"n_constructed":0,"n_ran_checked":3,"n_instrument":0,"n_unverified":3,"n_honours":0,"n_violates":0,"n_no_contract":3,"n_pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","sample_list":"/paper/can-pre-trained-vision-and-language-models#ran","syntology_url":"https://syntology.ai/paper/2302.11713","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2302.11713"}},"official":{"repos":["edchengg/infoseek_eval"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":3,"ran_from_kinds":["official"]}}},{"url":"/paper/connecting-vision-and-language-with-video","slug":"connecting-vision-and-language-with-video","title":"Connecting Vision and Language with Video Localized Narratives","date":"2023-02-22","arxiv_id":"2302.11217","repositories_listed":1,"syntology":{"n":4,"n_ran":4,"n_constructed":0,"n_ran_checked":4,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":4,"n_pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/connecting-vision-and-language-with-video#ran","syntology_url":"https://syntology.ai/paper/2302.11217","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2302.11217"}},"official":{"repos":["google/video-localized-narratives"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/chatgpt-jack-of-all-trades-master-of-none","slug":"chatgpt-jack-of-all-trades-master-of-none","title":"ChatGPT: Jack of all trades, master of none","date":"2023-02-21","arxiv_id":"2302.10724","repositories_listed":1,"syntology":{"n":2,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/chatgpt-jack-of-all-trades-master-of-none#ran","syntology_url":"https://syntology.ai/paper/2302.10724","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2302.10724"}},"official":{"repos":["clarin-pl/chatgpt-evaluation-01-2023"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/hyena-hierarchy-towards-larger-convolutional","slug":"hyena-hierarchy-towards-larger-convolutional","title":"Hyena Hierarchy: Towards Larger Convolutional Language Models","date":"2023-02-21","arxiv_id":"2302.10866","repositories_listed":7,"syntology":{"n":5,"n_ran":5,"n_constructed":0,"n_ran_checked":0,"n_instrument":5,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":4,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 5 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/hyena-hierarchy-towards-larger-convolutional#ran","syntology_url":"https://syntology.ai/paper/2302.10866","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2302.10866"}},"official":{"repos":["hazyresearch/safari"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed","official"]}}},{"url":"/paper/semantic-uncertainty-linguistic-invariances","slug":"semantic-uncertainty-linguistic-invariances","title":"Semantic Uncertainty: Linguistic Invariances for Uncertainty Estimation in Natural Language Generation","date":"2023-02-19","arxiv_id":"2302.09664","repositories_listed":3,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":1,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/semantic-uncertainty-linguistic-invariances#ran","syntology_url":"https://syntology.ai/paper/2302.09664","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2302.09664"}},"official":{"repos":["lorenzkuhn/semantic_uncertainty"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/compositional-exemplars-for-in-context","slug":"compositional-exemplars-for-in-context","title":"Compositional Exemplars for In-context Learning","date":"2023-02-11","arxiv_id":"2302.05698","repositories_listed":1,"syntology":{"n":7,"n_ran":5,"n_constructed":1,"n_ran_checked":2,"n_instrument":3,"n_unverified":2,"n_honours":1,"n_violates":0,"n_no_contract":1,"n_pointer_only":1,"phrase":"5 ran (of which 1 constructed an object rather than computing a result; 2 with no instrument failure: 1 honoured, 0 violated, 1 with no contract checked; 3 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/compositional-exemplars-for-in-context#ran","syntology_url":"https://syntology.ai/paper/2302.05698","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2302.05698"}},"official":{"repos":["hkunlp/icl-ceil"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":1,"n_ran_no_instrument_failure":2,"n_unverified":2,"ran_from_kinds":["official","unlocated"]}}},{"url":"/paper/exploring-the-benefits-of-training-expert","slug":"exploring-the-benefits-of-training-expert","title":"Exploring the Benefits of Training Expert Language Models over Instruction Tuning","date":"2023-02-07","arxiv_id":"2302.03202","repositories_listed":2,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":0,"n_instrument":3,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":3,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 3 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/exploring-the-benefits-of-training-expert#ran","syntology_url":"https://syntology.ai/paper/2302.03202","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2302.03202"}},"official":{"repos":["joeljang/elm"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official","unlocated"]}}},{"url":"/paper/liquid-a-framework-for-list-question","slug":"liquid-a-framework-for-list-question","title":"LIQUID: A Framework for List Question Answering Dataset Generation","date":"2023-02-03","arxiv_id":"2302.01691","repositories_listed":1,"syntology":{"n":2,"n_ran":2,"n_constructed":0,"n_ran_checked":0,"n_instrument":2,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/liquid-a-framework-for-list-question#ran","syntology_url":"https://syntology.ai/paper/2302.01691","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2302.01691"}},"official":{"repos":["dmis-lab/liquid"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/faithful-chain-of-thought-reasoning","slug":"faithful-chain-of-thought-reasoning","title":"Faithful Chain-of-Thought Reasoning","date":"2023-01-31","arxiv_id":"2301.13379","repositories_listed":1,"syntology":{"n":9,"n_ran":8,"n_constructed":0,"n_ran_checked":8,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":8,"n_pointer_only":1,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 0 violated, 8 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/faithful-chain-of-thought-reasoning#ran","syntology_url":"https://syntology.ai/paper/2301.13379","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2301.13379"}},"official":{"repos":["veronica320/faithful-cot"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/replug-retrieval-augmented-black-box-language","slug":"replug-retrieval-augmented-black-box-language","title":"REPLUG: Retrieval-Augmented Black-Box Language Models","date":"2023-01-30","arxiv_id":"2301.12652","repositories_listed":3,"syntology":{"n":13,"n_ran":10,"n_constructed":0,"n_ran_checked":10,"n_instrument":0,"n_unverified":3,"n_honours":0,"n_violates":0,"n_no_contract":10,"n_pointer_only":13,"phrase":"10 ran (of which 0 constructed an object rather than computing a result; 10 with no instrument failure: 0 honoured, 0 violated, 10 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","sample_list":"/paper/replug-retrieval-augmented-black-box-language#ran","syntology_url":"https://syntology.ai/paper/2301.12652","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2301.12652"}},"official":null}},{"url":"/paper/mqag-multiple-choice-question-answering-and","slug":"mqag-multiple-choice-question-answering-and","title":"MQAG: Multiple-choice Question Answering and Generation for Assessing Information Consistency in Summarization","date":"2023-01-28","arxiv_id":"2301.12307","repositories_listed":2,"syntology":{"n":2,"n_ran":1,"n_constructed":0,"n_ran_checked":1,"n_instrument":0,"n_unverified":1,"n_honours":1,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/mqag-multiple-choice-question-answering-and#ran","syntology_url":"https://syntology.ai/paper/2301.12307","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2301.12307"}},"official":{"repos":["potsawee/mqag0"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/thoughtsource-a-central-hub-for-large","slug":"thoughtsource-a-central-hub-for-large","title":"ThoughtSource: A central hub for large language model reasoning data","date":"2023-01-27","arxiv_id":"2301.11596","repositories_listed":1,"syntology":{"n":7,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":5,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 5 unverified","sample_list":"/paper/thoughtsource-a-central-hub-for-large#ran","syntology_url":"https://syntology.ai/paper/2301.11596","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2301.11596"}},"official":{"repos":["openbiolink/thoughtsource"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":5,"ran_from_kinds":["official"]}}},{"url":"/paper/videberta-a-powerful-pre-trained-language","slug":"videberta-a-powerful-pre-trained-language","title":"ViDeBERTa: A powerful pre-trained language model for Vietnamese","date":"2023-01-25","arxiv_id":"2301.10439","repositories_listed":1,"syntology":{"n":4,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/videberta-a-powerful-pre-trained-language#ran","syntology_url":"https://syntology.ai/paper/2301.10439","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2301.10439"}},"official":{"repos":["hysonlab/videberta"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/massive-language-models-can-be-accurately","slug":"massive-language-models-can-be-accurately","title":"SparseGPT: Massive Language Models Can Be Accurately Pruned in One-Shot","date":"2023-01-02","arxiv_id":"2301.00774","repositories_listed":6,"syntology":{"n":12,"n_ran":9,"n_constructed":0,"n_ran_checked":7,"n_instrument":2,"n_unverified":3,"n_honours":0,"n_violates":0,"n_no_contract":7,"n_pointer_only":9,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 2 where Syntology's instrument failed) · 3 unverified","sample_list":"/paper/massive-language-models-can-be-accurately#ran","syntology_url":"https://syntology.ai/paper/2301.00774","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2301.00774"}},"official":{"repos":["ist-daslab/sparsegpt"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":2,"ran_from_kinds":["listed","official"]}}},{"url":"/paper/hungry-hungry-hippos-towards-language","slug":"hungry-hungry-hippos-towards-language","title":"Hungry Hungry Hippos: Towards Language Modeling with State Space Models","date":"2022-12-28","arxiv_id":"2212.14052","repositories_listed":3,"syntology":{"n":15,"n_ran":7,"n_constructed":0,"n_ran_checked":3,"n_instrument":4,"n_unverified":8,"n_honours":1,"n_violates":0,"n_no_contract":2,"n_pointer_only":0,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 1 honoured, 0 violated, 2 with no contract checked; 4 where Syntology's instrument failed) · 8 unverified","sample_list":"/paper/hungry-hungry-hippos-towards-language#ran","syntology_url":"https://syntology.ai/paper/2212.14052","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2212.14052"}},"official":{"repos":["hazyresearch/h3"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed","official"]}}},{"url":"/paper/textbox-2-0-a-text-generation-library-with","slug":"textbox-2-0-a-text-generation-library-with","title":"TextBox 2.0: A Text Generation Library with Pre-trained Language Models","date":"2022-12-26","arxiv_id":"2212.13005","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":1,"n_ran_checked":1,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":0,"phrase":"1 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified; the one sample that ran constructed an object rather than computing a result","sample_list":"/paper/textbox-2-0-a-text-generation-library-with#ran","syntology_url":"https://syntology.ai/paper/2212.13005","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2212.13005"}},"official":{"repos":["RUCAIBox/TextBox"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/improving-complex-knowledge-base-question","slug":"improving-complex-knowledge-base-question","title":"Improving Complex Knowledge Base Question Answering via Question-to-Action and Question-to-Question Alignment","date":"2022-12-26","arxiv_id":"2212.13036","repositories_listed":1,"syntology":{"n":4,"n_ran":1,"n_constructed":0,"n_ran_checked":0,"n_instrument":1,"n_unverified":3,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":4,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 3 unverified","sample_list":"/paper/improving-complex-knowledge-base-question#ran","syntology_url":"https://syntology.ai/paper/2212.13036","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2212.13036"}},"official":{"repos":["tttttttty/alcqa"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":3,"ran_from_kinds":["official"]}}},{"url":"/paper/parallel-context-windows-improve-in-context","slug":"parallel-context-windows-improve-in-context","title":"Parallel Context Windows for Large Language Models","date":"2022-12-21","arxiv_id":"2212.10947","repositories_listed":1,"syntology":{"n":4,"n_ran":3,"n_constructed":0,"n_ran_checked":3,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":3,"n_pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/parallel-context-windows-improve-in-context#ran","syntology_url":"https://syntology.ai/paper/2212.10947","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2212.10947"}},"official":{"repos":["AI21Labs/Parallel-Context-Windows"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/language-models-are-better-than-humans-at","slug":"language-models-are-better-than-humans-at","title":"Language models are better than humans at next-token prediction","date":"2022-12-21","arxiv_id":"2212.11281","repositories_listed":1,"syntology":{"n":6,"n_ran":6,"n_constructed":0,"n_ran_checked":1,"n_instrument":5,"n_unverified":0,"n_honours":1,"n_violates":0,"n_no_contract":0,"n_pointer_only":6,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 5 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/language-models-are-better-than-humans-at#ran","syntology_url":"https://syntology.ai/paper/2212.11281","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2212.11281"}},"official":{"repos":["FabienRoger/lm-game-analysis-main"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/are-deep-neural-networks-smarter-than-second","slug":"are-deep-neural-networks-smarter-than-second","title":"Are Deep Neural Networks SMARTer than Second Graders?","date":"2022-12-20","arxiv_id":"2212.09993","repositories_listed":1,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":3,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":3,"n_pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/are-deep-neural-networks-smarter-than-second#ran","syntology_url":"https://syntology.ai/paper/2212.09993","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2212.09993"}},"official":{"repos":["merlresearch/SMART"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/toward-a-unified-framework-for-unsupervised","slug":"toward-a-unified-framework-for-unsupervised","title":"Optimization Techniques for Unsupervised Complex Table Reasoning via Self-Training Framework","date":"2022-12-20","arxiv_id":"2212.10097","repositories_listed":2,"syntology":{"n":20,"n_ran":17,"n_constructed":0,"n_ran_checked":17,"n_instrument":0,"n_unverified":3,"n_honours":0,"n_violates":0,"n_no_contract":17,"n_pointer_only":0,"phrase":"17 ran (of which 0 constructed an object rather than computing a result; 17 with no instrument failure: 0 honoured, 0 violated, 17 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","sample_list":"/paper/toward-a-unified-framework-for-unsupervised#ran","syntology_url":"https://syntology.ai/paper/2212.10097","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2212.10097"}},"official":{"repos":["leezythu/uctr","leezythu/uctr-st"],"state":"official (archive's flag): 17 ran","n_ran":17,"n_constructed":0,"n_ran_no_instrument_failure":17,"n_unverified":3,"ran_from_kinds":["official"]}}},{"url":"/paper/interleaving-retrieval-with-chain-of-thought","slug":"interleaving-retrieval-with-chain-of-thought","title":"Interleaving Retrieval with Chain-of-Thought Reasoning for Knowledge-Intensive Multi-Step Questions","date":"2022-12-20","arxiv_id":"2212.10509","repositories_listed":2,"syntology":{"n":2,"n_ran":2,"n_constructed":0,"n_ran_checked":1,"n_instrument":1,"n_unverified":0,"n_honours":0,"n_violates":1,"n_no_contract":0,"n_pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 1 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/interleaving-retrieval-with-chain-of-thought#ran","syntology_url":"https://syntology.ai/paper/2212.10509","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2212.10509"}},"official":{"repos":["stonybrooknlp/ircot"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/don-t-generate-discriminate-a-proposal-for","slug":"don-t-generate-discriminate-a-proposal-for","title":"Don't Generate, Discriminate: A Proposal for Grounding Language Models to Real-World Environments","date":"2022-12-19","arxiv_id":"2212.09736","repositories_listed":2,"syntology":{"n":3,"n_ran":1,"n_constructed":0,"n_ran_checked":0,"n_instrument":1,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":3,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/don-t-generate-discriminate-a-proposal-for#ran","syntology_url":"https://syntology.ai/paper/2212.09736","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2212.09736"}},"official":{"repos":["dki-lab/pangu"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/attributed-question-answering-evaluation-and","slug":"attributed-question-answering-evaluation-and","title":"Attributed Question Answering: Evaluation and Modeling for Attributed Large Language Models","date":"2022-12-15","arxiv_id":"2212.08037","repositories_listed":1,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":0,"n_instrument":3,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 3 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/attributed-question-answering-evaluation-and#ran","syntology_url":"https://syntology.ai/paper/2212.08037","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2212.08037"}},"official":{"repos":["google-research-datasets/attributed-qa"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/from-clozing-to-comprehending-retrofitting","slug":"from-clozing-to-comprehending-retrofitting","title":"From Cloze to Comprehension: Retrofitting Pre-trained Masked Language Model to Pre-trained Machine Reader","date":"2022-12-09","arxiv_id":"2212.04755","repositories_listed":1,"syntology":{"n":4,"n_ran":3,"n_constructed":2,"n_ran_checked":3,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":3,"n_pointer_only":4,"phrase":"3 ran (of which 2 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/from-clozing-to-comprehending-retrofitting#ran","syntology_url":"https://syntology.ai/paper/2212.04755","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2212.04755"}},"official":{"repos":["damo-nlp-sg/pmr"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":2,"n_ran_no_instrument_failure":3,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/discovering-latent-knowledge-in-language","slug":"discovering-latent-knowledge-in-language","title":"Discovering Latent Knowledge in Language Models Without Supervision","date":"2022-12-07","arxiv_id":"2212.03827","repositories_listed":1,"syntology":{"n":2,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/discovering-latent-knowledge-in-language#ran","syntology_url":"https://syntology.ai/paper/2212.03827","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2212.03827"}},"official":{"repos":["collin-burns/discovering_latent_knowledge"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/unikgqa-unified-retrieval-and-reasoning-for","slug":"unikgqa-unified-retrieval-and-reasoning-for","title":"UniKGQA: Unified Retrieval and Reasoning for Solving Multi-hop Question Answering Over Knowledge Graph","date":"2022-12-02","arxiv_id":"2212.00959","repositories_listed":1,"syntology":{"n":17,"n_ran":9,"n_constructed":0,"n_ran_checked":9,"n_instrument":0,"n_unverified":8,"n_honours":0,"n_violates":0,"n_no_contract":9,"n_pointer_only":0,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 9 with no instrument failure: 0 honoured, 0 violated, 9 with no contract checked; 0 where Syntology's instrument failed) · 8 unverified","sample_list":"/paper/unikgqa-unified-retrieval-and-reasoning-for#ran","syntology_url":"https://syntology.ai/paper/2212.00959","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2212.00959"}},"official":{"repos":["rucaibox/unikgqa"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":9,"n_unverified":8,"ran_from_kinds":["official"]}}},{"url":"/paper/super-clevr-a-virtual-benchmark-to-diagnose","slug":"super-clevr-a-virtual-benchmark-to-diagnose","title":"Super-CLEVR: A Virtual Benchmark to Diagnose Domain Robustness in Visual Reasoning","date":"2022-12-01","arxiv_id":"2212.00259","repositories_listed":2,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":0,"n_instrument":3,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":3,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 3 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/super-clevr-a-virtual-benchmark-to-diagnose#ran","syntology_url":"https://syntology.ai/paper/2212.00259","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2212.00259"}},"official":{"repos":["lizw14/super-clevr"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/frustratingly-easy-label-projection-for-cross","slug":"frustratingly-easy-label-projection-for-cross","title":"Frustratingly Easy Label Projection for Cross-lingual Transfer","date":"2022-11-28","arxiv_id":"2211.15613","repositories_listed":2,"syntology":{"n":9,"n_ran":6,"n_constructed":0,"n_ran_checked":6,"n_instrument":0,"n_unverified":3,"n_honours":0,"n_violates":0,"n_no_contract":6,"n_pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","sample_list":"/paper/frustratingly-easy-label-projection-for-cross#ran","syntology_url":"https://syntology.ai/paper/2211.15613","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2211.15613"}},"official":{"repos":["edchengg/easyproject"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":3,"ran_from_kinds":["official"]}}},{"url":"/paper/self-supervised-vision-language-pretraining","slug":"self-supervised-vision-language-pretraining","title":"Self-supervised vision-language pretraining for Medical visual question answering","date":"2022-11-24","arxiv_id":"2211.13594","repositories_listed":2,"syntology":{"n":10,"n_ran":6,"n_constructed":0,"n_ran_checked":1,"n_instrument":5,"n_unverified":4,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":3,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 5 where Syntology's instrument failed) · 4 unverified","sample_list":"/paper/self-supervised-vision-language-pretraining#ran","syntology_url":"https://syntology.ai/paper/2211.13594","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2211.13594"}},"official":{"repos":["pengfeiliheu/m2i2"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":4,"ran_from_kinds":["official"]}}},{"url":"/paper/visual-programming-compositional-visual","slug":"visual-programming-compositional-visual","title":"Visual Programming: Compositional visual reasoning without training","date":"2022-11-18","arxiv_id":"2211.11559","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":0,"n_instrument":1,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/visual-programming-compositional-visual#ran","syntology_url":"https://syntology.ai/paper/2211.11559","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2211.11559"}},"official":null}},{"url":"/paper/i-can-t-believe-there-s-no-images-learning","slug":"i-can-t-believe-there-s-no-images-learning","title":"I Can't Believe There's No Images! Learning Visual Tasks Using only Language Supervision","date":"2022-11-17","arxiv_id":"2211.09778","repositories_listed":1,"syntology":{"n":6,"n_ran":6,"n_constructed":0,"n_ran_checked":5,"n_instrument":1,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":5,"n_pointer_only":1,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/i-can-t-believe-there-s-no-images-learning#ran","syntology_url":"https://syntology.ai/paper/2211.09778","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2211.09778"}},"official":{"repos":["allenai/close"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/summarizing-community-based-question-answer","slug":"summarizing-community-based-question-answer","title":"Summarizing Community-based Question-Answer Pairs","date":"2022-11-17","arxiv_id":"2211.09892","repositories_listed":0,"syntology":{"n":8,"n_ran":7,"n_constructed":0,"n_ran_checked":6,"n_instrument":1,"n_unverified":1,"n_honours":0,"n_violates":1,"n_no_contract":5,"n_pointer_only":2,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 1 violated, 5 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/summarizing-community-based-question-answer#ran","syntology_url":"https://syntology.ai/paper/2211.09892","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2211.09892"}},"official":null}},{"url":"/paper/galactica-a-large-language-model-for-science-1","slug":"galactica-a-large-language-model-for-science-1","title":"Galactica: A Large Language Model for Science","date":"2022-11-16","arxiv_id":"2211.09085","repositories_listed":1,"syntology":{"n":2,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/galactica-a-large-language-model-for-science-1#ran","syntology_url":"https://syntology.ai/paper/2211.09085","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2211.09085"}},"official":{"repos":["paperswithcode/galai"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/large-language-models-struggle-to-learn-long","slug":"large-language-models-struggle-to-learn-long","title":"Large Language Models Struggle to Learn Long-Tail Knowledge","date":"2022-11-15","arxiv_id":"2211.08411","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":1,"n_instrument":0,"n_unverified":0,"n_honours":1,"n_violates":0,"n_no_contract":0,"n_pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/large-language-models-struggle-to-learn-long#ran","syntology_url":"https://syntology.ai/paper/2211.08411","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2211.08411"}},"official":{"repos":["nkandpa2/long_tail_knowledge"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/disentqa-disentangling-parametric-and","slug":"disentqa-disentangling-parametric-and","title":"DisentQA: Disentangling Parametric and Contextual Knowledge with Counterfactual Question Answering","date":"2022-11-10","arxiv_id":"2211.05655","repositories_listed":1,"syntology":{"n":5,"n_ran":4,"n_constructed":0,"n_ran_checked":4,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":4,"n_pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/disentqa-disentangling-parametric-and#ran","syntology_url":"https://syntology.ai/paper/2211.05655","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2211.05655"}},"official":{"repos":["ellaneeman/disent_qa"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/cripp-vqa-counterfactual-reasoning-about","slug":"cripp-vqa-counterfactual-reasoning-about","title":"CRIPP-VQA: Counterfactual Reasoning about Implicit Physical Properties via Video Question Answering","date":"2022-11-07","arxiv_id":"2211.03779","repositories_listed":1,"syntology":{"n":7,"n_ran":6,"n_constructed":0,"n_ran_checked":2,"n_instrument":4,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 4 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/cripp-vqa-counterfactual-reasoning-about#ran","syntology_url":"https://syntology.ai/paper/2211.03779","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2211.03779"}},"official":{"repos":["maitreyapatel/cripp-vqa"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":["found_in_text"]}}},{"url":"/paper/kglm-integrating-knowledge-graph-structure-in","slug":"kglm-integrating-knowledge-graph-structure-in","title":"KGLM: Integrating Knowledge Graph Structure in Language Models for Link Prediction","date":"2022-11-04","arxiv_id":"2211.02744","repositories_listed":1,"syntology":{"n":5,"n_ran":4,"n_constructed":0,"n_ran_checked":4,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":4,"n_pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/kglm-integrating-knowledge-graph-structure-in#ran","syntology_url":"https://syntology.ai/paper/2211.02744","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2211.02744"}},"official":{"repos":["ibpa/kglm"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/crosslingual-generalization-through-multitask","slug":"crosslingual-generalization-through-multitask","title":"Crosslingual Generalization through Multitask Finetuning","date":"2022-11-03","arxiv_id":"2211.01786","repositories_listed":1,"syntology":{"n":4,"n_ran":4,"n_constructed":0,"n_ran_checked":4,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":4,"n_pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/crosslingual-generalization-through-multitask#ran","syntology_url":"https://syntology.ai/paper/2211.01786","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2211.01786"}},"official":{"repos":["bigscience-workshop/xmtf"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/transfer-learning-with-synthetic-corpora-for","slug":"transfer-learning-with-synthetic-corpora-for","title":"Transfer Learning with Synthetic Corpora for Spatial Role Labeling and Reasoning","date":"2022-10-30","arxiv_id":"2210.16952","repositories_listed":1,"syntology":{"n":12,"n_ran":10,"n_constructed":0,"n_ran_checked":10,"n_instrument":0,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":10,"n_pointer_only":0,"phrase":"10 ran (of which 0 constructed an object rather than computing a result; 10 with no instrument failure: 0 honoured, 0 violated, 10 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/transfer-learning-with-synthetic-corpora-for#ran","syntology_url":"https://syntology.ai/paper/2210.16952","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2210.16952"}},"official":{"repos":["hlr/spartun"],"state":"official (archive's flag): 10 ran","n_ran":10,"n_constructed":0,"n_ran_no_instrument_failure":10,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/two-is-better-than-many-binary-classification","slug":"two-is-better-than-many-binary-classification","title":"Two is Better than Many? Binary Classification as an Effective Approach to Multi-Choice Question Answering","date":"2022-10-29","arxiv_id":"2210.16495","repositories_listed":1,"syntology":{"n":7,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":5,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 5 unverified","sample_list":"/paper/two-is-better-than-many-binary-classification#ran","syntology_url":"https://syntology.ai/paper/2210.16495","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2210.16495"}},"official":{"repos":["declare-lab/team"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":5,"ran_from_kinds":["official"]}}},{"url":"/paper/morphte-injecting-morphology-in-tensorized","slug":"morphte-injecting-morphology-in-tensorized","title":"MorphTE: Injecting Morphology in Tensorized Embeddings","date":"2022-10-27","arxiv_id":"2210.15379","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":1,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/morphte-injecting-morphology-in-tensorized#ran","syntology_url":"https://syntology.ai/paper/2210.15379","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2210.15379"}},"official":{"repos":["bigganbing/Fairseq_MorphTE"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/dyrex-dynamic-query-representation-for","slug":"dyrex-dynamic-query-representation-for","title":"DyREx: Dynamic Query Representation for Extractive Question Answering","date":"2022-10-26","arxiv_id":"2210.15048","repositories_listed":1,"syntology":{"n":3,"n_ran":2,"n_constructed":0,"n_ran_checked":1,"n_instrument":1,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":3,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/dyrex-dynamic-query-representation-for#ran","syntology_url":"https://syntology.ai/paper/2210.15048","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2210.15048"}},"official":{"repos":["urchade/dyrex"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official","unlocated"]}}},{"url":"/paper/vlc-bert-visual-question-answering-with","slug":"vlc-bert-visual-question-answering-with","title":"VLC-BERT: Visual Question Answering with Contextualized Commonsense Knowledge","date":"2022-10-24","arxiv_id":"2210.13626","repositories_listed":1,"syntology":{"n":2,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/vlc-bert-visual-question-answering-with#ran","syntology_url":"https://syntology.ai/paper/2210.13626","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2210.13626"}},"official":{"repos":["aditya10/vlc-bert"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/leveraging-large-language-models-for-multiple","slug":"leveraging-large-language-models-for-multiple","title":"Leveraging Large Language Models for Multiple Choice Question Answering","date":"2022-10-22","arxiv_id":"2210.12353","repositories_listed":1,"syntology":{"n":7,"n_ran":6,"n_constructed":3,"n_ran_checked":3,"n_instrument":3,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":3,"n_pointer_only":0,"phrase":"6 ran (of which 3 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 3 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/leveraging-large-language-models-for-multiple#ran","syntology_url":"https://syntology.ai/paper/2210.12353","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2210.12353"}},"official":{"repos":["byu-pccl/leveraging-llms-for-mcqa"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":3,"n_ran_no_instrument_failure":3,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/prompt-tuning-can-be-much-better-than-fine","slug":"prompt-tuning-can-be-much-better-than-fine","title":"Prompt-Tuning Can Be Much Better Than Fine-Tuning on Cross-lingual Understanding With Multilingual Language Models","date":"2022-10-22","arxiv_id":"2210.12360","repositories_listed":2,"syntology":{"n":6,"n_ran":2,"n_constructed":0,"n_ran_checked":1,"n_instrument":1,"n_unverified":4,"n_honours":1,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 4 unverified","sample_list":"/paper/prompt-tuning-can-be-much-better-than-fine#ran","syntology_url":"https://syntology.ai/paper/2210.12360","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2210.12360"}},"official":{"repos":["salesforce/mpt"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":4,"ran_from_kinds":["official"]}}},{"url":"/paper/reastap-injecting-table-reasoning-skills","slug":"reastap-injecting-table-reasoning-skills","title":"ReasTAP: Injecting Table Reasoning Skills During Pre-training via Synthetic Reasoning Examples","date":"2022-10-22","arxiv_id":"2210.12374","repositories_listed":1,"syntology":{"n":11,"n_ran":8,"n_constructed":0,"n_ran_checked":8,"n_instrument":0,"n_unverified":3,"n_honours":0,"n_violates":0,"n_no_contract":8,"n_pointer_only":0,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 0 violated, 8 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","sample_list":"/paper/reastap-injecting-table-reasoning-skills#ran","syntology_url":"https://syntology.ai/paper/2210.12374","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2210.12374"}},"official":{"repos":["yale-lily/reastap"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":3,"ran_from_kinds":["official"]}}},{"url":"/paper/graphnet-graph-neural-networks-for-neutrino","slug":"graphnet-graph-neural-networks-for-neutrino","title":"GraphNeT: Graph neural networks for neutrino telescope event reconstruction","date":"2022-10-21","arxiv_id":"2210.12194","repositories_listed":1,"syntology":{"n":4,"n_ran":4,"n_constructed":0,"n_ran_checked":4,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":4,"n_pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/graphnet-graph-neural-networks-for-neutrino#ran","syntology_url":"https://syntology.ai/paper/2210.12194","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2210.12194"}},"official":{"repos":["graphnet-team/graphnet"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/scaling-instruction-finetuned-language-models","slug":"scaling-instruction-finetuned-language-models","title":"Scaling Instruction-Finetuned Language Models","date":"2022-10-20","arxiv_id":"2210.11416","repositories_listed":9,"syntology":{"n":17,"n_ran":8,"n_constructed":1,"n_ran_checked":1,"n_instrument":7,"n_unverified":9,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":2,"phrase":"8 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 7 where Syntology's instrument failed) · 9 unverified","sample_list":"/paper/scaling-instruction-finetuned-language-models#ran","syntology_url":"https://syntology.ai/paper/2210.11416","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2210.11416"}},"official":null}},{"url":"/paper/qa-domain-adaptation-using-hidden-space","slug":"qa-domain-adaptation-using-hidden-space","title":"QA Domain Adaptation using Hidden Space Augmentation and Self-Supervised Contrastive Adaptation","date":"2022-10-19","arxiv_id":"2210.10861","repositories_listed":1,"syntology":{"n":16,"n_ran":12,"n_constructed":2,"n_ran_checked":8,"n_instrument":4,"n_unverified":4,"n_honours":5,"n_violates":1,"n_no_contract":2,"n_pointer_only":16,"phrase":"12 ran (of which 2 constructed an object rather than computing a result; 8 with no instrument failure: 5 honoured, 1 violated, 2 with no contract checked; 4 where Syntology's instrument failed) · 4 unverified","sample_list":"/paper/qa-domain-adaptation-using-hidden-space#ran","syntology_url":"https://syntology.ai/paper/2210.10861","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2210.10861"}},"official":{"repos":["yueeeeeeee/self-supervised-qa"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":2,"n_ran_no_instrument_failure":5,"n_unverified":4,"ran_from_kinds":["official","unlocated"]}}},{"url":"/paper/elastic-numerical-reasoning-with-adaptive","slug":"elastic-numerical-reasoning-with-adaptive","title":"ELASTIC: Numerical Reasoning with Adaptive Symbolic Compiler","date":"2022-10-18","arxiv_id":"2210.10105","repositories_listed":1,"syntology":{"n":2,"n_ran":2,"n_constructed":1,"n_ran_checked":2,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":0,"phrase":"2 ran (of which 1 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/elastic-numerical-reasoning-with-adaptive#ran","syntology_url":"https://syntology.ai/paper/2210.10105","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2210.10105"}},"official":{"repos":["neurasearch/neurips-2022-submission-3358"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":1,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/attributed-text-generation-via-post-hoc","slug":"attributed-text-generation-via-post-hoc","title":"RARR: Researching and Revising What Language Models Say, Using Language Models","date":"2022-10-17","arxiv_id":"2210.08726","repositories_listed":2,"syntology":{"n":2,"n_ran":2,"n_constructed":0,"n_ran_checked":0,"n_instrument":2,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/attributed-text-generation-via-post-hoc#ran","syntology_url":"https://syntology.ai/paper/2210.08726","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2210.08726"}},"official":{"repos":["anthonywchen/rarr"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/plug-and-play-vqa-zero-shot-vqa-by-conjoining","slug":"plug-and-play-vqa-zero-shot-vqa-by-conjoining","title":"Plug-and-Play VQA: Zero-shot VQA by Conjoining Large Pretrained Models with Zero Training","date":"2022-10-17","arxiv_id":"2210.08773","repositories_listed":3,"syntology":{"n":3,"n_ran":2,"n_constructed":0,"n_ran_checked":0,"n_instrument":2,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":3,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/plug-and-play-vqa-zero-shot-vqa-by-conjoining#ran","syntology_url":"https://syntology.ai/paper/2210.08773","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2210.08773"}},"official":{"repos":["salesforce/lavis"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"url":"/paper/deep-bidirectional-language-knowledge-graph","slug":"deep-bidirectional-language-knowledge-graph","title":"Deep Bidirectional Language-Knowledge Graph Pretraining","date":"2022-10-17","arxiv_id":"2210.09338","repositories_listed":2,"syntology":{"n":18,"n_ran":9,"n_constructed":0,"n_ran_checked":8,"n_instrument":1,"n_unverified":9,"n_honours":1,"n_violates":0,"n_no_contract":7,"n_pointer_only":1,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 1 honoured, 0 violated, 7 with no contract checked; 1 where Syntology's instrument failed) · 9 unverified","sample_list":"/paper/deep-bidirectional-language-knowledge-graph#ran","syntology_url":"https://syntology.ai/paper/2210.09338","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2210.09338"}},"official":{"repos":["michiyasunaga/dragon"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":9,"ran_from_kinds":["official"]}}},{"url":"/paper/sqa3d-situated-question-answering-in-3d","slug":"sqa3d-situated-question-answering-in-3d","title":"SQA3D: Situated Question Answering in 3D Scenes","date":"2022-10-14","arxiv_id":"2210.07474","repositories_listed":1,"syntology":{"n":8,"n_ran":7,"n_constructed":6,"n_ran_checked":7,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":7,"n_pointer_only":0,"phrase":"7 ran (of which 6 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/sqa3d-situated-question-answering-in-3d#ran","syntology_url":"https://syntology.ai/paper/2210.07474","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2210.07474"}},"official":{"repos":["SilongYong/SQA3D"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":6,"n_ran_no_instrument_failure":7,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/mapl-parameter-efficient-adaptation-of","slug":"mapl-parameter-efficient-adaptation-of","title":"MAPL: Parameter-Efficient Adaptation of Unimodal Pre-Trained Models for Vision-Language Few-Shot Prompting","date":"2022-10-13","arxiv_id":"2210.07179","repositories_listed":1,"syntology":{"n":2,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":1,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/mapl-parameter-efficient-adaptation-of#ran","syntology_url":"https://syntology.ai/paper/2210.07179","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2210.07179"}},"official":{"repos":["mair-lab/mapl"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/long-form-video-language-pre-training-with","slug":"long-form-video-language-pre-training-with","title":"Long-Form Video-Language Pre-Training with Multimodal Temporal Contrastive Learning","date":"2022-10-12","arxiv_id":"2210.06031","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":0,"n_instrument":1,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/long-form-video-language-pre-training-with#ran","syntology_url":"https://syntology.ai/paper/2210.06031","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2210.06031"}},"official":{"repos":["microsoft/xpretrain"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/ernie-layout-layout-knowledge-enhanced-pre","slug":"ernie-layout-layout-knowledge-enhanced-pre","title":"ERNIE-Layout: Layout Knowledge Enhanced Pre-training for Visually-rich Document Understanding","date":"2022-10-12","arxiv_id":"2210.06155","repositories_listed":2,"syntology":{"n":7,"n_ran":7,"n_constructed":0,"n_ran_checked":6,"n_instrument":1,"n_unverified":0,"n_honours":0,"n_violates":1,"n_no_contract":5,"n_pointer_only":1,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 1 violated, 5 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/ernie-layout-layout-knowledge-enhanced-pre#ran","syntology_url":"https://syntology.ai/paper/2210.06155","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2210.06155"}},"official":{"repos":["PaddlePaddle/PaddleNLP"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"url":"/paper/mixed-modality-representation-learning-and","slug":"mixed-modality-representation-learning-and","title":"Mixed-modality Representation Learning and Pre-training for Joint Table-and-Text Retrieval in OpenQA","date":"2022-10-11","arxiv_id":"2210.05197","repositories_listed":1,"syntology":{"n":8,"n_ran":5,"n_constructed":0,"n_ran_checked":5,"n_instrument":0,"n_unverified":3,"n_honours":0,"n_violates":0,"n_no_contract":5,"n_pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","sample_list":"/paper/mixed-modality-representation-learning-and#ran","syntology_url":"https://syntology.ai/paper/2210.05197","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2210.05197"}},"official":{"repos":["jun-jie-huang/otter"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":3,"ran_from_kinds":["official"]}}},{"url":"/paper/towards-robust-visual-question-answering","slug":"towards-robust-visual-question-answering","title":"Towards Robust Visual Question Answering: Making the Most of Biased Samples via Contrastive Learning","date":"2022-10-10","arxiv_id":"2210.04563","repositories_listed":1,"syntology":{"n":2,"n_ran":2,"n_constructed":0,"n_ran_checked":1,"n_instrument":1,"n_unverified":0,"n_honours":0,"n_violates":1,"n_no_contract":0,"n_pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 1 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/towards-robust-visual-question-answering#ran","syntology_url":"https://syntology.ai/paper/2210.04563","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2210.04563"}},"official":{"repos":["phoebussi/mmbs"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official","unlocated"]}}}],"record_sha256":"1572b5a273b2a57a9dc780922de5450207d16cc86a486d6e3c7cca1a2ac4d48d","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}