{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/task/vision-and-language-navigation/papers/ran/1","list_of":"/task/vision-and-language-navigation","task":"Vision and Language Navigation","archive":{"snapshot":"2025-07-28"},"key_notes":{"n_ran_checked":"legacy name, kept unchanged so existing readers do not break: it counts the samples that ran with no instrument failure (honoured, violated, and ran with no contract checked); it does not mean a contract was checked, and the pages print it as 'K with no instrument failure', not 'K checked'","n_constructed":"a sub-count of the samples that ran, never subtracted from them and never a failure: an executed sample whose run returned an instance of its own class (fixture_out_type equals the entry name): the run built an object and did not compute a result (Syntology's RAN record, counts.constructed)"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"ran","order_definition":"only papers where Syntology ran at least one harvested sample; date (newest first), ties by arXiv id","caption":"We ran code from the paper's repository; we did not run it on this task or check it against the task's benchmarks.","absence":"A paper missing from this list is not a recorded non-run: it may have no arXiv id, no harvested code, or only samples that have not run yet.","page":1,"pages_in_order":1,"rows_per_page":100,"rows":[1,52],"of":52,"counts":{"archive_papers_tagged":223,"with_a_code_link":114,"where_syntology_ran_a_sample":52,"not_listed_spam_title":0,"listed":223,"listed_where_code_ran":52,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":48,"every_run_a_failure_of_syntologys_instrument":4,"listed_with_a_run_with_no_instrument_failure":48,"listed_every_run_a_failure_of_syntologys_instrument":4,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/task/vision-and-language-navigation/papers/ran/1","prev":null,"next":null,"papers":[{"url":"/paper/navmorph-a-self-evolving-world-model-for","slug":"navmorph-a-self-evolving-world-model-for","title":"NavMorph: A Self-Evolving World Model for Vision-and-Language Navigation in Continuous Environments","date":"2025-06-30","arxiv_id":"2506.23468","repositories_listed":1,"syntology":{"n":12,"n_ran":11,"n_constructed":0,"n_ran_checked":10,"n_instrument":1,"n_unverified":1,"n_honours":0,"n_violates":1,"n_no_contract":9,"n_pointer_only":12,"phrase":"11 ran (of which 0 constructed an object rather than computing a result; 10 with no instrument failure: 0 honoured, 1 violated, 9 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/navmorph-a-self-evolving-world-model-for#ran","syntology_url":"https://syntology.ai/paper/2506.23468","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2506.23468"}},"official":{"repos":["feliciaxyao/navmorph"],"state":"official (archive's flag): 11 ran","n_ran":11,"n_constructed":0,"n_ran_no_instrument_failure":10,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/general-scene-adaptation-for-vision-and","slug":"general-scene-adaptation-for-vision-and","title":"General Scene Adaptation for Vision-and-Language Navigation","date":"2025-01-29","arxiv_id":"2501.17403","repositories_listed":1,"syntology":{"n":14,"n_ran":2,"n_constructed":1,"n_ran_checked":1,"n_instrument":1,"n_unverified":12,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":0,"phrase":"2 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 1 where Syntology's instrument failed) · 12 unverified","sample_list":"/paper/general-scene-adaptation-for-vision-and#ran","syntology_url":"https://syntology.ai/paper/2501.17403","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2501.17403"}},"official":{"repos":["honghd16/gsa-vln"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":12,"ran_from_kinds":["official"]}}},{"url":"/paper/g3d-lf-generalizable-3d-language-feature","slug":"g3d-lf-generalizable-3d-language-feature","title":"g3D-LF: Generalizable 3D-Language Feature Fields for Embodied Tasks","date":"2024-11-26","arxiv_id":"2411.17030","repositories_listed":1,"syntology":{"n":3,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":3,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/g3d-lf-generalizable-3d-language-feature#ran","syntology_url":"https://syntology.ai/paper/2411.17030","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2411.17030"}},"official":{"repos":["MrZihan/g3D-LF"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/flame-learning-to-navigate-with-multimodal","slug":"flame-learning-to-navigate-with-multimodal","title":"FLAME: Learning to Navigate with Multimodal LLM in Urban Environments","date":"2024-08-20","arxiv_id":"2408.11051","repositories_listed":1,"syntology":{"n":7,"n_ran":6,"n_constructed":0,"n_ran_checked":2,"n_instrument":4,"n_unverified":1,"n_honours":0,"n_violates":1,"n_no_contract":1,"n_pointer_only":2,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 1 violated, 1 with no contract checked; 4 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/flame-learning-to-navigate-with-multimodal#ran","syntology_url":"https://syntology.ai/paper/2408.11051","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2408.11051"}},"official":{"repos":["xyz9911/FLAME"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/navgpt-2-unleashing-navigational-reasoning","slug":"navgpt-2-unleashing-navigational-reasoning","title":"NavGPT-2: Unleashing Navigational Reasoning Capability for Large Vision-Language Models","date":"2024-07-17","arxiv_id":"2407.12366","repositories_listed":1,"syntology":{"n":14,"n_ran":12,"n_constructed":0,"n_ran_checked":9,"n_instrument":3,"n_unverified":2,"n_honours":1,"n_violates":1,"n_no_contract":7,"n_pointer_only":4,"phrase":"12 ran (of which 0 constructed an object rather than computing a result; 9 with no instrument failure: 1 honoured, 1 violated, 7 with no contract checked; 3 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/navgpt-2-unleashing-navigational-reasoning#ran","syntology_url":"https://syntology.ai/paper/2407.12366","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2407.12366"}},"official":{"repos":["gengzezhou/navgpt-2"],"state":"official (archive's flag): 12 ran","n_ran":12,"n_constructed":0,"n_ran_no_instrument_failure":9,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/human-aware-vision-and-language-navigation","slug":"human-aware-vision-and-language-navigation","title":"Human-Aware Vision-and-Language Navigation: Bridging Simulation to Reality with Dynamic Human Interactions","date":"2024-06-27","arxiv_id":"2406.19236","repositories_listed":1,"syntology":{"n":10,"n_ran":9,"n_constructed":0,"n_ran_checked":9,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":9,"n_pointer_only":10,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 9 with no instrument failure: 0 honoured, 0 violated, 9 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/human-aware-vision-and-language-navigation#ran","syntology_url":"https://syntology.ai/paper/2406.19236","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2406.19236"}},"official":{"repos":["lpercc/ha3d_simulator"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":7,"n_unverified":1,"ran_from_kinds":["found_in_text","official"]}}},{"url":"/paper/citynav-language-goal-aerial-navigation","slug":"citynav-language-goal-aerial-navigation","title":"CityNav: Language-Goal Aerial Navigation Dataset with Geographic Information","date":"2024-06-20","arxiv_id":"2406.14240","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":1,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/citynav-language-goal-aerial-navigation#ran","syntology_url":"https://syntology.ai/paper/2406.14240","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2406.14240"}},"official":null}},{"url":"/paper/sim-to-real-transfer-via-3d-feature-fields","slug":"sim-to-real-transfer-via-3d-feature-fields","title":"Sim-to-Real Transfer via 3D Feature Fields for Vision-and-Language Navigation","date":"2024-06-14","arxiv_id":"2406.09798","repositories_listed":1,"syntology":{"n":3,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":3,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/sim-to-real-transfer-via-3d-feature-fields#ran","syntology_url":"https://syntology.ai/paper/2406.09798","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2406.09798"}},"official":{"repos":["MrZihan/Sim2Real-VLN-3DFF"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/vision-and-language-navigation-via-causal","slug":"vision-and-language-navigation-via-causal","title":"Vision-and-Language Navigation via Causal Learning","date":"2024-04-16","arxiv_id":"2404.10241","repositories_listed":1,"syntology":{"n":15,"n_ran":12,"n_constructed":2,"n_ran_checked":9,"n_instrument":3,"n_unverified":3,"n_honours":0,"n_violates":0,"n_no_contract":9,"n_pointer_only":5,"phrase":"12 ran (of which 2 constructed an object rather than computing a result; 9 with no instrument failure: 0 honoured, 0 violated, 9 with no contract checked; 3 where Syntology's instrument failed) · 3 unverified","sample_list":"/paper/vision-and-language-navigation-via-causal#ran","syntology_url":"https://syntology.ai/paper/2404.10241","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2404.10241"}},"official":{"repos":["crystalsixone/vln-goat"],"state":"official (archive's flag): 12 ran","n_ran":12,"n_constructed":2,"n_ran_no_instrument_failure":9,"n_unverified":3,"ran_from_kinds":["official"]}}},{"url":"/paper/lookahead-exploration-with-neural-radiance","slug":"lookahead-exploration-with-neural-radiance","title":"Lookahead Exploration with Neural Radiance Representation for Continuous Vision-Language Navigation","date":"2024-04-02","arxiv_id":"2404.01943","repositories_listed":1,"syntology":{"n":12,"n_ran":11,"n_constructed":0,"n_ran_checked":10,"n_instrument":1,"n_unverified":1,"n_honours":0,"n_violates":1,"n_no_contract":9,"n_pointer_only":12,"phrase":"11 ran (of which 0 constructed an object rather than computing a result; 10 with no instrument failure: 0 honoured, 1 violated, 9 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/lookahead-exploration-with-neural-radiance#ran","syntology_url":"https://syntology.ai/paper/2404.01943","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2404.01943"}},"official":{"repos":["mrzihan/hnr-vln"],"state":"official (archive's flag): 11 ran","n_ran":11,"n_constructed":0,"n_ran_no_instrument_failure":10,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/navcot-boosting-llm-based-vision-and-language","slug":"navcot-boosting-llm-based-vision-and-language","title":"NavCoT: Boosting LLM-Based Vision-and-Language Navigation via Learning Disentangled Reasoning","date":"2024-03-12","arxiv_id":"2403.07376","repositories_listed":1,"syntology":{"n":13,"n_ran":8,"n_constructed":0,"n_ran_checked":6,"n_instrument":2,"n_unverified":5,"n_honours":1,"n_violates":0,"n_no_contract":5,"n_pointer_only":13,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 1 honoured, 0 violated, 5 with no contract checked; 2 where Syntology's instrument failed) · 5 unverified","sample_list":"/paper/navcot-boosting-llm-based-vision-and-language#ran","syntology_url":"https://syntology.ai/paper/2403.07376","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2403.07376"}},"official":{"repos":["expectorlin/navcot"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":5,"ran_from_kinds":["official"]}}},{"url":"/paper/weblinx-real-world-website-navigation-with","slug":"weblinx-real-world-website-navigation-with","title":"WebLINX: Real-World Website Navigation with Multi-Turn Dialogue","date":"2024-02-08","arxiv_id":"2402.05930","repositories_listed":2,"syntology":{"n":27,"n_ran":21,"n_constructed":0,"n_ran_checked":20,"n_instrument":1,"n_unverified":6,"n_honours":0,"n_violates":1,"n_no_contract":19,"n_pointer_only":0,"phrase":"21 ran (of which 0 constructed an object rather than computing a result; 20 with no instrument failure: 0 honoured, 1 violated, 19 with no contract checked; 1 where Syntology's instrument failed) · 6 unverified","sample_list":"/paper/weblinx-real-world-website-navigation-with#ran","syntology_url":"https://syntology.ai/paper/2402.05930","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.05930"}},"official":{"repos":["McGill-NLP/weblinx","McGill-NLP/webllama"],"state":"official (archive's flag): 21 ran","n_ran":21,"n_constructed":0,"n_ran_no_instrument_failure":20,"n_unverified":6,"ran_from_kinds":["official"]}}},{"url":"/paper/webvln-vision-and-language-navigation-on","slug":"webvln-vision-and-language-navigation-on","title":"WebVLN: Vision-and-Language Navigation on Websites","date":"2023-12-25","arxiv_id":"2312.15820","repositories_listed":1,"syntology":{"n":8,"n_ran":8,"n_constructed":0,"n_ran_checked":7,"n_instrument":1,"n_unverified":0,"n_honours":1,"n_violates":0,"n_no_contract":6,"n_pointer_only":8,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 1 honoured, 0 violated, 6 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/webvln-vision-and-language-navigation-on#ran","syntology_url":"https://syntology.ai/paper/2312.15820","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2312.15820"}},"official":{"repos":["webvln/webvln"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":7,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/test-time-adaptive-vision-and-language","slug":"test-time-adaptive-vision-and-language","title":"Fast-Slow Test-Time Adaptation for Online Vision-and-Language Navigation","date":"2023-11-22","arxiv_id":"2311.13209","repositories_listed":1,"syntology":{"n":6,"n_ran":6,"n_constructed":1,"n_ran_checked":1,"n_instrument":5,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":6,"phrase":"6 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 5 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/test-time-adaptive-vision-and-language#ran","syntology_url":"https://syntology.ai/paper/2311.13209","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2311.13209"}},"official":{"repos":["feliciaxyao/icml2024-fstta"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/grounded-entity-landmark-adaptive-pre","slug":"grounded-entity-landmark-adaptive-pre","title":"Grounded Entity-Landmark Adaptive Pre-training for Vision-and-Language Navigation","date":"2023-08-24","arxiv_id":"2308.12587","repositories_listed":1,"syntology":{"n":30,"n_ran":23,"n_constructed":14,"n_ran_checked":20,"n_instrument":3,"n_unverified":7,"n_honours":0,"n_violates":0,"n_no_contract":20,"n_pointer_only":30,"phrase":"23 ran (of which 14 constructed an object rather than computing a result; 20 with no instrument failure: 0 honoured, 0 violated, 20 with no contract checked; 3 where Syntology's instrument failed) · 7 unverified","sample_list":"/paper/grounded-entity-landmark-adaptive-pre#ran","syntology_url":"https://syntology.ai/paper/2308.12587","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2308.12587"}},"official":{"repos":["csir1996/vln-gela"],"state":"official (archive's flag): 23 ran","n_ran":23,"n_constructed":14,"n_ran_no_instrument_failure":20,"n_unverified":7,"ran_from_kinds":["official"]}}},{"url":"/paper/aerialvln-vision-and-language-navigation-for","slug":"aerialvln-vision-and-language-navigation-for","title":"AerialVLN: Vision-and-Language Navigation for UAVs","date":"2023-08-13","arxiv_id":"2308.06735","repositories_listed":1,"syntology":{"n":9,"n_ran":7,"n_constructed":0,"n_ran_checked":7,"n_instrument":0,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":7,"n_pointer_only":9,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/aerialvln-vision-and-language-navigation-for#ran","syntology_url":"https://syntology.ai/paper/2308.06735","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2308.06735"}},"official":{"repos":["airvln/airvln"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":7,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/scaling-data-generation-in-vision-and","slug":"scaling-data-generation-in-vision-and","title":"Scaling Data Generation in Vision-and-Language Navigation","date":"2023-07-28","arxiv_id":"2307.15644","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":1,"n_instrument":0,"n_unverified":0,"n_honours":1,"n_violates":0,"n_no_contract":0,"n_pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/scaling-data-generation-in-vision-and#ran","syntology_url":"https://syntology.ai/paper/2307.15644","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2307.15644"}},"official":{"repos":["wz0919/scalevln"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/gridmm-grid-memory-map-for-vision-and","slug":"gridmm-grid-memory-map-for-vision-and","title":"GridMM: Grid Memory Map for Vision-and-Language Navigation","date":"2023-07-24","arxiv_id":"2307.12907","repositories_listed":1,"syntology":{"n":5,"n_ran":5,"n_constructed":0,"n_ran_checked":3,"n_instrument":2,"n_unverified":0,"n_honours":1,"n_violates":1,"n_no_contract":1,"n_pointer_only":5,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 1 honoured, 1 violated, 1 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/gridmm-grid-memory-map-for-vision-and#ran","syntology_url":"https://syntology.ai/paper/2307.12907","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2307.12907"}},"official":{"repos":["mrzihan/gridmm"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/learning-vision-and-language-navigation-from","slug":"learning-vision-and-language-navigation-from","title":"Learning Vision-and-Language Navigation from YouTube Videos","date":"2023-07-22","arxiv_id":"2307.11984","repositories_listed":1,"syntology":{"n":19,"n_ran":15,"n_constructed":0,"n_ran_checked":13,"n_instrument":2,"n_unverified":4,"n_honours":1,"n_violates":0,"n_no_contract":12,"n_pointer_only":2,"phrase":"15 ran (of which 0 constructed an object rather than computing a result; 13 with no instrument failure: 1 honoured, 0 violated, 12 with no contract checked; 2 where Syntology's instrument failed) · 4 unverified","sample_list":"/paper/learning-vision-and-language-navigation-from#ran","syntology_url":"https://syntology.ai/paper/2307.11984","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2307.11984"}},"official":{"repos":["jeremylinky/youtube-vln"],"state":"official (archive's flag): 15 ran","n_ran":15,"n_constructed":0,"n_ran_no_instrument_failure":13,"n_unverified":4,"ran_from_kinds":["official"]}}},{"url":"/paper/behavioral-analysis-of-vision-and-language-1","slug":"behavioral-analysis-of-vision-and-language-1","title":"Behavioral Analysis of Vision-and-Language Navigation Agents","date":"2023-07-20","arxiv_id":"2307.10790","repositories_listed":1,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":3,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":3,"n_pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/behavioral-analysis-of-vision-and-language-1#ran","syntology_url":"https://syntology.ai/paper/2307.10790","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2307.10790"}},"official":{"repos":["yoark/vln-behave"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/velma-verbalization-embodiment-of-llm-agents","slug":"velma-verbalization-embodiment-of-llm-agents","title":"VELMA: Verbalization Embodiment of LLM Agents for Vision and Language Navigation in Street View","date":"2023-07-12","arxiv_id":"2307.06082","repositories_listed":1,"syntology":{"n":6,"n_ran":6,"n_constructed":0,"n_ran_checked":6,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":6,"n_pointer_only":6,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/velma-verbalization-embodiment-of-llm-agents#ran","syntology_url":"https://syntology.ai/paper/2307.06082","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2307.06082"}},"official":{"repos":["raphael-sch/velma"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/navgpt-explicit-reasoning-in-vision-and","slug":"navgpt-explicit-reasoning-in-vision-and","title":"NavGPT: Explicit Reasoning in Vision-and-Language Navigation with Large Language Models","date":"2023-05-26","arxiv_id":"2305.16986","repositories_listed":2,"syntology":{"n":5,"n_ran":4,"n_constructed":0,"n_ran_checked":2,"n_instrument":2,"n_unverified":1,"n_honours":2,"n_violates":0,"n_no_contract":0,"n_pointer_only":2,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 2 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/navgpt-explicit-reasoning-in-vision-and#ran","syntology_url":"https://syntology.ai/paper/2305.16986","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2305.16986"}},"official":{"repos":["gengzezhou/navgpt"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["listed","official"]}}},{"url":"/paper/a-dual-semantic-aware-recurrent-global","slug":"a-dual-semantic-aware-recurrent-global","title":"A Dual Semantic-Aware Recurrent Global-Adaptive Network For Vision-and-Language Navigation","date":"2023-05-05","arxiv_id":"2305.03602","repositories_listed":1,"syntology":{"n":24,"n_ran":10,"n_constructed":5,"n_ran_checked":6,"n_instrument":4,"n_unverified":14,"n_honours":0,"n_violates":0,"n_no_contract":6,"n_pointer_only":0,"phrase":"10 ran (of which 5 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 4 where Syntology's instrument failed) · 14 unverified","sample_list":"/paper/a-dual-semantic-aware-recurrent-global#ran","syntology_url":"https://syntology.ai/paper/2305.03602","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2305.03602"}},"official":{"repos":["crystalsixone/dsrg"],"state":"official (archive's flag): 10 ran","n_ran":10,"n_constructed":5,"n_ran_no_instrument_failure":6,"n_unverified":14,"ran_from_kinds":["official"]}}},{"url":"/paper/kerm-knowledge-enhanced-reasoning-for-vision","slug":"kerm-knowledge-enhanced-reasoning-for-vision","title":"KERM: Knowledge Enhanced Reasoning for Vision-and-Language Navigation","date":"2023-03-28","arxiv_id":"2303.15796","repositories_listed":1,"syntology":{"n":10,"n_ran":1,"n_constructed":0,"n_ran_checked":1,"n_instrument":0,"n_unverified":9,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 9 unverified","sample_list":"/paper/kerm-knowledge-enhanced-reasoning-for-vision#ran","syntology_url":"https://syntology.ai/paper/2303.15796","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2303.15796"}},"official":{"repos":["xiangyangli20/kerm"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":9,"ran_from_kinds":["official"]}}},{"url":"/paper/simple-and-effective-synthesis-of-indoor-3d","slug":"simple-and-effective-synthesis-of-indoor-3d","title":"Simple and Effective Synthesis of Indoor 3D Scenes","date":"2022-04-06","arxiv_id":"2204.02960","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":1,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/simple-and-effective-synthesis-of-indoor-3d#ran","syntology_url":"https://syntology.ai/paper/2204.02960","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2204.02960"}},"official":{"repos":["google-research/se3ds"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/envedit-environment-editing-for-vision-and","slug":"envedit-environment-editing-for-vision-and","title":"EnvEdit: Environment Editing for Vision-and-Language Navigation","date":"2022-03-29","arxiv_id":"2203.15685","repositories_listed":1,"syntology":{"n":5,"n_ran":4,"n_constructed":0,"n_ran_checked":3,"n_instrument":1,"n_unverified":1,"n_honours":1,"n_violates":0,"n_no_contract":2,"n_pointer_only":1,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 1 honoured, 0 violated, 2 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/envedit-environment-editing-for-vision-and#ran","syntology_url":"https://syntology.ai/paper/2203.15685","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2203.15685"}},"official":{"repos":["jialuli-luka/envedit"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/fedvln-privacy-preserving-federated-vision","slug":"fedvln-privacy-preserving-federated-vision","title":"FedVLN: Privacy-preserving Federated Vision-and-Language Navigation","date":"2022-03-28","arxiv_id":"2203.14936","repositories_listed":1,"syntology":{"n":10,"n_ran":4,"n_constructed":2,"n_ran_checked":3,"n_instrument":1,"n_unverified":6,"n_honours":0,"n_violates":0,"n_no_contract":3,"n_pointer_only":0,"phrase":"4 ran (of which 2 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 1 where Syntology's instrument failed) · 6 unverified","sample_list":"/paper/fedvln-privacy-preserving-federated-vision#ran","syntology_url":"https://syntology.ai/paper/2203.14936","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2203.14936"}},"official":{"repos":["eric-ai-lab/fedvln"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":2,"n_ran_no_instrument_failure":3,"n_unverified":6,"ran_from_kinds":["official"]}}},{"url":"/paper/analyzing-generalization-of-vision-and","slug":"analyzing-generalization-of-vision-and","title":"Analyzing Generalization of Vision and Language Navigation to Unseen Outdoor Areas","date":"2022-03-25","arxiv_id":"2203.13838","repositories_listed":1,"syntology":{"n":4,"n_ran":3,"n_constructed":0,"n_ran_checked":3,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":3,"n_pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/analyzing-generalization-of-vision-and#ran","syntology_url":"https://syntology.ai/paper/2203.13838","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2203.13838"}},"official":{"repos":["raphael-sch/map2seq_vln"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/hop-history-and-order-aware-pre-training-for","slug":"hop-history-and-order-aware-pre-training-for","title":"HOP: History-and-Order Aware Pre-training for Vision-and-Language Navigation","date":"2022-03-22","arxiv_id":"2203.11591","repositories_listed":1,"syntology":{"n":12,"n_ran":8,"n_constructed":0,"n_ran_checked":5,"n_instrument":3,"n_unverified":4,"n_honours":1,"n_violates":0,"n_no_contract":4,"n_pointer_only":2,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 1 honoured, 0 violated, 4 with no contract checked; 3 where Syntology's instrument failed) · 4 unverified","sample_list":"/paper/hop-history-and-order-aware-pre-training-for#ran","syntology_url":"https://syntology.ai/paper/2203.11591","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2203.11591"}},"official":{"repos":["yanyuanqiao/hop-vln"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":4,"ran_from_kinds":["official"]}}},{"url":"/paper/cross-modal-map-learning-for-vision-and","slug":"cross-modal-map-learning-for-vision-and","title":"Cross-modal Map Learning for Vision and Language Navigation","date":"2022-03-10","arxiv_id":"2203.05137","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":1,"n_ran_checked":1,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":1,"phrase":"1 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified; the one sample that ran constructed an object rather than computing a result","sample_list":"/paper/cross-modal-map-learning-for-vision-and#ran","syntology_url":"https://syntology.ai/paper/2203.05137","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2203.05137"}},"official":{"repos":["ggeorgak11/cm2"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/bridging-the-gap-between-learning-in-discrete","slug":"bridging-the-gap-between-learning-in-discrete","title":"Bridging the Gap Between Learning in Discrete and Continuous Environments for Vision-and-Language Navigation","date":"2022-03-05","arxiv_id":"2203.02764","repositories_listed":1,"syntology":{"n":2,"n_ran":1,"n_constructed":0,"n_ran_checked":1,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/bridging-the-gap-between-learning-in-discrete#ran","syntology_url":"https://syntology.ai/paper/2203.02764","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2203.02764"}},"official":{"repos":["yiconghong/discrete-continuous-vln"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/think-global-act-local-dual-scale-graph","slug":"think-global-act-local-dual-scale-graph","title":"Think Global, Act Local: Dual-scale Graph Transformer for Vision-and-Language Navigation","date":"2022-02-23","arxiv_id":"2202.11742","repositories_listed":1,"syntology":{"n":5,"n_ran":4,"n_constructed":0,"n_ran_checked":2,"n_instrument":2,"n_unverified":1,"n_honours":0,"n_violates":1,"n_no_contract":1,"n_pointer_only":5,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 1 violated, 1 with no contract checked; 2 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/think-global-act-local-dual-scale-graph#ran","syntology_url":"https://syntology.ai/paper/2202.11742","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2202.11742"}},"official":null}},{"url":"/paper/history-aware-multimodal-transformer-for","slug":"history-aware-multimodal-transformer-for","title":"History Aware Multimodal Transformer for Vision-and-Language Navigation","date":"2021-10-25","arxiv_id":"2110.13309","repositories_listed":1,"syntology":{"n":15,"n_ran":7,"n_constructed":3,"n_ran_checked":7,"n_instrument":0,"n_unverified":8,"n_honours":0,"n_violates":0,"n_no_contract":7,"n_pointer_only":0,"phrase":"7 ran (of which 3 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 0 where Syntology's instrument failed) · 8 unverified","sample_list":"/paper/history-aware-multimodal-transformer-for#ran","syntology_url":"https://syntology.ai/paper/2110.13309","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2110.13309"}},"official":null}},{"url":"/paper/sasra-semantically-aware-spatio-temporal","slug":"sasra-semantically-aware-spatio-temporal","title":"SASRA: Semantically-aware Spatio-temporal Reasoning Agent for Vision-and-Language Navigation in Continuous Environments","date":"2021-08-26","arxiv_id":"2108.11945","repositories_listed":1,"syntology":{"n":2,"n_ran":1,"n_constructed":0,"n_ran_checked":1,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/sasra-semantically-aware-spatio-temporal#ran","syntology_url":"https://syntology.ai/paper/2108.11945","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2108.11945"}},"official":null}},{"url":"/paper/airbert-in-domain-pretraining-for-vision-and","slug":"airbert-in-domain-pretraining-for-vision-and","title":"Airbert: In-domain Pretraining for Vision-and-Language Navigation","date":"2021-08-20","arxiv_id":"2108.09105","repositories_listed":2,"syntology":{"n":14,"n_ran":10,"n_constructed":0,"n_ran_checked":8,"n_instrument":2,"n_unverified":4,"n_honours":1,"n_violates":0,"n_no_contract":7,"n_pointer_only":2,"phrase":"10 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 1 honoured, 0 violated, 7 with no contract checked; 2 where Syntology's instrument failed) · 4 unverified","sample_list":"/paper/airbert-in-domain-pretraining-for-vision-and#ran","syntology_url":"https://syntology.ai/paper/2108.09105","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2108.09105"}},"official":{"repos":["airbert-vln/airbert"],"state":"official (archive's flag): 10 ran","n_ran":10,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":4,"ran_from_kinds":["official"]}}},{"url":"/paper/neighbor-view-enhanced-model-for-vision-and","slug":"neighbor-view-enhanced-model-for-vision-and","title":"Neighbor-view Enhanced Model for Vision and Language Navigation","date":"2021-07-15","arxiv_id":"2107.07201","repositories_listed":1,"syntology":{"n":4,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/neighbor-view-enhanced-model-for-vision-and#ran","syntology_url":"https://syntology.ai/paper/2107.07201","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2107.07201"}},"official":{"repos":["MarSaKi/NvEM"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/how-much-can-clip-benefit-vision-and-language","slug":"how-much-can-clip-benefit-vision-and-language","title":"How Much Can CLIP Benefit Vision-and-Language Tasks?","date":"2021-07-13","arxiv_id":"2107.06383","repositories_listed":4,"syntology":{"n":10,"n_ran":7,"n_constructed":0,"n_ran_checked":2,"n_instrument":5,"n_unverified":3,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":9,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 5 where Syntology's instrument failed) · 3 unverified","sample_list":"/paper/how-much-can-clip-benefit-vision-and-language#ran","syntology_url":"https://syntology.ai/paper/2107.06383","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2107.06383"}},"official":{"repos":["clip-vil/CLIP-ViL"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":["listed"]}}},{"url":"/paper/visitron-visual-semantics-aligned","slug":"visitron-visual-semantics-aligned","title":"VISITRON: Visual Semantics-Aligned Interactively Trained Object-Navigator","date":"2021-05-25","arxiv_id":"2105.11589","repositories_listed":1,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":1,"n_instrument":2,"n_unverified":0,"n_honours":1,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/visitron-visual-semantics-aligned#ran","syntology_url":"https://syntology.ai/paper/2105.11589","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2105.11589"}},"official":{"repos":["alexa/visitron"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/pathdreamer-a-world-model-for-indoor","slug":"pathdreamer-a-world-model-for-indoor","title":"Pathdreamer: A World Model for Indoor Navigation","date":"2021-05-18","arxiv_id":"2105.08756","repositories_listed":1,"syntology":{"n":7,"n_ran":7,"n_constructed":0,"n_ran_checked":7,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":7,"n_pointer_only":0,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/pathdreamer-a-world-model-for-indoor#ran","syntology_url":"https://syntology.ai/paper/2105.08756","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2105.08756"}},"official":{"repos":["google-research/pathdreamer"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":7,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/episodic-transformer-for-vision-and-language","slug":"episodic-transformer-for-vision-and-language","title":"Episodic Transformer for Vision-and-Language Navigation","date":"2021-05-13","arxiv_id":"2105.06453","repositories_listed":1,"syntology":{"n":6,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":4,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","sample_list":"/paper/episodic-transformer-for-vision-and-language#ran","syntology_url":"https://syntology.ai/paper/2105.06453","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2105.06453"}},"official":{"repos":["alexpashevich/E.T."],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":4,"ran_from_kinds":["official"]}}},{"url":"/paper/room-across-room-multilingual-vision-and","slug":"room-across-room-multilingual-vision-and","title":"Room-Across-Room: Multilingual Vision-and-Language Navigation with Dense Spatiotemporal Grounding","date":"2020-10-15","arxiv_id":"2010.07954","repositories_listed":3,"syntology":{"n":4,"n_ran":3,"n_constructed":0,"n_ran_checked":3,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":3,"n_pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/room-across-room-multilingual-vision-and#ran","syntology_url":"https://syntology.ai/paper/2010.07954","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2010.07954"}},"official":{"repos":["google-research-datasets/RxR"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"url":"/paper/babywalk-going-farther-in-vision-and-language","slug":"babywalk-going-farther-in-vision-and-language","title":"BabyWalk: Going Farther in Vision-and-Language Navigation by Taking Baby Steps","date":"2020-05-10","arxiv_id":"2005.04625","repositories_listed":1,"syntology":{"n":5,"n_ran":4,"n_constructed":1,"n_ran_checked":1,"n_instrument":3,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":0,"phrase":"4 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 3 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/babywalk-going-farther-in-vision-and-language#ran","syntology_url":"https://syntology.ai/paper/2005.04625","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2005.04625"}},"official":{"repos":["Sha-Lab/babywalk"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/sub-instruction-aware-vision-and-language","slug":"sub-instruction-aware-vision-and-language","title":"Sub-Instruction Aware Vision-and-Language Navigation","date":"2020-04-06","arxiv_id":"2004.02707","repositories_listed":1,"syntology":{"n":2,"n_ran":2,"n_constructed":0,"n_ran_checked":0,"n_instrument":2,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/sub-instruction-aware-vision-and-language#ran","syntology_url":"https://syntology.ai/paper/2004.02707","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2004.02707"}},"official":{"repos":["YicongHong/Fine-Grained-R2R"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/towards-learning-a-generic-agent-for-vision","slug":"towards-learning-a-generic-agent-for-vision","title":"Towards Learning a Generic Agent for Vision-and-Language Navigation via Pre-training","date":"2020-02-25","arxiv_id":"2002.10638","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":1,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/towards-learning-a-generic-agent-for-vision#ran","syntology_url":"https://syntology.ai/paper/2002.10638","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2002.10638"}},"official":{"repos":["weituo12321/PREVALENT"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/retouchdown-adding-touchdown-to-streetlearn","slug":"retouchdown-adding-touchdown-to-streetlearn","title":"Retouchdown: Adding Touchdown to StreetLearn as a Shareable Resource for Language Grounding Tasks in Street View","date":"2020-01-10","arxiv_id":"2001.03671","repositories_listed":4,"syntology":{"n":8,"n_ran":8,"n_constructed":0,"n_ran_checked":7,"n_instrument":1,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":7,"n_pointer_only":1,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/retouchdown-adding-touchdown-to-streetlearn#ran","syntology_url":"https://syntology.ai/paper/2001.03671","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2001.03671"}},"official":{"repos":["google-research/valan"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"url":"/paper/perceive-transform-and-act-multi-modal","slug":"perceive-transform-and-act-multi-modal","title":"Multimodal Attention Networks for Low-Level Vision-and-Language Navigation","date":"2019-11-27","arxiv_id":"1911.12377","repositories_listed":1,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":3,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":3,"n_pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/perceive-transform-and-act-multi-modal#ran","syntology_url":"https://syntology.ai/paper/1911.12377","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1911.12377"}},"official":{"repos":["aimagelab/perceive-transform-and-act"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/chasing-ghosts-instruction-following-as","slug":"chasing-ghosts-instruction-following-as","title":"Chasing Ghosts: Instruction Following as Bayesian State Tracking","date":"2019-07-03","arxiv_id":"1907.02022","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":0,"n_instrument":1,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/chasing-ghosts-instruction-following-as#ran","syntology_url":"https://syntology.ai/paper/1907.02022","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1907.02022"}},"official":{"repos":["batra-mlp-lab/vln-chasing-ghosts"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/rerere-remote-embodied-referring-expressions","slug":"rerere-remote-embodied-referring-expressions","title":"REVERIE: Remote Embodied Visual Referring Expression in Real Indoor Environments","date":"2019-04-23","arxiv_id":"1904.10151","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":0,"n_instrument":1,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/rerere-remote-embodied-referring-expressions#ran","syntology_url":"https://syntology.ai/paper/1904.10151","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1904.10151"}},"official":null}},{"url":"/paper/tactical-rewind-self-correction-via","slug":"tactical-rewind-self-correction-via","title":"Tactical Rewind: Self-Correction via Backtracking in Vision-and-Language Navigation","date":"2019-03-06","arxiv_id":"1903.02547","repositories_listed":1,"syntology":{"n":4,"n_ran":4,"n_constructed":0,"n_ran_checked":1,"n_instrument":3,"n_unverified":0,"n_honours":1,"n_violates":0,"n_no_contract":0,"n_pointer_only":4,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 3 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/tactical-rewind-self-correction-via#ran","syntology_url":"https://syntology.ai/paper/1903.02547","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1903.02547"}},"official":{"repos":["Kelym/FAST"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/the-regretful-agent-heuristic-aided","slug":"the-regretful-agent-heuristic-aided","title":"The Regretful Agent: Heuristic-Aided Navigation through Progress Estimation","date":"2019-03-05","arxiv_id":"1903.01602","repositories_listed":3,"syntology":{"n":11,"n_ran":10,"n_constructed":0,"n_ran_checked":10,"n_instrument":0,"n_unverified":1,"n_honours":1,"n_violates":0,"n_no_contract":9,"n_pointer_only":0,"phrase":"10 ran (of which 0 constructed an object rather than computing a result; 10 with no instrument failure: 1 honoured, 0 violated, 9 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/the-regretful-agent-heuristic-aided#ran","syntology_url":"https://syntology.ai/paper/1903.01602","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1903.01602"}},"official":{"repos":["chihyaoma/regretful-agent"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["listed","official"]}}},{"url":"/paper/self-monitoring-navigation-agent-via","slug":"self-monitoring-navigation-agent-via","title":"Self-Monitoring Navigation Agent via Auxiliary Progress Estimation","date":"2019-01-10","arxiv_id":"1901.03035","repositories_listed":2,"syntology":{"n":6,"n_ran":5,"n_constructed":0,"n_ran_checked":5,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":5,"n_pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/self-monitoring-navigation-agent-via#ran","syntology_url":"https://syntology.ai/paper/1901.03035","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1901.03035"}},"official":{"repos":["chihyaoma/selfmonitoring-agent"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/speaker-follower-models-for-vision-and","slug":"speaker-follower-models-for-vision-and","title":"Speaker-Follower Models for Vision-and-Language Navigation","date":"2018-06-07","arxiv_id":"1806.02724","repositories_listed":1,"syntology":{"n":2,"n_ran":2,"n_constructed":0,"n_ran_checked":0,"n_instrument":2,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/speaker-follower-models-for-vision-and#ran","syntology_url":"https://syntology.ai/paper/1806.02724","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1806.02724"}},"official":null}}],"record_sha256":"376eee18f32701fe4abdfceac7237130b491a252a9ec0105a7bfb068b5bf8436","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}