{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/task/video-instance-segmentation/papers/ran/1","list_of":"/task/video-instance-segmentation","task":"Video Instance Segmentation","archive":{"snapshot":"2025-07-28"},"key_notes":{"n_ran_checked":"legacy name, kept unchanged so existing readers do not break: it counts the samples that ran with no instrument failure (honoured, violated, and ran with no contract checked); it does not mean a contract was checked, and the pages print it as 'K with no instrument failure', not 'K checked'","n_constructed":"a sub-count of the samples that ran, never subtracted from them and never a failure: an executed sample whose run returned an instance of its own class (fixture_out_type equals the entry name): the run built an object and did not compute a result (Syntology's RAN record, counts.constructed)"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"ran","order_definition":"only papers where Syntology ran at least one harvested sample; date (newest first), ties by arXiv id","caption":"We ran code from the paper's repository; we did not run it on this task or check it against the task's benchmarks.","absence":"A paper missing from this list is not a recorded non-run: it may have no arXiv id, no harvested code, or only samples that have not run yet.","page":1,"pages_in_order":1,"rows_per_page":100,"rows":[1,30],"of":30,"counts":{"archive_papers_tagged":148,"with_a_code_link":94,"where_syntology_ran_a_sample":30,"not_listed_spam_title":0,"listed":148,"listed_where_code_ran":30,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":28,"every_run_a_failure_of_syntologys_instrument":2,"listed_with_a_run_with_no_instrument_failure":28,"listed_every_run_a_failure_of_syntologys_instrument":2,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/task/video-instance-segmentation/papers/ran/1","prev":null,"next":null,"papers":[{"url":"/paper/dvis-daq-improving-video-segmentation-via","slug":"dvis-daq-improving-video-segmentation-via","title":"DVIS-DAQ: Improving Video Segmentation via Dynamic Anchor Queries","date":"2024-03-29","arxiv_id":"2404.00086","repositories_listed":3,"syntology":{"n":5,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":3,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":3,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","sample_list":"/paper/dvis-daq-improving-video-segmentation-via#ran","syntology_url":"https://syntology.ai/paper/2404.00086","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2404.00086"}},"official":{"repos":["skyworkai/daq-vs"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"url":"/paper/univs-unified-and-universal-video","slug":"univs-unified-and-universal-video","title":"UniVS: Unified and Universal Video Segmentation with Prompts as Queries","date":"2024-02-28","arxiv_id":"2402.18115","repositories_listed":1,"syntology":{"n":14,"n_ran":12,"n_constructed":0,"n_ran_checked":12,"n_instrument":0,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":12,"n_pointer_only":14,"phrase":"12 ran (of which 0 constructed an object rather than computing a result; 12 with no instrument failure: 0 honoured, 0 violated, 12 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/univs-unified-and-universal-video#ran","syntology_url":"https://syntology.ai/paper/2402.18115","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2402.18115"}},"official":{"repos":["minghanli/univs"],"state":"official (archive's flag): 12 ran","n_ran":12,"n_constructed":0,"n_ran_no_instrument_failure":12,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/general-object-foundation-model-for-images","slug":"general-object-foundation-model-for-images","title":"General Object Foundation Model for Images and Videos at Scale","date":"2023-12-14","arxiv_id":"2312.09158","repositories_listed":1,"syntology":{"n":13,"n_ran":11,"n_constructed":0,"n_ran_checked":9,"n_instrument":2,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":9,"n_pointer_only":3,"phrase":"11 ran (of which 0 constructed an object rather than computing a result; 9 with no instrument failure: 0 honoured, 0 violated, 9 with no contract checked; 2 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/general-object-foundation-model-for-images#ran","syntology_url":"https://syntology.ai/paper/2312.09158","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2312.09158"}},"official":{"repos":["FoundationVision/GLEE"],"state":"official (archive's flag): 11 ran","n_ran":11,"n_constructed":0,"n_ran_no_instrument_failure":9,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/tmt-vis-taxonomy-aware-multi-dataset-joint-1","slug":"tmt-vis-taxonomy-aware-multi-dataset-joint-1","title":"TMT-VIS: Taxonomy-aware Multi-dataset Joint Training for Video Instance Segmentation","date":"2023-12-11","arxiv_id":"2312.06630","repositories_listed":1,"syntology":{"n":7,"n_ran":6,"n_constructed":0,"n_ran_checked":6,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":6,"n_pointer_only":7,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/tmt-vis-taxonomy-aware-multi-dataset-joint-1#ran","syntology_url":"https://syntology.ai/paper/2312.06630","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2312.06630"}},"official":{"repos":["rkzheng99/tmt-vis"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/videocutler-surprisingly-simple-unsupervised","slug":"videocutler-surprisingly-simple-unsupervised","title":"VideoCutLER: Surprisingly Simple Unsupervised Video Instance Segmentation","date":"2023-08-28","arxiv_id":"2308.14710","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":0,"n_instrument":1,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/videocutler-surprisingly-simple-unsupervised#ran","syntology_url":"https://syntology.ai/paper/2308.14710","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2308.14710"}},"official":{"repos":["facebookresearch/cutler"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/ctvis-consistent-training-for-online-video","slug":"ctvis-consistent-training-for-online-video","title":"CTVIS: Consistent Training for Online Video Instance Segmentation","date":"2023-07-24","arxiv_id":"2307.12616","repositories_listed":1,"syntology":{"n":7,"n_ran":4,"n_constructed":0,"n_ran_checked":4,"n_instrument":0,"n_unverified":3,"n_honours":3,"n_violates":0,"n_no_contract":1,"n_pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 3 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","sample_list":"/paper/ctvis-consistent-training-for-online-video#ran","syntology_url":"https://syntology.ai/paper/2307.12616","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2307.12616"}},"official":{"repos":["kainingying/ctvis"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":3,"ran_from_kinds":["official"]}}},{"url":"/paper/dq-det-learning-dynamic-query-combinations","slug":"dq-det-learning-dynamic-query-combinations","title":"Learning Dynamic Query Combinations for Transformer-based Object Detection and Segmentation","date":"2023-07-23","arxiv_id":"2307.12239","repositories_listed":1,"syntology":{"n":5,"n_ran":3,"n_constructed":0,"n_ran_checked":2,"n_instrument":1,"n_unverified":2,"n_honours":0,"n_violates":1,"n_no_contract":1,"n_pointer_only":2,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 1 violated, 1 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/dq-det-learning-dynamic-query-combinations#ran","syntology_url":"https://syntology.ai/paper/2307.12239","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2307.12239"}},"official":{"repos":["bytedance/dq-det"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/segment-anything-meets-point-tracking","slug":"segment-anything-meets-point-tracking","title":"Segment Anything Meets Point Tracking","date":"2023-07-03","arxiv_id":"2307.01197","repositories_listed":1,"syntology":{"n":12,"n_ran":7,"n_constructed":0,"n_ran_checked":7,"n_instrument":0,"n_unverified":5,"n_honours":0,"n_violates":0,"n_no_contract":7,"n_pointer_only":1,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 0 where Syntology's instrument failed) · 5 unverified","sample_list":"/paper/segment-anything-meets-point-tracking#ran","syntology_url":"https://syntology.ai/paper/2307.01197","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2307.01197"}},"official":{"repos":["syscv/sam-pt"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":7,"n_unverified":5,"ran_from_kinds":["official"]}}},{"url":"/paper/mask-free-video-instance-segmentation","slug":"mask-free-video-instance-segmentation","title":"Mask-Free Video Instance Segmentation","date":"2023-03-28","arxiv_id":"2303.15904","repositories_listed":1,"syntology":{"n":8,"n_ran":8,"n_constructed":0,"n_ran_checked":8,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":8,"n_pointer_only":0,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 0 violated, 8 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/mask-free-video-instance-segmentation#ran","syntology_url":"https://syntology.ai/paper/2303.15904","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2303.15904"}},"official":{"repos":["syscv/maskfreevis"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/mdqe-mining-discriminative-query-embeddings","slug":"mdqe-mining-discriminative-query-embeddings","title":"MDQE: Mining Discriminative Query Embeddings to Segment Occluded Instances on Challenging Videos","date":"2023-03-25","arxiv_id":"2303.14395","repositories_listed":1,"syntology":{"n":4,"n_ran":3,"n_constructed":2,"n_ran_checked":2,"n_instrument":1,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":4,"phrase":"3 ran (of which 2 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/mdqe-mining-discriminative-query-embeddings#ran","syntology_url":"https://syntology.ai/paper/2303.14395","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2303.14395"}},"official":{"repos":["minghanli/mdqe_cvpr2023"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":2,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/universal-instance-perception-as-object","slug":"universal-instance-perception-as-object","title":"Universal Instance Perception as Object Discovery and Retrieval","date":"2023-03-12","arxiv_id":"2303.06674","repositories_listed":1,"syntology":{"n":4,"n_ran":4,"n_constructed":0,"n_ran_checked":2,"n_instrument":2,"n_unverified":0,"n_honours":1,"n_violates":0,"n_no_contract":1,"n_pointer_only":2,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 1 honoured, 0 violated, 1 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/universal-instance-perception-as-object#ran","syntology_url":"https://syntology.ai/paper/2303.06674","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2303.06674"}},"official":{"repos":["MasterBin-IIAU/UNINEXT"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/tarvis-a-unified-approach-for-target-based","slug":"tarvis-a-unified-approach-for-target-based","title":"TarViS: A Unified Approach for Target-based Video Segmentation","date":"2023-01-06","arxiv_id":"2301.02657","repositories_listed":1,"syntology":{"n":6,"n_ran":4,"n_constructed":0,"n_ran_checked":4,"n_instrument":0,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":4,"n_pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/tarvis-a-unified-approach-for-target-based#ran","syntology_url":"https://syntology.ai/paper/2301.02657","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2301.02657"}},"official":{"repos":["Ali2500/TarViS"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/video-mask-transfiner-for-high-quality-video","slug":"video-mask-transfiner-for-high-quality-video","title":"Video Mask Transfiner for High-Quality Video Instance Segmentation","date":"2022-07-28","arxiv_id":"2207.14012","repositories_listed":1,"syntology":{"n":7,"n_ran":5,"n_constructed":0,"n_ran_checked":2,"n_instrument":3,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 3 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/video-mask-transfiner-for-high-quality-video#ran","syntology_url":"https://syntology.ai/paper/2207.14012","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2207.14012"}},"official":null}},{"url":"/paper/devis-making-deformable-transformers-work-for","slug":"devis-making-deformable-transformers-work-for","title":"DeVIS: Making Deformable Transformers Work for Video Instance Segmentation","date":"2022-07-22","arxiv_id":"2207.11103","repositories_listed":1,"syntology":{"n":11,"n_ran":9,"n_constructed":0,"n_ran_checked":8,"n_instrument":1,"n_unverified":2,"n_honours":1,"n_violates":0,"n_no_contract":7,"n_pointer_only":11,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 1 honoured, 0 violated, 7 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/devis-making-deformable-transformers-work-for#ran","syntology_url":"https://syntology.ai/paper/2207.11103","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2207.11103"}},"official":{"repos":["acaelles97/devis"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/in-defense-of-online-models-for-video","slug":"in-defense-of-online-models-for-video","title":"In Defense of Online Models for Video Instance Segmentation","date":"2022-07-21","arxiv_id":"2207.10661","repositories_listed":2,"syntology":{"n":5,"n_ran":3,"n_constructed":0,"n_ran_checked":0,"n_instrument":3,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 3 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/in-defense-of-online-models-for-video#ran","syntology_url":"https://syntology.ai/paper/2207.10661","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2207.10661"}},"official":{"repos":["wjf5203/vnext"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"url":"/paper/vita-video-instance-segmentation-via-object","slug":"vita-video-instance-segmentation-via-object","title":"VITA: Video Instance Segmentation via Object Token Association","date":"2022-06-09","arxiv_id":"2206.04403","repositories_listed":1,"syntology":{"n":7,"n_ran":6,"n_constructed":0,"n_ran_checked":6,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":6,"n_pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/vita-video-instance-segmentation-via-object#ran","syntology_url":"https://syntology.ai/paper/2206.04403","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2206.04403"}},"official":{"repos":["sukjunhwang/vita"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/temporally-efficient-vision-transformer-for","slug":"temporally-efficient-vision-transformer-for","title":"Temporally Efficient Vision Transformer for Video Instance Segmentation","date":"2022-04-18","arxiv_id":"2204.08412","repositories_listed":3,"syntology":{"n":2,"n_ran":1,"n_constructed":0,"n_ran_checked":1,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/temporally-efficient-vision-transformer-for#ran","syntology_url":"https://syntology.ai/paper/2204.08412","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2204.08412"}},"official":{"repos":["hustvl/tevit"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/mlseg-image-and-video-segmentation-as-multi","slug":"mlseg-image-and-video-segmentation-as-multi","title":"RankSeg: Adaptive Pixel Classification with Image Category Ranking for Segmentation","date":"2022-03-08","arxiv_id":"2203.04187","repositories_listed":2,"syntology":{"n":11,"n_ran":9,"n_constructed":0,"n_ran_checked":7,"n_instrument":2,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":7,"n_pointer_only":0,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 2 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/mlseg-image-and-video-segmentation-as-multi#ran","syntology_url":"https://syntology.ai/paper/2203.04187","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2203.04187"}},"official":{"repos":["openseg-group/mlseg","openseg-group/rankseg"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":1,"ran_from_kinds":["found_in_text","official"]}}},{"url":"/paper/language-as-queries-for-referring-video","slug":"language-as-queries-for-referring-video","title":"Language as Queries for Referring Video Object Segmentation","date":"2022-01-03","arxiv_id":"2201.00487","repositories_listed":1,"syntology":{"n":8,"n_ran":7,"n_constructed":0,"n_ran_checked":5,"n_instrument":2,"n_unverified":1,"n_honours":1,"n_violates":0,"n_no_contract":4,"n_pointer_only":8,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 1 honoured, 0 violated, 4 with no contract checked; 2 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/language-as-queries-for-referring-video#ran","syntology_url":"https://syntology.ai/paper/2201.00487","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2201.00487"}},"official":{"repos":["wjn922/referformer"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/mask2former-for-video-instance-segmentation","slug":"mask2former-for-video-instance-segmentation","title":"Mask2Former for Video Instance Segmentation","date":"2021-12-20","arxiv_id":"2112.10764","repositories_listed":6,"syntology":{"n":7,"n_ran":6,"n_constructed":0,"n_ran_checked":4,"n_instrument":2,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":4,"n_pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 2 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/mask2former-for-video-instance-segmentation#ran","syntology_url":"https://syntology.ai/paper/2112.10764","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2112.10764"}},"official":{"repos":["facebookresearch/Mask2Former"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"url":"/paper/visolo-grid-based-space-time-aggregation-for","slug":"visolo-grid-based-space-time-aggregation-for","title":"VISOLO: Grid-Based Space-Time Aggregation for Efficient Online Video Instance Segmentation","date":"2021-12-08","arxiv_id":"2112.04177","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":1,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/visolo-grid-based-space-time-aggregation-for#ran","syntology_url":"https://syntology.ai/paper/2112.04177","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2112.04177"}},"official":{"repos":["suhohan95/visolo"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/do-different-tracking-tasks-require-different","slug":"do-different-tracking-tasks-require-different","title":"Do Different Tracking Tasks Require Different Appearance Models?","date":"2021-07-05","arxiv_id":"2107.02156","repositories_listed":1,"syntology":{"n":17,"n_ran":6,"n_constructed":4,"n_ran_checked":4,"n_instrument":2,"n_unverified":11,"n_honours":0,"n_violates":0,"n_no_contract":4,"n_pointer_only":0,"phrase":"6 ran (of which 4 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 2 where Syntology's instrument failed) · 11 unverified","sample_list":"/paper/do-different-tracking-tasks-require-different#ran","syntology_url":"https://syntology.ai/paper/2107.02156","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2107.02156"}},"official":{"repos":["Zhongdao/UniTrack"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":4,"n_ran_no_instrument_failure":4,"n_unverified":11,"ran_from_kinds":["official"]}}},{"url":"/paper/prototypical-cross-attention-networks-for","slug":"prototypical-cross-attention-networks-for","title":"Prototypical Cross-Attention Networks for Multiple Object Tracking and Segmentation","date":"2021-06-22","arxiv_id":"2106.11958","repositories_listed":1,"syntology":{"n":5,"n_ran":5,"n_constructed":0,"n_ran_checked":5,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":5,"n_pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/prototypical-cross-attention-networks-for#ran","syntology_url":"https://syntology.ai/paper/2106.11958","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2106.11958"}},"official":{"repos":["SysCV/pcan"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/video-instance-segmentation-using-inter-frame","slug":"video-instance-segmentation-using-inter-frame","title":"Video Instance Segmentation using Inter-Frame Communication Transformers","date":"2021-06-07","arxiv_id":"2106.03299","repositories_listed":1,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":3,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":1,"n_no_contract":2,"n_pointer_only":2,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 1 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/video-instance-segmentation-using-inter-frame#ran","syntology_url":"https://syntology.ai/paper/2106.03299","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2106.03299"}},"official":{"repos":["sukjunhwang/IFC"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/queryinst-parallelly-supervised-mask-query","slug":"queryinst-parallelly-supervised-mask-query","title":"Instances as Queries","date":"2021-05-05","arxiv_id":"2105.01928","repositories_listed":5,"syntology":{"n":3,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":1,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/queryinst-parallelly-supervised-mask-query#ran","syntology_url":"https://syntology.ai/paper/2105.01928","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2105.01928"}},"official":{"repos":["hustvl/QueryInst"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"url":"/paper/video-instance-segmentation-with-a-propose","slug":"video-instance-segmentation-with-a-propose","title":"Video Instance Segmentation with a Propose-Reduce Paradigm","date":"2021-03-25","arxiv_id":"2103.13746","repositories_listed":1,"syntology":{"n":5,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":3,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":5,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","sample_list":"/paper/video-instance-segmentation-with-a-propose#ran","syntology_url":"https://syntology.ai/paper/2103.13746","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2103.13746"}},"official":{"repos":["dvlab-research/proposereduce"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":3,"ran_from_kinds":["official"]}}},{"url":"/paper/track-to-detect-and-segment-an-online-multi","slug":"track-to-detect-and-segment-an-online-multi","title":"Track to Detect and Segment: An Online Multi-Object Tracker","date":"2021-03-16","arxiv_id":"2103.08808","repositories_listed":1,"syntology":{"n":4,"n_ran":4,"n_constructed":0,"n_ran_checked":1,"n_instrument":3,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 3 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/track-to-detect-and-segment-an-online-multi#ran","syntology_url":"https://syntology.ai/paper/2103.08808","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2103.08808"}},"official":{"repos":["JialianW/TraDeS"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/learning-monocular-depth-in-dynamic-scenes","slug":"learning-monocular-depth-in-dynamic-scenes","title":"Learning Monocular Depth in Dynamic Scenes via Instance-Aware Projection Consistency","date":"2021-02-04","arxiv_id":"2102.02629","repositories_listed":1,"syntology":{"n":8,"n_ran":4,"n_constructed":0,"n_ran_checked":2,"n_instrument":2,"n_unverified":4,"n_honours":2,"n_violates":0,"n_no_contract":0,"n_pointer_only":8,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 2 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 4 unverified","sample_list":"/paper/learning-monocular-depth-in-dynamic-scenes#ran","syntology_url":"https://syntology.ai/paper/2102.02629","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2102.02629"}},"official":{"repos":["SeokjuLee/Insta-DM"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":4,"ran_from_kinds":["official"]}}},{"url":"/paper/video-instance-segmentation","slug":"video-instance-segmentation","title":"Video Instance Segmentation","date":"2019-05-12","arxiv_id":"1905.04804","repositories_listed":6,"syntology":{"n":11,"n_ran":9,"n_constructed":0,"n_ran_checked":8,"n_instrument":1,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":8,"n_pointer_only":1,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 0 violated, 8 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/video-instance-segmentation#ran","syntology_url":"https://syntology.ai/paper/1905.04804","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1905.04804"}},"official":{"repos":["Epiphqny/VisTR"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"url":"/paper/simple-online-and-realtime-tracking-with-a","slug":"simple-online-and-realtime-tracking-with-a","title":"Simple Online and Realtime Tracking with a Deep Association Metric","date":"2017-03-21","arxiv_id":"1703.07402","repositories_listed":75,"syntology":{"n":40,"n_ran":29,"n_constructed":0,"n_ran_checked":24,"n_instrument":5,"n_unverified":11,"n_honours":4,"n_violates":3,"n_no_contract":17,"n_pointer_only":9,"phrase":"29 ran (of which 0 constructed an object rather than computing a result; 24 with no instrument failure: 4 honoured, 3 violated, 17 with no contract checked; 5 where Syntology's instrument failed) · 11 unverified","sample_list":"/paper/simple-online-and-realtime-tracking-with-a#ran","syntology_url":"https://syntology.ai/paper/1703.07402","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1703.07402"}},"official":{"repos":["nwojke/deep_sort"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}}],"record_sha256":"c735bafa8ba10f7df465bb712b43451b32f42460c2e6bfbe41608c31af07f650","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}