{"url":"/dataset/segtrack-v2-1","name":"SegTrack-v2","full_name":null,"description_markdown":"SegTrack v2 is a video segmentation dataset with full pixel-level annotations on multiple objects at each frame within each video.\r\n\r\nSource: [https://web.engr.oregonstate.edu/~lif/SegTrack2/dataset.html](https://web.engr.oregonstate.edu/~lif/SegTrack2/dataset.html)\r\nImage Source: [https://www.researchgate.net/publication/325842926_Semantic_Video_Segmentation_A_Review_on_Recent_Approaches](https://www.researchgate.net/publication/325842926_Semantic_Video_Segmentation_A_Review_on_Recent_Approaches)","description_withheld":null,"homepage":"https://web.engr.oregonstate.edu/~lif/SegTrack2/dataset.html","introduced_date":"2013-01-01","introduced_date_note":null,"introduced_by":{"paper":null,"title":"Video Segmentation by Tracking Many Figure-Ground Segments","first_author":null,"url":"https://doi.org/10.1109/ICCV.2013.273"},"license":{"name":"Unknown","url":null},"modalities":[{"name":"Images","url":"/datasets/modality/images"},{"name":"Videos","url":"/datasets/modality/videos"}],"tasks":[{"name":"Semantic Segmentation","url":"/task/semantic-segmentation","datasets_with_task":"/datasets/task/semantic-segmentation"},{"name":"Video Object Segmentation","url":"/task/video-object-segmentation","datasets_with_task":"/datasets/task/video-object-segmentation"},{"name":"Video Salient Object Detection","url":"/task/video-salient-object-detection","datasets_with_task":"/datasets/task/video-salient-object-detection"},{"name":"Unsupervised Object Segmentation","url":"/task/unsupervised-object-segmentation","datasets_with_task":"/datasets/task/unsupervised-object-segmentation"},{"name":"Unsupervised Video Object Segmentation","url":"/task/unsupervised-video-object-segmentation","datasets_with_task":"/datasets/task/unsupervised-video-object-segmentation"},{"name":"Video Semantic Segmentation","url":"/task/video-semantic-segmentation","datasets_with_task":"/datasets/task/video-semantic-segmentation"},{"name":"Video Segmentation","url":"/task/video-segmentation","datasets_with_task":"/datasets/task/video-segmentation"}],"languages":[],"variants":["SegTrack v2","SegTrack-v2"],"data_loaders":[],"num_papers_in_archive":107,"source":{"archive":"pwc-archive (Hugging Face), CC BY-SA 4.0","snapshot":"2025-07-28"},"benchmarks":[{"leaderboard":"/sota/unsupervised-object-segmentation-on-segtrack","task":"Unsupervised Object Segmentation","dataset_variant":"SegTrack-v2","rows":8,"metrics":["mIoU"],"first_row_in_archive_order":{"model":"RCF (with post-processing)","paper":"/paper/bootstrapping-objectness-from-videos-by","metrics":{"mIoU":"79.6"},"code_links":[{"title":"TonyLianLong/RCF-UnsupVideoSeg","url":"https://github.com/TonyLianLong/RCF-UnsupVideoSeg"}]},"note":"rows are the archive's own order at snapshot; nothing here re-ranks them"},{"leaderboard":"/sota/video-salient-object-detection-on-segtrack-v2","task":"Video Salient Object Detection","dataset_variant":"SegTrack v2","rows":8,"metrics":["S-Measure","max E-measure","MAX F-MEASURE","AVERAGE MAE"],"first_row_in_archive_order":{"model":"UFO","paper":"/paper/a-unified-transformer-framework-for-group","metrics":{"AVERAGE MAE":"0.022","MAX F-MEASURE":"0.863","S-Measure":"0.892"},"code_links":[{"title":"suyukun666/UFO","url":"https://github.com/suyukun666/UFO"}]},"note":"rows are the archive's own order at snapshot; nothing here re-ranks them"},{"leaderboard":"/sota/unsupervised-video-object-segmentation-on-3","task":"Unsupervised Video Object Segmentation","dataset_variant":"SegTrack v2","rows":4,"metrics":["Mean IoU","Jaccard (Mean)"],"first_row_in_archive_order":{"model":"FrameSelect","paper":"/paper/mask-selection-and-propagation-for","metrics":{"Mean IoU":"72.2"},"code_links":[{"title":"vidit98/FrameSelect","url":"https://github.com/vidit98/FrameSelect"}]},"note":"rows are the archive's own order at snapshot; nothing here re-ranks them"},{"leaderboard":"/sota/video-object-segmentation-on-segtrack-v2-1","task":"Video Object Segmentation","dataset_variant":"SegTrack-v2","rows":1,"metrics":["mIoU"],"first_row_in_archive_order":{"model":"LOCATE","paper":"/paper/locate-self-supervised-object-discovery-via","metrics":{"mIoU":"79.9"},"code_links":[{"title":"silky1708/locate","url":"https://github.com/silky1708/locate"}]},"note":"rows are the archive's own order at snapshot; nothing here re-ranks them"},{"leaderboard":"/sota/video-segmentation-on-segtrack-v2","task":"Video Segmentation","dataset_variant":"SegTrack v2","rows":1,"metrics":["Accuracy"],"first_row_in_archive_order":{"model":"GDHF","paper":"/paper/geodesic-distance-histogram-feature-for-video","metrics":{"Accuracy":"86.86"},"code_links":[]},"note":"rows are the archive's own order at snapshot; nothing here re-ranks them"}],"papers_with_a_benchmark_row":[{"paper":"/paper/locate-self-supervised-object-discovery-via","title":"LOCATE: Self-supervised Object Discovery via Flow-guided Graph-cut and Bootstrapped Self-training","date":"2023-08-22","rows_on_this_dataset":1,"code_links":1,"syntology":{"read_at":"2026-09-24T18:15:14+00:00","samples_harvested":6,"samples_ran":5,"samples_unverified":1,"pointer_only_for_licence":6,"claim":"Per-sample execution on synthesized fixtures; not a correctness claim."}},{"paper":"/paper/bootstrapping-objectness-from-videos-by","title":"Bootstrapping Objectness from Videos by Relaxed Common Fate and Visual Grouping","date":"2023-04-17","rows_on_this_dataset":2,"code_links":1,"syntology":null},{"paper":"/paper/motion-inductive-self-supervised-object","title":"Motion-inductive Self-supervised Object Discovery in Videos","date":"2022-10-01","rows_on_this_dataset":1,"code_links":0,"syntology":null},{"paper":"/paper/a-simple-and-powerful-global-optimization-for","title":"A Simple and Powerful Global Optimization for Unsupervised Video Object Segmentation","date":"2022-09-19","rows_on_this_dataset":1,"code_links":1,"syntology":{"read_at":"2026-09-24T18:15:14+00:00","samples_harvested":10,"samples_ran":4,"samples_unverified":6,"pointer_only_for_licence":0,"claim":"Per-sample execution on synthesized fixtures; not a correctness claim."}},{"paper":"/paper/tokencut-segmenting-objects-in-images-and","title":"TokenCut: Segmenting Objects in Images and Videos with Self-supervised Transformer and Normalized Cut","date":"2022-09-01","rows_on_this_dataset":1,"code_links":0,"syntology":null},{"paper":"/paper/segmenting-moving-objects-via-an-object","title":"Segmenting Moving Objects via an Object-Centric Layered Representation","date":"2022-07-05","rows_on_this_dataset":1,"code_links":1,"syntology":{"read_at":"2026-09-24T18:15:14+00:00","samples_harvested":21,"samples_ran":16,"samples_unverified":5,"pointer_only_for_licence":0,"claim":"Per-sample execution on synthesized fixtures; not a correctness claim."}},{"paper":"/paper/guess-what-moves-unsupervised-video-and-image","title":"Guess What Moves: Unsupervised Video and Image Segmentation by Anticipating Motion","date":"2022-05-16","rows_on_this_dataset":1,"code_links":0,"syntology":null},{"paper":"/paper/a-unified-transformer-framework-for-group","title":"A Unified Transformer Framework for Group-based Segmentation: Co-Segmentation, Co-Saliency Detection and Video Salient Object Detection","date":"2022-03-09","rows_on_this_dataset":1,"code_links":1,"syntology":{"read_at":"2026-09-24T18:15:14+00:00","samples_harvested":13,"samples_ran":0,"samples_unverified":13,"pointer_only_for_licence":0,"claim":"Per-sample execution on synthesized fixtures; not a correctness claim."}},{"paper":"/paper/the-emergence-of-objectness-learning-zero","title":"The Emergence of Objectness: Learning Zero-Shot Segmentation from Videos","date":"2021-11-11","rows_on_this_dataset":1,"code_links":1,"syntology":null},{"paper":"/paper/mask-selection-and-propagation-for","title":"Mask Selection and Propagation for Unsupervised Video Object Segmentation","date":"2021-01-05","rows_on_this_dataset":1,"code_links":1,"syntology":null},{"paper":"/paper/exploiting-geometric-constraints-on-dense","title":"EpO-Net: Exploiting Geometric Constraints on Dense Trajectories for Motion Saliency","date":"2019-09-29","rows_on_this_dataset":1,"code_links":3,"syntology":null},{"paper":"/paper/time-masking-leveraging-temporal-information","title":"Time Masking: Leveraging Temporal Information in Spoken Dialogue Systems","date":"2019-07-25","rows_on_this_dataset":1,"code_links":0,"syntology":null},{"paper":"/paper/shifting-more-attention-to-video-salient","title":"Shifting More Attention to Video Salient Object Detection","date":"2019-06-01","rows_on_this_dataset":1,"code_links":1,"syntology":null},{"paper":"/paper/unsupervised-video-object-segmentation-with-1","title":"Unsupervised Video Object Segmentation with Motion-based Bilateral Networks","date":"2018-09-01","rows_on_this_dataset":1,"code_links":0,"syntology":null},{"paper":"/paper/pyramid-dilated-deeper-convlstm-for-video","title":"Pyramid Dilated Deeper ConvLSTM for Video Salient Object Detection","date":"2018-09-01","rows_on_this_dataset":1,"code_links":0,"syntology":null},{"paper":"/paper/learning-video-object-segmentation-with","title":"Learning Video Object Segmentation with Visual Memory","date":"2017-04-19","rows_on_this_dataset":1,"code_links":0,"syntology":null},{"paper":"/paper/geodesic-distance-histogram-feature-for-video","title":"Geodesic Distance Histogram Feature for Video Segmentation","date":"2017-03-31","rows_on_this_dataset":1,"code_links":0,"syntology":null},{"paper":"/paper/real-time-salient-object-detection-with-a","title":"Real-Time Salient Object Detection With a Minimum Spanning Tree","date":"2016-06-01","rows_on_this_dataset":1,"code_links":1,"syntology":null},{"paper":"/paper/minimum-barrier-salient-object-detection-at","title":"Minimum Barrier Salient Object Detection at 80 FPS","date":"2015-12-01","rows_on_this_dataset":1,"code_links":0,"syntology":null},{"paper":"/paper/saliency-aware-geodesic-video-object","title":"Saliency-Aware Geodesic Video Object Segmentation","date":"2015-06-01","rows_on_this_dataset":1,"code_links":0,"syntology":null}],"syntology_totals":{"read_at":"2026-09-24T18:15:14+00:00","papers_with_samples":4,"samples_harvested":50,"samples_ran":25,"samples_unverified":25,"pointer_only_for_licence":6,"papers_with_no_sample_that_ran":1,"note":"the per-paper counts above, summed; not a rate"},"papers_note":"The archive never published its papers-using-dataset list; these are papers with a leaderboard row on this dataset's benchmarks."}