{"url":"/dataset/condensed-movies","name":"Condensed Movies","full_name":null,"description_markdown":"A large-scale video dataset, featuring clips from movies with detailed captions.","description_withheld":null,"homepage":"https://www.robots.ox.ac.uk/~vgg/data/condensed-movies/","introduced_date":"2020-05-08","introduced_date_note":null,"introduced_by":{"paper":"/paper/condensed-movies-story-based-retrieval-with","title":"Condensed Movies: Story Based Retrieval with Contextual Embeddings","first_author":"Max Bain","url":null},"license":null,"modalities":[{"name":"Videos","url":"/datasets/modality/videos"},{"name":"Texts","url":"/datasets/modality/texts"}],"tasks":[{"name":"Video Retrieval","url":"/task/video-retrieval","datasets_with_task":"/datasets/task/video-retrieval"},{"name":"Text Retrieval","url":"/task/text-retrieval","datasets_with_task":"/datasets/task/text-retrieval"}],"languages":[],"variants":["Condensed Movies"],"data_loaders":[],"num_papers_in_archive":15,"source":{"archive":"pwc-archive (Hugging Face), CC BY-SA 4.0","snapshot":"2025-07-28"},"benchmarks":[{"leaderboard":"/sota/video-retrieval-on-condensed-movies","task":"Video Retrieval","dataset_variant":"Condensed Movies","rows":3,"metrics":["text-to-video R@1","text-to-video R@5","text-to-video R@10"],"first_row_in_archive_order":{"model":"TESTA (ViT-B/16)","paper":"/paper/testa-temporal-spatial-token-aggregation-for","metrics":{"text-to-video R@1":"24.9","text-to-video R@10":"55.1","text-to-video R@5":"46.5"},"code_links":[{"title":"renshuhuai-andy/testa","url":"https://github.com/renshuhuai-andy/testa"}]},"note":"rows are the archive's own order at snapshot; nothing here re-ranks them"}],"papers_with_a_benchmark_row":[{"paper":"/paper/testa-temporal-spatial-token-aggregation-for","title":"TESTA: Temporal-Spatial Token Aggregation for Long-form Video-Language Understanding","date":"2023-10-29","rows_on_this_dataset":1,"code_links":1,"syntology":{"read_at":"2026-09-24T18:15:14+00:00","samples_harvested":15,"samples_ran":11,"samples_unverified":4,"pointer_only_for_licence":0,"claim":"Per-sample execution on synthesized fixtures; not a correctness claim."}},{"paper":"/paper/vindlu-a-recipe-for-effective-video-and","title":"VindLU: A Recipe for Effective Video-and-Language Pretraining","date":"2022-12-09","rows_on_this_dataset":1,"code_links":1,"syntology":null},{"paper":"/paper/long-form-video-language-pre-training-with","title":"Long-Form Video-Language Pre-Training with Multimodal Temporal Contrastive Learning","date":"2022-10-12","rows_on_this_dataset":1,"code_links":1,"syntology":{"read_at":"2026-09-24T18:15:14+00:00","samples_harvested":1,"samples_ran":1,"samples_unverified":0,"pointer_only_for_licence":1,"claim":"Per-sample execution on synthesized fixtures; not a correctness claim."}}],"syntology_totals":{"read_at":"2026-09-24T18:15:14+00:00","papers_with_samples":2,"samples_harvested":16,"samples_ran":12,"samples_unverified":4,"pointer_only_for_licence":1,"papers_with_no_sample_that_ran":0,"note":"the per-paper counts above, summed; not a rate"},"papers_note":"The archive never published its papers-using-dataset list; these are papers with a leaderboard row on this dataset's benchmarks."}