{"about":{"site":"https://codewithpapers.app","non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page"},"url":"/paper/audio-visual-scene-aware-dialog-avsd","title":"Audio Visual Scene-Aware Dialog (AVSD) Challenge at DSTC7","arxiv_id":"1806.00525","date":"2018-06-01","proceeding":null,"authors":["Huda Alamri","Vincent Cartillier","Raphael Gontijo Lopes","Abhishek Das","Jue Wang","Irfan Essa","Dhruv Batra","Devi Parikh","Anoop Cherian","Tim K. Marks","Chiori Hori"],"abstract":"Scene-aware dialog systems will be able to have conversations with users\nabout the objects and events around them. Progress on such systems can be made\nby integrating state-of-the-art technologies from multiple research areas\nincluding end-to-end dialog systems visual dialog, and video description. We\nintroduce the Audio Visual Scene Aware Dialog (AVSD) challenge and dataset. In\nthis challenge, which is one track of the 7th Dialog System Technology\nChallenges (DSTC7) workshop1, the task is to build a system that generates\nresponses in a dialog about an input video","url_abs":"http://arxiv.org/abs/1806.00525v1","url_pdf":"http://arxiv.org/pdf/1806.00525v1.pdf","source":{"archive":"pwc-archive (Hugging Face), CC BY-SA 4.0","snapshot":"2025-07-28","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","row_kind":"abstracts"},"code_links":[{"paper_slug":"audio-visual-scene-aware-dialog-avsd","repo_url":"https://github.com/batra-mlp-lab/avsd","is_official":0,"mentioned_in_paper":0,"mentioned_in_github":1,"framework":"pytorch","reach":{"status":"ok"}},{"paper_slug":"audio-visual-scene-aware-dialog-avsd","repo_url":"https://github.com/hudaAlamri/DSTC7-Audio-Visual-Scene-Aware-Dialog-AVSD-Challenge","is_official":0,"mentioned_in_paper":0,"mentioned_in_github":1,"framework":"pytorch","reach":{"status":"ok","spdx":"MIT"}},{"paper_slug":"audio-visual-scene-aware-dialog-avsd","repo_url":"https://github.com/lemuria-wchen/DialogVED","is_official":0,"mentioned_in_paper":0,"mentioned_in_github":1,"framework":"pytorch","reach":{"status":"ok"}},{"paper_slug":"audio-visual-scene-aware-dialog-avsd","repo_url":"https://github.com/vincentcartillier/AVSD","is_official":0,"mentioned_in_paper":0,"mentioned_in_github":1,"framework":"pytorch","reach":{"status":"ok"}}],"tasks":[{"task_slug":"video-description","task_name":"Video Description"},{"task_slug":"visual-dialogue","task_name":"Visual Dialog"}],"methods":[],"datasets_introduced":[{"slug":"avsd","name":"AVSD","full_name":"Audio-Visual Scene-Aware Dialog"}],"methods_introduced":[],"results":[],"syntology":{"syntology_url":null,"atlas_url":"https://app.syntology.ai/?focus=1806.00525","mcp":null,"developers":"https://syntology.ai/developers"},"arxiv_metadata":null,"syntology_extracted_results":null}