{"about":{"site":"https://codewithpapers.app","non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page"},"url":"/paper/simpleds-a-simple-deep-reinforcement-learning","title":"SimpleDS: A Simple Deep Reinforcement Learning Dialogue System","arxiv_id":"1601.04574","date":"2016-01-18","proceeding":null,"authors":["Heriberto Cuayáhuitl"],"abstract":"This paper presents 'SimpleDS', a simple and publicly available dialogue\nsystem trained with deep reinforcement learning. In contrast to previous\nreinforcement learning dialogue systems, this system avoids manual feature\nengineering by performing action selection directly from raw text of the last\nsystem and (noisy) user responses. Our initial results, in the restaurant\ndomain, show that it is indeed possible to induce reasonable dialogue behaviour\nwith an approach that aims for high levels of automation in dialogue control\nfor intelligent interactive agents.","url_abs":"http://arxiv.org/abs/1601.04574v1","url_pdf":"http://arxiv.org/pdf/1601.04574v1.pdf","source":{"archive":"pwc-archive (Hugging Face), CC BY-SA 4.0","snapshot":"2025-07-28","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","row_kind":"abstracts"},"code_links":[{"paper_slug":"simpleds-a-simple-deep-reinforcement-learning","repo_url":"https://github.com/cuayahuitl/SimpleDS","is_official":1,"mentioned_in_paper":1,"mentioned_in_github":1,"framework":"none","reach":{"status":"ok","spdx":"GPL-3.0"}}],"tasks":[{"task_slug":"deep-reinforcement-learning","task_name":"Deep Reinforcement Learning"},{"task_slug":"feature-engineering","task_name":"Feature Engineering"},{"task_slug":"reinforcement-learning","task_name":"Reinforcement Learning"},{"task_slug":"reinforcement-learning-1","task_name":"Reinforcement Learning (RL)"},{"task_slug":"reinforcement-learning-2","task_name":"reinforcement-learning"}],"methods":[],"datasets_introduced":[],"methods_introduced":[],"results":[],"syntology":{"atlas_url":null,"mcp":null,"developers":"https://syntology.ai/developers"},"arxiv_metadata":null,"syntology_extracted_results":null}