{"url":"/task/conversational-response-selection","name":"Conversational Response Selection","slug":"conversational-response-selection","description_markdown":"Conversational response selection refers to the task of identifying the most relevant response to a given input sentence from a collection of sentences.","categories":[{"name":"Natural Language Processing","url":"/area/natural-language-processing"}],"source":{"archive":"pwc-archive (Hugging Face), CC BY-SA 4.0","snapshot":"2025-07-28","slug_source":"archive_url"},"counts":{"papers_tagged":46,"papers_with_code":36,"benchmarks":14,"benchmark_tables_in_archive":14,"benchmark_tables_shown":14,"benchmark_tables_withheld_as_spam":0,"benchmark_definition":"a leaderboard table with at least one row; benchmark_tables_shown also counts the zero-row tables; benchmark_tables_in_archive adds the tables withheld as spam","datasets":12,"subtasks":0,"parent_tasks":0},"benchmarks":[{"leaderboard":"/sota/conversational-response-selection-on-ubuntu-1","slug":"conversational-response-selection-on-ubuntu-1","dataset":"Ubuntu Dialogue (v1, Ranking)","dataset_url":"/dataset/ubuntu-dialogue-corpus","rows_in_archive":25,"metrics":["R10@1","R10@2","R10@5","R2@1"],"first_row_in_archive_order":{"model":"Dial-MAE","paper_title":"Dial-MAE: ConTextual Masked Auto-Encoder for Retrieval-based Dialogue Systems","paper_url":"/paper/contextual-masked-auto-encoder-for-retrieval","paper_date":"2023-06-07","arxiv_id":"2306.04357","code_links":[{"title":"suu990901/Dial-MAE","url":"https://github.com/suu990901/Dial-MAE"}],"syntology":null}},{"leaderboard":"/sota/conversational-response-selection-on-douban-1","slug":"conversational-response-selection-on-douban-1","dataset":"Douban","dataset_url":"/dataset/douban","rows_in_archive":16,"metrics":["MAP","MRR","P@1","R10@1","R10@2","R10@5"],"first_row_in_archive_order":{"model":"SEMSOL(W/o utterances)","paper_title":"Knowledge-aware response selection with semantics underlying multi-turn open-domain conversations","paper_url":"/paper/knowledge-aware-response-selection-with","paper_date":"2023-07-27","arxiv_id":null,"code_links":[{"title":"losmes/SemSol","url":"https://github.com/losmes/SemSol"}],"syntology":null}},{"leaderboard":"/sota/conversational-response-selection-on-e","slug":"conversational-response-selection-on-e","dataset":"E-commerce","dataset_url":"/dataset/e-commerce-1","rows_in_archive":15,"metrics":["R10@1","R10@2","R10@5"],"first_row_in_archive_order":{"model":"BERT-FP+EDHNS","paper_title":"Efficient Dynamic Hard Negative Sampling for Dialogue Selection","paper_url":"/paper/efficient-dynamic-hard-negative-sampling-for","paper_date":"2024-08-16","arxiv_id":null,"code_links":[{"title":"hanjanghoon/EDHNS","url":"https://github.com/hanjanghoon/EDHNS"}],"syntology":null}},{"leaderboard":"/sota/conversational-response-selection-on-rrs","slug":"conversational-response-selection-on-rrs","dataset":"RRS","dataset_url":"/dataset/rrs","rows_in_archive":7,"metrics":["MAP","MRR","P@1","R10@1","R10@2","R10@5"],"first_row_in_archive_order":{"model":"BERT-FP","paper_title":"Fine-grained Post-training for Improving Retrieval-based Dialogue Systems","paper_url":"/paper/fine-grained-post-training-for-improving","paper_date":"2021-05-24","arxiv_id":null,"code_links":[{"title":"hanjanghoon/BERT_FP","url":"https://github.com/hanjanghoon/BERT_FP"}],"syntology":null}},{"leaderboard":"/sota/conversational-response-selection-on-dstc7","slug":"conversational-response-selection-on-dstc7","dataset":"DSTC7 Ubuntu","dataset_url":"/dataset/dstc7-task-1","rows_in_archive":5,"metrics":["1-of-100 Accuracy"],"first_row_in_archive_order":{"model":"Multi-context ConveRT","paper_title":"ConveRT: Efficient and Accurate Conversational Representations from Transformers","paper_url":"/paper/convert-efficient-and-accurate-conversational","paper_date":"2019-11-09","arxiv_id":"1911.03688","code_links":[{"title":"golsun/dialogrpt","url":"https://github.com/golsun/dialogrpt"},{"title":"davidalami/convert","url":"https://github.com/davidalami/convert"},{"title":"jordiclive/Convert-PolyAI-Torch","url":"https://github.com/jordiclive/Convert-PolyAI-Torch"},{"title":"koujm/convert-tf","url":"https://github.com/koujm/convert-tf"},{"title":"phamnam-mta/ConveRT-PolyAI-Vietnamese","url":"https://github.com/phamnam-mta/ConveRT-PolyAI-Vietnamese"}],"syntology":{"n":9,"n_ran":0,"n_unverified":9,"n_pointer_only":0}}},{"leaderboard":"/sota/conversational-response-selection-on-polyai","slug":"conversational-response-selection-on-polyai","dataset":"PolyAI Reddit","dataset_url":"/dataset/reddit","rows_in_archive":5,"metrics":["1-of-100 Accuracy"],"first_row_in_archive_order":{"model":"Multi-context ConveRT","paper_title":"ConveRT: Efficient and Accurate Conversational Representations from Transformers","paper_url":"/paper/convert-efficient-and-accurate-conversational","paper_date":"2019-11-09","arxiv_id":"1911.03688","code_links":[{"title":"golsun/dialogrpt","url":"https://github.com/golsun/dialogrpt"},{"title":"davidalami/convert","url":"https://github.com/davidalami/convert"},{"title":"jordiclive/Convert-PolyAI-Torch","url":"https://github.com/jordiclive/Convert-PolyAI-Torch"},{"title":"koujm/convert-tf","url":"https://github.com/koujm/convert-tf"},{"title":"phamnam-mta/ConveRT-PolyAI-Vietnamese","url":"https://github.com/phamnam-mta/ConveRT-PolyAI-Vietnamese"}],"syntology":{"n":9,"n_ran":0,"n_unverified":9,"n_pointer_only":0}}},{"leaderboard":"/sota/conversational-response-selection-on-ubuntu-3","slug":"conversational-response-selection-on-ubuntu-3","dataset":"Ubuntu IRC","dataset_url":"/dataset/ubuntu-irc-1","rows_in_archive":5,"metrics":["Accuracy"],"first_row_in_archive_order":{"model":"MPC-BERT","paper_title":"MPC-BERT: A Pre-Trained Language Model for Multi-Party Conversation Understanding","paper_url":"/paper/mpc-bert-a-pre-trained-language-model-for","paper_date":"2021-06-03","arxiv_id":"2106.01541","code_links":[{"title":"JasonForJoy/MPC-BERT","url":"https://github.com/JasonForJoy/MPC-BERT"}],"syntology":null}},{"leaderboard":"/sota/conversational-response-selection-on-rrs-1","slug":"conversational-response-selection-on-rrs-1","dataset":"RRS Ranking Test","dataset_url":"/dataset/rrs-ranking-test","rows_in_archive":4,"metrics":["NDCG@3","NDCG@5"],"first_row_in_archive_order":{"model":"Poly-encoder","paper_title":"Poly-encoders: Transformer Architectures and Pre-training Strategies for Fast and Accurate Multi-sentence Scoring","paper_url":"/paper/190501969","paper_date":"2019-04-22","arxiv_id":"1905.01969","code_links":[{"title":"sfzhou5678/PolyEncoder","url":"https://github.com/sfzhou5678/PolyEncoder"},{"title":"chijames/Poly-Encoder","url":"https://github.com/chijames/Poly-Encoder"},{"title":"csong27/collision-bert","url":"https://github.com/csong27/collision-bert"},{"title":"llStringll/Poly-encoders","url":"https://github.com/llStringll/Poly-encoders"},{"title":"fangrouli/Document-embedding-generation-models","url":"https://github.com/fangrouli/Document-embedding-generation-models"},{"title":"i2r-simmc/i2r-simmc-2020","url":"https://github.com/i2r-simmc/i2r-simmc-2020"},{"title":"Alexey-Borisov/3_course_diary","url":"https://github.com/Alexey-Borisov/3_course_diary"}],"syntology":null}},{"leaderboard":"/sota/conversational-response-selection-on-persona","slug":"conversational-response-selection-on-persona","dataset":"Persona-Chat","dataset_url":null,"rows_in_archive":2,"metrics":["MRR","R20@1"],"first_row_in_archive_order":{"model":"Uni-Encoder","paper_title":"Uni-Encoder: A Fast and Accurate Response Selection Paradigm for Generation-Based Dialogue Systems","paper_url":"/paper/global-selector-a-new-benchmark-dataset-and","paper_date":"2021-06-02","arxiv_id":"2106.01263","code_links":[{"title":"dll-wu/uni-encoder","url":"https://github.com/dll-wu/uni-encoder"}],"syntology":null}},{"leaderboard":"/sota/conversational-response-selection-on-polyai-2","slug":"conversational-response-selection-on-polyai-2","dataset":"PolyAI AmazonQA","dataset_url":null,"rows_in_archive":2,"metrics":["1-of-100 Accuracy"],"first_row_in_archive_order":{"model":"ConveRT","paper_title":"ConveRT: Efficient and Accurate Conversational Representations from Transformers","paper_url":"/paper/convert-efficient-and-accurate-conversational","paper_date":"2019-11-09","arxiv_id":"1911.03688","code_links":[{"title":"golsun/dialogrpt","url":"https://github.com/golsun/dialogrpt"},{"title":"davidalami/convert","url":"https://github.com/davidalami/convert"},{"title":"jordiclive/Convert-PolyAI-Torch","url":"https://github.com/jordiclive/Convert-PolyAI-Torch"},{"title":"koujm/convert-tf","url":"https://github.com/koujm/convert-tf"},{"title":"phamnam-mta/ConveRT-PolyAI-Vietnamese","url":"https://github.com/phamnam-mta/ConveRT-PolyAI-Vietnamese"}],"syntology":{"n":9,"n_ran":0,"n_unverified":9,"n_pointer_only":0}}},{"leaderboard":"/sota/conversational-response-selection-on","slug":"conversational-response-selection-on","dataset":"personachat","dataset_url":null,"rows_in_archive":1,"metrics":["R20@1"],"first_row_in_archive_order":{"model":"P5","paper_title":"P5: Plug-and-Play Persona Prompting for Personalized Response Selection","paper_url":"/paper/p5-plug-and-play-persona-prompting-for","paper_date":"2023-10-10","arxiv_id":"2310.06390","code_links":[{"title":"rungjoo/plug-and-play-prompt-persona","url":"https://github.com/rungjoo/plug-and-play-prompt-persona"}],"syntology":{"n":1,"n_ran":0,"n_unverified":1,"n_pointer_only":1}}},{"leaderboard":"/sota/conversational-response-selection-on-advising","slug":"conversational-response-selection-on-advising","dataset":"Advising Corpus","dataset_url":"/dataset/advising-corpus","rows_in_archive":1,"metrics":["R@1","R@10","R@50"],"first_row_in_archive_order":{"model":"CtxDec & -Rev","paper_title":"Sequential Attention-based Network for Noetic End-to-End Response Selection","paper_url":"/paper/sequential-attention-based-network-for-noetic","paper_date":"2019-01-09","arxiv_id":"1901.02609","code_links":[{"title":"alibaba/esim-response-selection","url":"https://github.com/alibaba/esim-response-selection"},{"title":"swe-zzf/esim-response-selection","url":"https://github.com/swe-zzf/esim-response-selection"},{"title":"jadler1/retrieval-chitchat","url":"https://github.com/jadler1/retrieval-chitchat"},{"title":"ArthurRizar/dialog_state_tracking_bert","url":"https://github.com/ArthurRizar/dialog_state_tracking_bert"}],"syntology":null}},{"leaderboard":"/sota/conversational-response-selection-on-polyai-1","slug":"conversational-response-selection-on-polyai-1","dataset":"PolyAI OpenSubtitles","dataset_url":null,"rows_in_archive":1,"metrics":["1-of-100 Accuracy"],"first_row_in_archive_order":{"model":"PolyAI Encoder","paper_title":"A Repository of Conversational Datasets","paper_url":"/paper/190406472","paper_date":"2019-04-13","arxiv_id":"1904.06472","code_links":[{"title":"PolyAI-LDN/conversational-datasets","url":"https://github.com/PolyAI-LDN/conversational-datasets"},{"title":"ACE-VSIT/ACE-Ampethatic_bot","url":"https://github.com/ACE-VSIT/ACE-Ampethatic_bot"},{"title":"SarthakVaswani/ace_bot","url":"https://github.com/SarthakVaswani/ace_bot"}],"syntology":{"n":3,"n_ran":0,"n_unverified":3,"n_pointer_only":0}}},{"leaderboard":"/sota/conversational-response-selection-on-ubuntu-2","slug":"conversational-response-selection-on-ubuntu-2","dataset":"Ubuntu Dialogue (v2, Ranking)","dataset_url":"/dataset/ubuntu-dialogue-corpus","rows_in_archive":1,"metrics":["R10@1","R10@2","R10@5"],"first_row_in_archive_order":{"model":"Uni-Encoder","paper_title":"Uni-Encoder: A Fast and Accurate Response Selection Paradigm for Generation-Based Dialogue Systems","paper_url":"/paper/global-selector-a-new-benchmark-dataset-and","paper_date":"2021-06-02","arxiv_id":"2106.01263","code_links":[{"title":"dll-wu/uni-encoder","url":"https://github.com/dll-wu/uni-encoder"}],"syntology":null}}],"datasets":[{"url":"/dataset/reddit","name":"Reddit","full_name":"","num_papers_in_archive":699},{"url":"/dataset/douban","name":"Douban","full_name":"Douban Conversation Corpus","num_papers_in_archive":81},{"url":"/dataset/ubuntu-dialogue-corpus","name":"UDC","full_name":"Ubuntu Dialogue Corpus","num_papers_in_archive":46},{"url":"/dataset/e-commerce-1","name":"E-commerce","full_name":"E-commerce","num_papers_in_archive":42},{"url":"/dataset/ubuntu-irc-1","name":"Ubuntu IRC","full_name":"","num_papers_in_archive":24},{"url":"/dataset/douban-conversation-corpus","name":"Douban Conversation Corpus","full_name":"Douban Conversation Corpus","num_papers_in_archive":22},{"url":"/dataset/dstc7-task-1","name":"DSTC7 Task 1","full_name":"Dialog System Technology Challenges Task 1","num_papers_in_archive":12},{"url":"/dataset/rrs","name":"RRS","full_name":"Restoration-200k for Response Selection","num_papers_in_archive":8},{"url":"/dataset/reddit-corpus","name":"Reddit Corpus","full_name":"","num_papers_in_archive":7},{"url":"/dataset/rrs-ranking-test","name":"RRS Ranking Test","full_name":"Restoration-200k for Response Selection with Ranking Test Set","num_papers_in_archive":5},{"url":"/dataset/advising-corpus","name":"Advising Corpus","full_name":"","num_papers_in_archive":4},{"url":"/dataset/bbai-dataset","name":"BBAI Dataset","full_name":"Black-box Agent Integration","num_papers_in_archive":1}],"subtasks":[],"parent_tasks":[],"papers":{"order":"repositories listed in the archive (desc), then date (desc); the archive holds no stars","population":"papers tagged with this task that list at least one repository in the archive","shown":30,"of":36,"tagged_in_all":46,"items":[{"url":"/paper/bert-pre-training-of-deep-bidirectional","title":"BERT: Pre-training of Deep Bidirectional Transformers for Language Understanding","date":"2018-10-11","arxiv_id":"1810.04805","repositories_listed":534,"syntology":{"n":659,"n_ran":204,"n_unverified":455,"n_pointer_only":149}},{"url":"/paper/deep-contextualized-word-representations","title":"Deep contextualized word representations","date":"2018-02-15","arxiv_id":"1802.05365","repositories_listed":46,"syntology":{"n":58,"n_ran":23,"n_unverified":35,"n_pointer_only":25}},{"url":"/paper/universal-sentence-encoder","title":"Universal Sentence Encoder","date":"2018-03-29","arxiv_id":"1803.11175","repositories_listed":24,"syntology":{"n":22,"n_ran":1,"n_unverified":21,"n_pointer_only":1}},{"url":"/paper/the-ubuntu-dialogue-corpus-a-large-dataset-1","title":"The Ubuntu Dialogue Corpus: A Large Dataset for Research in Unstructured Multi-Turn Dialogue Systems","date":"2015-06-30","arxiv_id":"1506.08909","repositories_listed":21,"syntology":{"n":26,"n_ran":1,"n_unverified":25,"n_pointer_only":1}},{"url":"/paper/personalizing-dialogue-agents-i-have-a-dog-do","title":"Personalizing Dialogue Agents: I have a dog, do you have pets too?","date":"2018-01-22","arxiv_id":"1801.07243","repositories_listed":15,"syntology":{"n":1,"n_ran":0,"n_unverified":1,"n_pointer_only":1}},{"url":"/paper/190501969","title":"Poly-encoders: Transformer Architectures and Pre-training Strategies for Fast and Accurate Multi-sentence Scoring","date":"2019-04-22","arxiv_id":"1905.01969","repositories_listed":7,"syntology":null},{"url":"/paper/convert-efficient-and-accurate-conversational","title":"ConveRT: Efficient and Accurate Conversational Representations from Transformers","date":"2019-11-09","arxiv_id":"1911.03688","repositories_listed":5,"syntology":{"n":9,"n_ran":0,"n_unverified":9,"n_pointer_only":0}},{"url":"/paper/sequential-attention-based-network-for-noetic","title":"Sequential Attention-based Network for Noetic End-to-End Response Selection","date":"2019-01-09","arxiv_id":"1901.02609","repositories_listed":4,"syntology":null},{"url":"/paper/190406472","title":"A Repository of Conversational Datasets","date":"2019-04-13","arxiv_id":"1904.06472","repositories_listed":3,"syntology":{"n":3,"n_ran":0,"n_unverified":3,"n_pointer_only":0}},{"url":"/paper/sequential-matching-network-a-new","title":"Sequential Matching Network: A New Architecture for Multi-turn Response Selection in Retrieval-based Chatbots","date":"2016-12-06","arxiv_id":"1612.01627","repositories_listed":3,"syntology":{"n":3,"n_ran":0,"n_unverified":3,"n_pointer_only":0}},{"url":"/paper/dialogue-response-ranking-training-with-large","title":"Dialogue Response Ranking Training with Large-Scale Human Feedback Data","date":"2020-09-15","arxiv_id":"2009.06978","repositories_listed":2,"syntology":{"n":6,"n_ran":1,"n_unverified":5,"n_pointer_only":0}},{"url":"/paper/speaker-aware-bert-for-multi-turn-response","title":"Speaker-Aware BERT for Multi-Turn Response Selection in Retrieval-Based Chatbots","date":"2020-04-07","arxiv_id":"2004.03588","repositories_listed":2,"syntology":null},{"url":"/paper/efficient-dynamic-hard-negative-sampling-for","title":"Efficient Dynamic Hard Negative Sampling for Dialogue Selection","date":"2024-08-16","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/p5-plug-and-play-persona-prompting-for","title":"P5: Plug-and-Play Persona Prompting for Personalized Response Selection","date":"2023-10-10","arxiv_id":"2310.06390","repositories_listed":1,"syntology":{"n":1,"n_ran":0,"n_unverified":1,"n_pointer_only":1}},{"url":"/paper/knowledge-aware-response-selection-with","title":"Knowledge-aware response selection with semantics underlying multi-turn open-domain conversations","date":"2023-07-27","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/contextual-masked-auto-encoder-for-retrieval","title":"Dial-MAE: ConTextual Masked Auto-Encoder for Retrieval-based Dialogue Systems","date":"2023-06-07","arxiv_id":"2306.04357","repositories_listed":1,"syntology":null},{"url":"/paper/learning-dialogue-representations-from","title":"Learning Dialogue Representations from Consecutive Utterances","date":"2022-05-26","arxiv_id":"2205.13568","repositories_listed":1,"syntology":null},{"url":"/paper/one-agent-to-rule-them-all-towards-multi-1","title":"One Agent To Rule Them All: Towards Multi-agent Conversational AI","date":"2022-03-15","arxiv_id":"2203.07665","repositories_listed":1,"syntology":null},{"url":"/paper/exploring-dense-retrieval-for-dialogue","title":"Exploring Dense Retrieval for Dialogue Response Selection","date":"2021-10-13","arxiv_id":"2110.06612","repositories_listed":1,"syntology":{"n":3,"n_ran":0,"n_unverified":3,"n_pointer_only":0}},{"url":"/paper/response-ranking-with-multi-types-of-deep","title":"Response Ranking with Multi-types of Deep Interactive Representations in Retrieval-based Dialogues","date":"2021-08-17","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/mpc-bert-a-pre-trained-language-model-for","title":"MPC-BERT: A Pre-Trained Language Model for Multi-Party Conversation Understanding","date":"2021-06-03","arxiv_id":"2106.01541","repositories_listed":1,"syntology":null},{"url":"/paper/global-selector-a-new-benchmark-dataset-and","title":"Uni-Encoder: A Fast and Accurate Response Selection Paradigm for Generation-Based Dialogue Systems","date":"2021-06-02","arxiv_id":"2106.01263","repositories_listed":1,"syntology":null},{"url":"/paper/fine-grained-post-training-for-improving","title":"Fine-grained Post-training for Improving Retrieval-based Dialogue Systems","date":"2021-05-24","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/open-domain-question-classification-and","title":"Open-domain question classification and completion in conversational information search","date":"2021-02-26","arxiv_id":"2102.13495","repositories_listed":1,"syntology":null},{"url":"/paper/dialogue-response-selection-with-hierarchical","title":"Dialogue Response Selection with Hierarchical Curriculum Learning","date":"2020-12-29","arxiv_id":"2012.14756","repositories_listed":1,"syntology":null},{"url":"/paper/do-response-selection-models-really-know-what","title":"Do Response Selection Models Really Know What's Next? Utterance Manipulation Strategies for Multi-turn Response Selection","date":"2020-09-10","arxiv_id":"2009.04703","repositories_listed":1,"syntology":null},{"url":"/paper/utterance-to-utterance-interactive-matching","title":"Utterance-to-Utterance Interactive Matching Network for Multi-Turn Response Selection in Retrieval-Based Chatbots","date":"2019-11-16","arxiv_id":"1911.06940","repositories_listed":1,"syntology":null},{"url":"/paper/multi-hop-selector-network-for-multi-turn","title":"Multi-hop Selector Network for Multi-turn Response Selection in Retrieval-based Chatbots","date":"2019-11-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/domain-adaptive-training-bert-for-response","title":"An Effective Domain Adaptive Post-Training Method for BERT in Response Selection","date":"2019-08-13","arxiv_id":"1908.04812","repositories_listed":1,"syntology":null},{"url":"/paper/one-time-of-interaction-may-not-be-enough-go","title":"One Time of Interaction May Not Be Enough: Go Deep with an Interaction-over-Interaction Network for Response Selection in Dialogues","date":"2019-07-01","arxiv_id":null,"repositories_listed":1,"syntology":null}],"syntology_records":11,"syntology_note":"a paper without a record is not a recorded non-run: it may lack an arXiv id or simply be absent from the graph layer"},"description_links":{"kept":0,"unwrapped_to_text":0,"bare_urls_linked":0,"relative_images_dropped":0,"rule":"internal links are kept only when the target slug exists in the catalog"},"syntology":{"read_at":"2026-09-24T18:15:14+00:00","claim":"Per-sample execution status on synthesized fixtures ('ran N of M samples'); not a correctness claim and not a ranking signal.","status_vocabulary":{"ran_honours":"ran, honoured the contract we drafted","ran_violates":"ran, violated the contract we drafted","ran_draft_wrong":"ran; our contract draft was wrong, not the code","ran_fixture":"ran; our fixture could not drive it","ran":"ran on a synthesized input","unverified":"unverified (harvested, no recorded run)"}},"not_shown":{"libraries":"the archive has no per-task library table","trend_sparklines":"the Trend column of the benchmarks table was a rendered image; it is not in the archive","social_and_latest_sorts":"stars and social signals are not in the archive"}}