{"url":"/method/mixtext","slug":"mixtext","name":"MixText","full_name":"MixText","full_name_withheld":false,"description_markdown":"**MixText** is a semi-supervised learning method for text classification, which uses a new data augmentation method called TMix. TMix creates a large amount of augmented training samples by interpolating text in hidden space. The technique leverages advances in data augmentation to guess low-entropy labels for unlabeled data, making them as easy to use as labeled data.","description_state":"present","introduced_year":null,"introduced_by":{"title":"MixText: Linguistically-Informed Interpolation of Hidden Space for Semi-Supervised Text Classification","paper":"/paper/mixtext-linguistically-informed-interpolation","first_author":"Jiaao Chen","n_authors":3,"url_abs":null,"archive_paper_url":"https://paperswithcode.com/paper/mixtext-linguistically-informed-interpolation"},"source":{"url":"https://arxiv.org/abs/2004.12239v1","title":"MixText: Linguistically-Informed Interpolation of Hidden Space for Semi-Supervised Text Classification","url_on_a_paper_host":true},"code_snippet_url":null,"code_snippet_url_on_a_code_host":false,"categories":[{"area":"Natural Language Processing","area_id":"natural-language-processing","collection":"Text Augmentation","url":"/methods/category/text-augmentation","pwc_aliases":[]},{"area":"Natural Language Processing","area_id":"natural-language-processing","collection":"Text Classification Models","url":"/methods/category/text-classification-models","pwc_aliases":[]},{"area":"General","area_id":"general","collection":"Semi-Supervised Learning Methods","url":"/methods/category/semi-supervised-learning-methods","pwc_aliases":[]}],"n_papers_tagged":3,"archive_num_papers":3,"papers_newest_first":[{"paper":null,"title":"FPMT: Enhanced Semi-Supervised Model for Traffic Incident Detection","date":"2024-09-12","arxiv_id":"2409.07839","n_code_links":0,"syntology":null},{"paper":"/paper/llm-as-a-coauthor-the-challenges-of-detecting","title":"LLM-as-a-Coauthor: Can Mixed Human-Written and Machine-Generated Text Be Detected?","date":"2024-01-11","arxiv_id":"2401.05952","n_code_links":2,"syntology":{"ran":6,"of":16,"unverified":10,"pointer_only":16}},{"paper":"/paper/mixtext-linguistically-informed-interpolation","title":"MixText: Linguistically-Informed Interpolation of Hidden Space for Semi-Supervised Text Classification","date":"2020-04-25","arxiv_id":"2004.12239","n_code_links":2,"syntology":null}],"papers_shown":3,"tasks":[{"task":"/task/data-augmentation","name":"Data Augmentation","papers":2},{"task":"/task/binary-text-classification","name":"Binary text classification","papers":1},{"task":"/task/classification-1","name":"Classification","papers":1},{"task":"/task/classification","name":"General Classification","papers":1},{"task":"/task/semi-supervised-text-classification-1","name":"Semi-Supervised Text Classification","papers":1},{"task":"/task/text-classification","name":"Text Classification","papers":1}],"tasks_shown":6,"n_tasks":6,"usage_by_year":[{"year":"2020","papers":1},{"year":"2024","papers":2}],"row_source":"methods_table","archive":{"source":"pwc-archive (Hugging Face), CC BY-SA 4.0","snapshot":"2025-07-28","archive_url":"https://paperswithcode.com/method/mixtext"},"syntology_read_at":"2026-09-24T18:15:14+00:00"}