{"url":"/method/cutmix","slug":"cutmix","name":"CutMix","full_name":"CutMix","full_name_withheld":false,"description_markdown":"**CutMix** is an image data augmentation strategy. Instead of simply removing pixels as in [Cutout](https://paperswithcode.com/method/cutout), we replace the removed regions with a patch from another image. The ground truth labels are also mixed proportionally to the number of pixels of combined images. The added patches further enhance localization ability by requiring the model to identify the object from a partial view.","description_state":"present","introduced_year":null,"introduced_by":{"title":"CutMix: Regularization Strategy to Train Strong Classifiers with Localizable Features","paper":"/paper/cutmix-regularization-strategy-to-train","first_author":"Sangdoo Yun","n_authors":6,"url_abs":null,"archive_paper_url":"https://paperswithcode.com/paper/cutmix-regularization-strategy-to-train"},"source":{"url":"https://arxiv.org/abs/1905.04899v2","title":"CutMix: Regularization Strategy to Train Strong Classifiers with Localizable Features","url_on_a_paper_host":true},"code_snippet_url":null,"code_snippet_url_on_a_code_host":false,"categories":[{"area":"Computer Vision","area_id":"computer-vision","collection":"Image Data Augmentation","url":"/methods/category/image-data-augmentation","pwc_aliases":[]}],"n_papers_tagged":208,"archive_num_papers":208,"papers_newest_first":[{"paper":null,"title":"Pattern-Based Phase-Separation of Tracer and Dispersed Phase Particles in Two-Phase Defocusing Particle Tracking Velocimetry","date":"2025-06-22","arxiv_id":"2506.18157","n_code_links":0,"syntology":null},{"paper":null,"title":"Compositional Attribute Imbalance in Vision Datasets","date":"2025-06-17","arxiv_id":"2506.14418","n_code_links":0,"syntology":null},{"paper":"/paper/saint-attention-based-modeling-of-sub-action","title":"SAINT: Attention-Based Modeling of Sub-Action Dependencies in Multi-Action Policies","date":"2025-05-17","arxiv_id":"2505.12109","n_code_links":1,"syntology":null},{"paper":"/paper/mediaug-exploring-visual-augmentation-in","title":"MediAug: Exploring Visual Augmentation in Medical Imaging","date":"2025-04-26","arxiv_id":"2504.18983","n_code_links":1,"syntology":null},{"paper":"/paper/event-based-crossing-dataset-ebcd","title":"Event-Based Crossing Dataset (EBCD)","date":"2025-03-21","arxiv_id":"2503.17499","n_code_links":1,"syntology":null},{"paper":"/paper/similarity-aware-token-pruning-your-vlm-but","title":"Similarity-Aware Token Pruning: Your VLM but Faster","date":"2025-03-14","arxiv_id":"2503.11549","n_code_links":1,"syntology":null},{"paper":"/paper/walnutdata-a-uav-remote-sensing-dataset-of","title":"WalnutData: A UAV Remote Sensing Dataset of Green Walnuts and Model Evaluation","date":"2025-02-27","arxiv_id":"2502.20092","n_code_links":1,"syntology":null},{"paper":null,"title":"Towards Understanding Why Data Augmentation Improves Generalization","date":"2025-02-13","arxiv_id":"2502.08940","n_code_links":0,"syntology":null},{"paper":null,"title":"Vision-Integrated LLMs for Autonomous Driving Assistance : Human Performance Comparison and Trust Evaluation","date":"2025-02-06","arxiv_id":"2502.06843","n_code_links":0,"syntology":null},{"paper":null,"title":"YOLOv4: A Breakthrough in Real-Time Object Detection","date":"2025-02-06","arxiv_id":"2502.04161","n_code_links":0,"syntology":null},{"paper":null,"title":"ASCenD-BDS: Adaptable, Stochastic and Context-aware framework for Detection of Bias, Discrimination and Stereotyping","date":"2025-02-04","arxiv_id":"2502.02072","n_code_links":0,"syntology":null},{"paper":null,"title":"SPFFNet: Strip Perception and Feature Fusion Spatial Pyramid Pooling for Fabric Defect Detection","date":"2025-02-03","arxiv_id":"2502.01445","n_code_links":0,"syntology":null},{"paper":null,"title":"Efficient Object Detection of Marine Debris using Pruned YOLO Model","date":"2025-01-27","arxiv_id":"2501.16571","n_code_links":0,"syntology":null},{"paper":"/paper/tdattenmix-top-down-attention-guided-mixup","title":"TdAttenMix: Top-Down Attention Guided Mixup","date":"2025-01-26","arxiv_id":"2501.15409","n_code_links":1,"syntology":null},{"paper":"/paper/deepifsa-deep-imputation-of-missing-values","title":"DeepIFSAC: Deep Imputation of Missing Values Using Feature and Sample Attention within Contrastive Framework","date":"2025-01-19","arxiv_id":"2501.10910","n_code_links":1,"syntology":null},{"paper":"/paper/computational-analysis-of-yaredawi-yezema","title":"Computational Analysis of Yaredawi YeZema Silt in Ethiopian Orthodox Tewahedo Church Chants","date":"2024-12-25","arxiv_id":"2412.18788","n_code_links":1,"syntology":null},{"paper":null,"title":"Object Detection Approaches to Identifying Hand Images with High Forensic Values","date":"2024-12-21","arxiv_id":"2412.16431","n_code_links":0,"syntology":null},{"paper":null,"title":"Exploring Machine Learning Engineering for Object Detection and Tracking by Unmanned Aerial Vehicle (UAV)","date":"2024-12-19","arxiv_id":"2412.15347","n_code_links":0,"syntology":null},{"paper":null,"title":"J-CaPA : Joint Channel and Pyramid Attention Improves Medical Image Segmentation","date":"2024-11-25","arxiv_id":"2411.16568","n_code_links":0,"syntology":null},{"paper":null,"title":"Provable Benefit of Cutout and CutMix for Feature Learning","date":"2024-10-31","arxiv_id":"2410.23672","n_code_links":0,"syntology":null},{"paper":null,"title":"A Survey on Deep Tabular Learning","date":"2024-10-15","arxiv_id":"2410.12034","n_code_links":0,"syntology":null},{"paper":null,"title":"Impact of Regularization on Calibration and Robustness: from the Representation Space Perspective","date":"2024-10-05","arxiv_id":"2410.03999","n_code_links":0,"syntology":null},{"paper":null,"title":"SAFLEX: Self-Adaptive Augmentation via Feature Label Extrapolation","date":"2024-10-03","arxiv_id":"2410.02512","n_code_links":0,"syntology":null},{"paper":null,"title":"UICE-MIRNet guided image enhancement for underwater object detection","date":"2024-09-24","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"title":"Privacy-Preserving Split Learning with Vision Transformers using Patch-Wise Random and Noisy CutMix","date":"2024-08-02","arxiv_id":"2408.01040","n_code_links":0,"syntology":null},{"paper":"/paper/sumix-mixup-with-semantic-and-uncertain","title":"SUMix: Mixup with Semantic and Uncertain Information","date":"2024-07-10","arxiv_id":"2407.07805","n_code_links":2,"syntology":{"ran":3,"of":6,"unverified":3,"pointer_only":0}},{"paper":"/paper/enhanced-long-tailed-recognition-with","title":"Enhanced Long-Tailed Recognition with Contrastive CutMix Augmentation","date":"2024-07-06","arxiv_id":"2407.04911","n_code_links":2,"syntology":null},{"paper":null,"title":"Semantic Compositions Enhance Vision-Language Contrastive Learning","date":"2024-07-01","arxiv_id":"2407.01408","n_code_links":0,"syntology":null},{"paper":"/paper/back-to-the-color-learning-depth-to-specific","title":"Back to the Color: Learning Depth to Specific Color Transformation for Unsupervised Depth Estimation","date":"2024-06-11","arxiv_id":"2406.07741","n_code_links":1,"syntology":null},{"paper":null,"title":"A Label Propagation Strategy for CutMix in Multi-Label Remote Sensing Image Classification","date":"2024-05-22","arxiv_id":"2405.13451","n_code_links":0,"syntology":null}],"papers_shown":30,"tasks":[{"task":"/task/object-detection","name":"Object Detection","papers":69},{"task":"/task/object-detection-1","name":"object-detection","papers":66},{"task":"/task/data-augmentation","name":"Data Augmentation","papers":63},{"task":"/task/object","name":"Object","papers":30},{"task":"/task/image-classification","name":"Image Classification","papers":28},{"task":"/task/semantic-segmentation","name":"Semantic Segmentation","papers":25},{"task":"/task/image-classification","name":"image-classification","papers":20},{"task":"/task/segmentation","name":"Segmentation","papers":14},{"task":"/task/classification-1","name":"Classification","papers":9},{"task":"/task/real-time-object-detection","name":"Real-Time Object Detection","papers":9},{"task":"/task/deep-learning","name":"Deep Learning","papers":8},{"task":null,"name":"GPU","papers":8},{"task":"/task/instance-segmentation","name":"Instance Segmentation","papers":8},{"task":"/task/knowledge-distillation","name":"Knowledge Distillation","papers":8},{"task":"/task/semi-supervised-semantic-segmentation","name":"Semi-Supervised Semantic Segmentation","papers":8},{"task":"/task/transfer-learning","name":"Transfer Learning","papers":8},{"task":"/task/contrastive-learning","name":"Contrastive Learning","papers":7},{"task":"/task/classification","name":"General Classification","papers":7},{"task":"/task/autonomous-driving","name":"Autonomous Driving","papers":6},{"task":"/task/image-segmentation","name":"Image Segmentation","papers":6}],"tasks_shown":20,"n_tasks":205,"usage_by_year":[{"year":"2016","papers":1},{"year":"2019","papers":2},{"year":"2020","papers":35},{"year":"2021","papers":43},{"year":"2022","papers":56},{"year":"2023","papers":34},{"year":"2024","papers":22},{"year":"2025","papers":15}],"row_source":"methods_table","archive":{"source":"pwc-archive (Hugging Face), CC BY-SA 4.0","snapshot":"2025-07-28","archive_url":"https://paperswithcode.com/method/cutmix"},"syntology_read_at":"2026-09-24T18:15:14+00:00"}