{"url":"/method/levit","slug":"levit","name":"LeVIT","full_name":"LeVIT","full_name_withheld":false,"description_markdown":"**LeVIT** is a hybrid neural network for fast inference image classification. LeViT is a stack of [transformer blocks](https://paperswithcode.com/method/transformer), with [pooling steps](https://paperswithcode.com/methods/category/pooling-operation) to reduce the resolution of the activation maps as in classical [convolutional architectures](https://paperswithcode.com/methods/category/convolutional-neural-networks). This replaces the uniform structure of a Transformer by a pyramid with pooling, similar to the [LeNet](https://paperswithcode.com/method/lenet) architecture","description_state":"present","introduced_year":null,"introduced_by":{"title":"LeViT: a Vision Transformer in ConvNet's Clothing for Faster Inference","paper":"/paper/levit-a-vision-transformer-in-convnet-s","first_author":"Ben Graham","n_authors":7,"url_abs":null,"archive_paper_url":"https://paperswithcode.com/paper/levit-a-vision-transformer-in-convnet-s"},"source":{"url":"https://arxiv.org/abs/2104.01136v2","title":"LeViT: a Vision Transformer in ConvNet's Clothing for Faster Inference","url_on_a_paper_host":true},"code_snippet_url":null,"code_snippet_url_on_a_code_host":false,"categories":[{"area":"Computer Vision","area_id":"computer-vision","collection":"Vision Transformers","url":"/methods/category/vision-transformers","pwc_aliases":["vision-transformer"]}],"n_papers_tagged":2,"archive_num_papers":2,"papers_newest_first":[{"paper":"/paper/learning-cortical-anomaly-through-masked","title":"Learning Cortical Anomaly through Masked Encoding for Unsupervised Heterogeneity Mapping","date":"2023-12-05","arxiv_id":"2312.02762","n_code_links":1,"syntology":null},{"paper":"/paper/levit-a-vision-transformer-in-convnet-s","title":"LeViT: a Vision Transformer in ConvNet's Clothing for Faster Inference","date":"2021-04-02","arxiv_id":"2104.01136","n_code_links":12,"syntology":{"ran":23,"of":30,"unverified":7,"pointer_only":0}}],"papers_shown":2,"tasks":[{"task":"/task/anomaly-detection","name":"Anomaly Detection","papers":1},{"task":null,"name":"CPU","papers":1},{"task":"/task/classification","name":"General Classification","papers":1},{"task":"/task/image-classification","name":"Image Classification","papers":1},{"task":"/task/image-classification","name":"image-classification","papers":1}],"tasks_shown":5,"n_tasks":5,"usage_by_year":[{"year":"2021","papers":1},{"year":"2023","papers":1}],"row_source":"methods_table","archive":{"source":"pwc-archive (Hugging Face), CC BY-SA 4.0","snapshot":"2025-07-28","archive_url":"https://paperswithcode.com/method/levit"},"syntology_read_at":"2026-09-24T18:15:14+00:00"}