{"url":"/method/fastformer","slug":"fastformer","name":"Fastformer","full_name":"Fastformer","full_name_withheld":false,"description_markdown":"**Fastformer** is an type of [Transformer](https://paperswithcode.com/method/transformer) which uses [additive attention](https://www.paperswithcode.com/method/additive-attention) as a building block. Instead of modeling the pair-wise interactions between tokens, [additive attention](https://paperswithcode.com/method/additive-attention) is used to model global contexts, and then each token representation is further transformed based on its interaction with global context representations.","description_state":"present","introduced_year":null,"introduced_by":{"title":"Fastformer: Additive Attention Can Be All You Need","paper":"/paper/fastformer-additive-attention-is-all-you-need","first_author":"Chuhan Wu","n_authors":5,"url_abs":null,"archive_paper_url":"https://paperswithcode.com/paper/fastformer-additive-attention-is-all-you-need"},"source":{"url":"https://arxiv.org/abs/2108.09084v6","title":"Fastformer: Additive Attention Can Be All You Need","url_on_a_paper_host":true},"code_snippet_url":null,"code_snippet_url_on_a_code_host":false,"categories":[{"area":"Natural Language Processing","area_id":"natural-language-processing","collection":"Transformers","url":"/methods/category/transformers","pwc_aliases":[]}],"n_papers_tagged":4,"archive_num_papers":4,"papers_newest_first":[{"paper":null,"title":"Multi-Granularity Vision Fastformer with Fusion Mechanism for Skin Lesion Segmentation","date":"2025-04-04","arxiv_id":"2504.03108","n_code_links":0,"syntology":null},{"paper":null,"title":"An Analysis of Linear Complexity Attention Substitutes with BEST-RQ","date":"2024-09-04","arxiv_id":"2409.02596","n_code_links":0,"syntology":null},{"paper":null,"title":"FUM: Fine-grained and Fast User Modeling for News Recommendation","date":"2022-04-10","arxiv_id":"2204.04727","n_code_links":0,"syntology":null},{"paper":"/paper/fastformer-additive-attention-is-all-you-need","title":"Fastformer: Additive Attention Can Be All You Need","date":"2021-08-20","arxiv_id":"2108.09084","n_code_links":13,"syntology":{"ran":4,"of":4,"unverified":0,"pointer_only":3}}],"papers_shown":4,"tasks":[{"task":"/task/news-recommendation","name":"News Recommendation","papers":2},{"task":"/task/all","name":"All","papers":1},{"task":"/task/image-segmentation","name":"Image Segmentation","papers":1},{"task":"/task/lesion-segmentation","name":"Lesion Segmentation","papers":1},{"task":"/task/mamba","name":"Mamba","papers":1},{"task":"/task/medical-image-segmentation","name":"Medical Image Segmentation","papers":1},{"task":"/task/segmentation","name":"Segmentation","papers":1},{"task":"/task/self-supervised-learning","name":"Self-Supervised Learning","papers":1},{"task":"/task/semantic-segmentation","name":"Semantic Segmentation","papers":1},{"task":"/task/skin-lesion-segmentation","name":"Skin Lesion Segmentation","papers":1},{"task":"/task/text-classification","name":"Text Classification","papers":1},{"task":"/task/text-summarization","name":"Text Summarization","papers":1}],"tasks_shown":12,"n_tasks":12,"usage_by_year":[{"year":"2021","papers":1},{"year":"2022","papers":1},{"year":"2024","papers":1},{"year":"2025","papers":1}],"row_source":"methods_table","archive":{"source":"pwc-archive (Hugging Face), CC BY-SA 4.0","snapshot":"2025-07-28","archive_url":"https://paperswithcode.com/method/fastformer"},"syntology_read_at":"2026-09-24T18:15:14+00:00"}