{"url":"/method/pipedream","slug":"pipedream","name":"PipeDream","full_name":"PipeDream","full_name_withheld":false,"description_markdown":"PipeDream is an asynchronous pipeline parallel strategy for training large neural networks. It adds inter-batch pipelining to intra-batch parallelism to further improve parallel training throughput, helping to better overlap computation with communication and reduce the amount of communication when possible.","description_state":"present","introduced_year":2019,"introduced_by":{"title":null,"paper":null,"first_author":null,"n_authors":0,"url_abs":null,"archive_paper_url":null},"source":{"url":null,"title":null,"url_on_a_paper_host":false},"code_snippet_url":null,"code_snippet_url_on_a_code_host":false,"categories":[{"area":"General","area_id":"general","collection":"Distributed Methods","url":"/methods/category/distributed-methods","pwc_aliases":[]}],"n_papers_tagged":5,"archive_num_papers":null,"papers_newest_first":[{"paper":null,"title":"GraphPipe: Improving Performance and Scalability of DNN Training with Graph Pipeline Parallelism","date":"2024-06-24","arxiv_id":"2406.17145","n_code_links":0,"syntology":null},{"paper":"/paper/pipeoptim-ensuring-effective-1f1b-schedule","title":"PipeOptim: Ensuring Effective 1F1B Schedule with Optimizer-Dependent Weight Prediction","date":"2023-12-01","arxiv_id":"2312.00839","n_code_links":1,"syntology":null},{"paper":null,"title":"Pipeline Parallelism for Inference on Heterogeneous Edge Computing","date":"2021-10-28","arxiv_id":"2110.14895","n_code_links":0,"syntology":null},{"paper":"/paper/group-based-interleaved-pipeline-parallelism","title":"Group-based Interleaved Pipeline Parallelism for Large-scale DNN Training","date":"2021-09-29","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"title":"LayerPipe: Accelerating Deep Neural Network Training by Intra-Layer and Inter-Layer Gradient Pipelining and Multiprocessor Scheduling","date":"2021-08-14","arxiv_id":"2108.06629","n_code_links":0,"syntology":null}],"papers_shown":5,"tasks":[{"task":"/task/edge-computing","name":"Edge-computing","papers":1},{"task":null,"name":"GPU","papers":1},{"task":"/task/image-classification","name":"Image Classification","papers":1},{"task":"/task/machine-translation","name":"Machine Translation","papers":1},{"task":"/task/prediction","name":"Prediction","papers":1},{"task":"/task/scheduling","name":"Scheduling","papers":1},{"task":"/task/sentiment-analysis","name":"Sentiment Analysis","papers":1},{"task":"/task/image-classification","name":"image-classification","papers":1}],"tasks_shown":8,"n_tasks":8,"usage_by_year":[{"year":"2021","papers":3},{"year":"2023","papers":1},{"year":"2024","papers":1}],"row_source":"embedded","archive":{"source":"pwc-archive (Hugging Face), CC BY-SA 4.0","snapshot":"2025-07-28","archive_url":"https://paperswithcode.com/method/pipedream"},"syntology_read_at":"2026-09-24T18:15:14+00:00"}