{"url":"/method/torchbeast","slug":"torchbeast","name":"TorchBeast","full_name":"TorchBeast","full_name_withheld":false,"description_markdown":"**TorchBeast** is a platform for reinforcement learning (RL) research in PyTorch. It implements a version of the popular [IMPALA](https://paperswithcode.com/method/impala) algorithm for fast, asynchronous, parallel training of RL agents.","description_state":"present","introduced_year":null,"introduced_by":{"title":"TorchBeast: A PyTorch Platform for Distributed RL","paper":"/paper/torchbeast-a-pytorch-platform-for-distributed","first_author":"Heinrich Küttler","n_authors":7,"url_abs":null,"archive_paper_url":"https://paperswithcode.com/paper/torchbeast-a-pytorch-platform-for-distributed"},"source":{"url":"https://arxiv.org/abs/1910.03552v1","title":"TorchBeast: A PyTorch Platform for Distributed RL","url_on_a_paper_host":true},"code_snippet_url":null,"code_snippet_url_on_a_code_host":false,"categories":[{"area":"Reinforcement Learning","area_id":"reinforcement-learning","collection":"Distributed Reinforcement Learning","url":"/methods/category/distributed-reinforcement-learning","pwc_aliases":[]},{"area":"General","area_id":"general","collection":"Distributed Methods","url":"/methods/category/distributed-methods","pwc_aliases":[]}],"n_papers_tagged":2,"archive_num_papers":2,"papers_newest_first":[{"paper":"/paper/cleanba-a-reproducible-and-efficient","title":"Cleanba: A Reproducible and Efficient Distributed Reinforcement Learning Platform","date":"2023-09-29","arxiv_id":"2310.00036","n_code_links":1,"syntology":{"ran":4,"of":8,"unverified":4,"pointer_only":8}},{"paper":"/paper/torchbeast-a-pytorch-platform-for-distributed","title":"TorchBeast: A PyTorch Platform for Distributed RL","date":"2019-10-08","arxiv_id":"1910.03552","n_code_links":3,"syntology":{"ran":4,"of":9,"unverified":5,"pointer_only":3}}],"papers_shown":2,"tasks":[{"task":"/task/reinforcement-learning","name":"Reinforcement Learning","papers":2},{"task":"/task/deep-reinforcement-learning","name":"Deep Reinforcement Learning","papers":1},{"task":"/task/openai-gym","name":"OpenAI Gym","papers":1},{"task":"/task/reinforcement-learning-1","name":"Reinforcement Learning (RL)","papers":1},{"task":"/task/reinforcement-learning-2","name":"reinforcement-learning","papers":1}],"tasks_shown":5,"n_tasks":5,"usage_by_year":[{"year":"2019","papers":1},{"year":"2023","papers":1}],"row_source":"methods_table","archive":{"source":"pwc-archive (Hugging Face), CC BY-SA 4.0","snapshot":"2025-07-28","archive_url":"https://paperswithcode.com/method/torchbeast"},"syntology_read_at":"2026-09-24T18:15:14+00:00"}