{"url":"/method/actkr","slug":"actkr","name":"ACTKR","full_name":"ACTKR","full_name_withheld":false,"description_markdown":"**ACKTR**, or **Actor Critic with Kronecker-factored Trust Region**, is an actor-critic method for reinforcement learning that applies [trust region optimization](https://paperswithcode.com/method/trpo) using a recently proposed Kronecker-factored approximation to the curvature. The method extends the framework of natural policy gradient and optimizes both the actor and the critic using Kronecker-factored approximate\r\ncurvature (K-FAC) with trust region.","description_state":"present","introduced_year":null,"introduced_by":{"title":"Scalable trust-region method for deep reinforcement learning using Kronecker-factored approximation","paper":"/paper/scalable-trust-region-method-for-deep","first_author":"Yuhuai Wu","n_authors":5,"url_abs":null,"archive_paper_url":"https://paperswithcode.com/paper/scalable-trust-region-method-for-deep"},"source":{"url":"http://arxiv.org/abs/1708.05144v2","title":"Scalable trust-region method for deep reinforcement learning using Kronecker-factored approximation","url_on_a_paper_host":true},"code_snippet_url":null,"code_snippet_url_on_a_code_host":false,"categories":[{"area":"Reinforcement Learning","area_id":"reinforcement-learning","collection":"Policy Gradient Methods","url":"/methods/category/policy-gradient-methods","pwc_aliases":[]}],"n_papers_tagged":2,"archive_num_papers":2,"papers_newest_first":[{"paper":null,"title":"myGym: Modular Toolkit for Visuomotor Robotic Tasks","date":"2020-12-21","arxiv_id":"2012.11643","n_code_links":0,"syntology":null},{"paper":"/paper/scalable-trust-region-method-for-deep","title":"Scalable trust-region method for deep reinforcement learning using Kronecker-factored approximation","date":"2017-08-17","arxiv_id":"1708.05144","n_code_links":8,"syntology":{"ran":1,"of":1,"unverified":0,"pointer_only":0}}],"papers_shown":2,"tasks":[{"task":"/task/reinforcement-learning-1","name":"Reinforcement Learning (RL)","papers":2},{"task":"/task/atari-games","name":"Atari Games","papers":1},{"task":"/task/continuous-control","name":"Continuous Control","papers":1},{"task":"/task/deep-reinforcement-learning","name":"Deep Reinforcement Learning","papers":1},{"task":"/task/imitation-learning","name":"Imitation Learning","papers":1},{"task":"/task/mujoco","name":"MuJoCo","papers":1},{"task":"/task/openai-gym","name":"OpenAI Gym","papers":1},{"task":"/task/reinforcement-learning","name":"Reinforcement Learning","papers":1},{"task":"/task/continuous-control","name":"continuous-control","papers":1},{"task":"/task/reinforcement-learning-2","name":"reinforcement-learning","papers":1}],"tasks_shown":10,"n_tasks":10,"usage_by_year":[{"year":"2017","papers":1},{"year":"2020","papers":1}],"row_source":"methods_table","archive":{"source":"pwc-archive (Hugging Face), CC BY-SA 4.0","snapshot":"2025-07-28","archive_url":"https://paperswithcode.com/method/actkr"},"syntology_read_at":"2026-09-24T18:15:14+00:00"}