@misc{indiciaee96949d3fe2a, title = {Fine-tuning Reinforcement Learning Models is Secretly a Forgetting Mitigation Problem}, author = {Maciej Wołczyk and Bartłomiej Cupiał and Mateusz Ostaszewski and Michał Bortkiewicz and Michał Zając and Razvan Pascanu and Łukasz Kuciński and Piotr Miłoś}, year = {2024}, url = {https://arxiv.org/abs/2402.02868}, note = {Source identifier: 2402.02868} }