@misc{indiciaecdb5e9c95330, title = {Reinforcement Learning from LLM Feedback to Counteract Goal Misgeneralization}, author = {Houda Nait El Barj and Theophile Sautory}, year = {2024}, url = {https://arxiv.org/abs/2401.07181}, note = {Source identifier: 2401.07181} }