@misc{indiciaed3567297641f, title = {Reinforcement Learning from Reflective Feedback (RLRF): Aligning and Improving LLMs via Fine-Grained Self-Reflection}, author = {Kyungjae Lee and Dasol Hwang and Sunghyun Park and Youngsoo Jang and Moontae Lee}, year = {2024}, url = {https://arxiv.org/abs/2403.14238}, note = {Source identifier: 2403.14238} }