@misc{indiciaededd33278998, title = {Enhancing Safety in Reinforcement Learning with Human Feedback via Rectified Policy Optimization}, author = {Xiyue Peng and Hengquan Guo and Jiawei Zhang and Dongqing Zou and Ziyu Shao and Honghao Wei and Xin Liu}, year = {2025}, url = {https://arxiv.org/abs/2410.19933}, note = {Source identifier: 2410.19933} }