@misc{indiciae55c1e7f9c1c9, title = {SAFE: Stable Alignment Finetuning with Entropy-Aware Predictive Control for Reinforcement Learning from Human Feedback (RLHF)}, author = {Dipan Maity}, year = {2026}, url = {https://arxiv.org/abs/2602.04651}, note = {Source identifier: 2602.04651} }