@misc{indiciae5d6f5c1435c9, title = {Reinforcement Learning from Human Feedback with High-Confidence Safety Constraints}, author = {Yaswanth Chittepu and Blossom Metevier and Will Schwarzer and Austin Hoag and Scott Niekum and Philip S. Thomas}, year = {2025}, url = {https://arxiv.org/abs/2506.08266}, note = {Source identifier: 2506.08266} }