@misc{indiciae2452bc2f346b, title = {Improved Regret Bound for Safe Reinforcement Learning via Tighter Cost Pessimism and Reward Optimism}, author = {Kihyun Yu and Duksang Lee and William Overman and Dabeen Lee}, year = {2024}, url = {https://arxiv.org/abs/2410.10158}, note = {Source identifier: 2410.10158} }