@misc{indiciae1baa13cefe26, title = {Safe Reinforcement Learning using Finite-Horizon Gradient-based Estimation}, author = {Juntao Dai and Yaodong Yang and Qian Zheng and Gang Pan}, year = {2024}, url = {https://arxiv.org/abs/2412.11138}, note = {Source identifier: 2412.11138} }