@misc{indiciae1cdf7011f89a, title = {CRPO: A New Approach for Safe Reinforcement Learning with Convergence Guarantee}, author = {Tengyu Xu and Yingbin Liang and Guanghui Lan}, year = {2021}, url = {https://arxiv.org/abs/2011.05869}, note = {Source identifier: 2011.05869} }