@misc{indiciaeabb018f4894b, title = {Risk-Sensitive Deep RL: Variance-Constrained Actor-Critic Provably Finds Globally Optimal Policy}, author = {Han Zhong and Xun Deng and Ethan X. Fang and Zhuoran Yang and Zhaoran Wang and Runze Li}, year = {2023}, url = {https://arxiv.org/abs/2012.14098}, note = {Source identifier: 2012.14098} }