@misc{indiciae4817dda3ec05, title = {Successive Convex Approximation Based Off-Policy Optimization for Constrained Reinforcement Learning}, author = {Chang Tian and An Liu and Guang Huang and Wu Luo}, year = {2021}, doi = {10.1109/tsp.2022.3158737}, url = {https://arxiv.org/abs/2105.12545}, note = {Source identifier: 2105.12545} }