@misc{indiciaed8bc0de7bbd3, title = {Towards Safe Reinforcement Learning Using NMPC and Policy Gradients: Part I - Stochastic case}, author = {Sebastien Gros and Mario Zanon}, year = {2019}, url = {https://arxiv.org/abs/1906.04057}, note = {Source identifier: 1906.04057} }