@misc{indiciae7029a8214a8e, title = {Smoothed functional-based gradient algorithms for off-policy reinforcement learning: A non-asymptotic viewpoint}, author = {Nithia Vijayan and Prashanth L. A}, year = {2024}, url = {https://arxiv.org/abs/2101.02137}, note = {Source identifier: 2101.02137} }