@misc{indiciae5823e97e8fb5, title = {Optimization for Reinforcement Learning: From Single Agent to Cooperative Agents}, author = {Donghwan Lee and Niao He and Parameswaran Kamalaruban and Volkan Cevher}, year = {2019}, doi = {10.1109/msp.2020.2976000}, url = {https://arxiv.org/abs/1912.00498}, note = {Source identifier: 1912.00498} }