@misc{indiciae6b4033faabe1, title = {Full Gradient DQN Reinforcement Learning: A Provably Convergent Scheme}, author = {K. E. Avrachenkov and V. S. Borkar and H. P. Dolhare and K. Patil}, year = {2021}, doi = {10.1007/978-3-030-76928-4\_10}, url = {https://arxiv.org/abs/2103.05981}, note = {Source identifier: 2103.05981} }