@misc{indiciae95ae1c101b9d, title = {The Local Optimality of Reinforcement Learning by Value Gradients, and its Relationship to Policy Gradient Learning}, author = {Michael Fairbank and Eduardo Alonso}, year = {2011}, url = {https://arxiv.org/abs/1101.0428}, note = {Source identifier: 1101.0428} }