@misc{indiciae895b1f7ed547, title = {Quantum Reinforcement Learning via Policy Iteration}, author = {El Amine Cherrat and Iordanis Kerenidis and Anupam Prakash}, year = {2022}, doi = {10.1007/s42484-023-00116-1}, url = {https://arxiv.org/abs/2203.01889}, note = {Source identifier: 2203.01889} }