@misc{indiciae71e1ab3c8ee9, title = {Reinforcement Learning for Finite-Horizon Restless Multi-Armed Multi-Action Bandits}, author = {Guojun Xiong and Jian Li and Rahul Singh}, year = {2022}, url = {https://arxiv.org/abs/2109.09855}, note = {Source identifier: 2109.09855} }