@misc{indiciae6a84b3de5d5a, title = {Reinforcement Learning of Speech Recognition System Based on Policy Gradient and Hypothesis Selection}, author = {Taku Kato and Takahiro Shinozaki}, year = {2017}, url = {https://arxiv.org/abs/1711.03689}, note = {Source identifier: 1711.03689} }