@misc{indiciae3d42fe5226ce, title = {Switch-based Active Deep Dyna-Q: Efficient Adaptive Planning for Task-Completion Dialogue Policy Learning}, author = {Yuexin Wu and Xiujun Li and Jingjing Liu and Jianfeng Gao and Yiming Yang}, year = {2018}, url = {https://arxiv.org/abs/1811.07550}, note = {Source identifier: 1811.07550} }