@misc{indiciaeecd2d7e68aa7, title = {Variational Reward Estimator Bottleneck: Learning Robust Reward Estimator for Multi-Domain Task-Oriented Dialog}, author = {Jeiyoon Park and Chanhee Lee and Kuekyeng Kim and Heuiseok Lim}, year = {2020}, url = {https://arxiv.org/abs/2006.00417}, note = {Source identifier: 2006.00417} }