@misc{indiciae6c5efc517526, title = {Disentangled Speech Representation Learning Based on Factorized Hierarchical Variational Autoencoder with Self-Supervised Objective}, author = {Yuying Xie and Thomas Arildsen and Zheng-Hua Tan}, year = {2022}, doi = {10.1109/mlsp52302.2021.9596320}, url = {https://arxiv.org/abs/2204.02166}, note = {Source identifier: 2204.02166} }