@misc{indiciae231eb64233c6, title = {Polylogarithmic width suffices for gradient descent to achieve arbitrarily small test error with shallow ReLU networks}, author = {Ziwei Ji and Matus Telgarsky}, year = {2020}, url = {https://arxiv.org/abs/1909.12292}, note = {Source identifier: 1909.12292} }