@inproceedings{324c511824544cd4b945f3e47405f0a9,
title = "Single channel voice separation for unknown number of speakers under reverberant and noisy settings",
abstract = "We present a unified network for voice separation of an unknown number of speakers. The proposed approach is composed of several separation heads optimized together with a speaker classification branch. The separation is carried out in the time domain, together with parameter sharing between all separation heads. The classification branch estimates the number of speakers while each head is specialized in separating a different number of speakers. We evaluate the proposed model under both clean and noisy reverberant settings. Results suggest that the proposed approach is superior to the baseline model by a significant margin. Additionally, we present a new noisy and reverberant dataset of up to five different speakers speaking simultaneously.",
keywords = "Source separation, Speaker classification, Speech processing",
author = "Chazan, {Shlomo E.} and Lior Wolf and Eliya Nachmani and Yossi Adi",
note = "Publisher Copyright: {\textcopyright} 2021 IEEE; 2021 IEEE International Conference on Acoustics, Speech, and Signal Processing, ICASSP 2021 ; Conference date: 06-06-2021 Through 11-06-2021",
year = "2021",
doi = "10.1109/ICASSP39728.2021.9413627",
language = "الإنجليزيّة",
series = "ICASSP, IEEE International Conference on Acoustics, Speech and Signal Processing - Proceedings",
publisher = "Institute of Electrical and Electronics Engineers Inc.",
pages = "3730--3734",
booktitle = "ICASSP, IEEE International Conference on Acoustics, Speech and Signal Processing - Proceedings",
address = "الولايات المتّحدة",
}