@inproceedings {PN3309,
	title = {Teaching a Humanoid Robot: Headset-Free Speech Interaction for Audio-Visual Association Learning},
	author = {Martin Ernst Heckmann AND Holger Brandl AND Jens Schm{\"u}dderich AND Xavier Domont AND Bram Bolder AND Inna Mikhailova AND Herbert Jan{\ss}en AND Michael Gienger AND Achim Bendig AND Tobias Rodemann AND Mark Dunn AND Frank Joublin AND Christian Goerick},
	year = {2009},
	abstract = {Based on inspirations from infant development we present a system which learns associations between acoustic labels and visual representations in interaction with its tutor. The system is integrated with a humanoid robot. Except for a few trigger phrases to start learning all acoustical representations are learned online and in interaction. Similar, for the visual domain the clusters are not predefined and fully learned online. In contrast to other interactive systems the interaction with the acoustic environment is solely based on the two microphones mounted on the robots head. In this paper we give an overview on all key elements of the system and focus on the challenges arising from the headset-free learning of speech labels. In particular we present a mechanism for auditory attention integrating bottom-up and top-down information for the segmentation of the acoustic stream. The performance of the system is evaluated based on offline tests of individual parts of the system and an analysis of the accompanying video showing the system in interaction.},
	publisher = {IEEE},
	booktitle = {Proceedings of the 18th IEEE International Symposium on Robot and Human Interactive Communication (RO-MAN)},
	address = {Toyama, Japan}
}
