Thesis Open Access

A Wavenet for Music Source Separation

Francesc Lluís Salvadó


DCAT Export

<?xml version='1.0' encoding='utf-8'?>
<rdf:RDF xmlns:rdf="http://www.w3.org/1999/02/22-rdf-syntax-ns#" xmlns:adms="http://www.w3.org/ns/adms#" xmlns:cnt="http://www.w3.org/2011/content#" xmlns:dc="http://purl.org/dc/elements/1.1/" xmlns:dct="http://purl.org/dc/terms/" xmlns:dctype="http://purl.org/dc/dcmitype/" xmlns:dcat="http://www.w3.org/ns/dcat#" xmlns:duv="http://www.w3.org/ns/duv#" xmlns:foaf="http://xmlns.com/foaf/0.1/" xmlns:frapo="http://purl.org/cerif/frapo/" xmlns:geo="http://www.w3.org/2003/01/geo/wgs84_pos#" xmlns:gsp="http://www.opengis.net/ont/geosparql#" xmlns:locn="http://www.w3.org/ns/locn#" xmlns:org="http://www.w3.org/ns/org#" xmlns:owl="http://www.w3.org/2002/07/owl#" xmlns:prov="http://www.w3.org/ns/prov#" xmlns:rdfs="http://www.w3.org/2000/01/rdf-schema#" xmlns:schema="http://schema.org/" xmlns:skos="http://www.w3.org/2004/02/skos/core#" xmlns:vcard="http://www.w3.org/2006/vcard/ns#" xmlns:wdrs="http://www.w3.org/2007/05/powder-s#">
  <rdf:Description rdf:about="https://doi.org/10.5281/zenodo.1475940">
    <rdf:type rdf:resource="http://www.w3.org/ns/dcat#Dataset"/>
    <dct:type rdf:resource="http://purl.org/dc/dcmitype/Text"/>
    <dct:identifier rdf:datatype="http://www.w3.org/2001/XMLSchema#anyURI">https://doi.org/10.5281/zenodo.1475940</dct:identifier>
    <foaf:page rdf:resource="https://doi.org/10.5281/zenodo.1475940"/>
    <dct:creator>
      <rdf:Description>
        <rdf:type rdf:resource="http://xmlns.com/foaf/0.1/Agent"/>
        <foaf:name>Francesc Lluís Salvadó</foaf:name>
      </rdf:Description>
    </dct:creator>
    <dct:title>A Wavenet for Music Source Separation</dct:title>
    <dct:publisher>
      <foaf:Agent>
        <foaf:name>Zenodo</foaf:name>
      </foaf:Agent>
    </dct:publisher>
    <dct:issued rdf:datatype="http://www.w3.org/2001/XMLSchema#gYear">2018</dct:issued>
    <dct:issued rdf:datatype="http://www.w3.org/2001/XMLSchema#date">2018-08-31</dct:issued>
    <owl:sameAs rdf:resource="https://zenodo.org/record/1475940"/>
    <adms:identifier>
      <adms:Identifier>
        <skos:notation rdf:datatype="http://www.w3.org/2001/XMLSchema#anyURI">https://zenodo.org/record/1475940</skos:notation>
      </adms:Identifier>
    </adms:identifier>
    <dct:isVersionOf rdf:resource="https://doi.org/10.5281/zenodo.1475939"/>
    <dct:isPartOf rdf:resource="https://zenodo.org/communities/smc-master"/>
    <dct:description>&lt;p&gt;Currently, most successful source separation techniques use magnitude spectrograms as input, and are therefore by default discarding part of the signal: the phase. In order to avoid discarding potentially useful information, we propose an end-to-end learning model based on Wavenet for music source separation. As a result, the model we propose directly operates over the waveform, enabling, in that way, to consider any information available in the raw audio signal. Provided that the original Wavenet model operates sequentially (i.e., is not parallelizable and hence slow), in this work we make use of a discriminative non-causal adaptation of Wavenet capable to predict more than one sample at a time, thus permitting to overcome the undesirable time-complexity that the original Wavenet model has. Further, we investigate several data augmentation techniques and architectural changes to provide some insights on which are the most sensitive hyper-parameters for this family of Wavenet-like models. Our experimental results show that it is possible to approach the problem of music source separation in a end-to-end learning fashion, since our model performs on par with DeepConvSep, a state-of-the-art method based on processing magnitude spectrograms.&lt;/p&gt;</dct:description>
    <dct:accessRights rdf:resource="http://publications.europa.eu/resource/authority/access-right/PUBLIC"/>
    <dct:accessRights>
      <dct:RightsStatement rdf:about="info:eu-repo/semantics/openAccess">
        <rdfs:label>Open Access</rdfs:label>
      </dct:RightsStatement>
    </dct:accessRights>
    <dcat:distribution>
      <dcat:Distribution>
        <dct:license rdf:resource="http://creativecommons.org/licenses/by/4.0/legalcode"/>
        <dcat:accessURL rdf:resource="https://doi.org/10.5281/zenodo.1475940"/>
      </dcat:Distribution>
    </dcat:distribution>
  </rdf:Description>
</rdf:RDF>
250
147
views
downloads
All versions This version
Views 250250
Downloads 147147
Data volume 222.8 MB222.8 MB
Unique views 223223
Unique downloads 131131

Share

Cite as