<?xml version="1.0" encoding="utf-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD JATS (Z39.96) Journal Publishing DTD v1.1d3 20150301//EN" "http://jats.nlm.nih.gov/publishing/1.1d3/JATS-journalpublishing1.dtd">
<article article-type="research-article" dtd-version="1.1d3" xml:lang="en" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink">
<front>
<journal-meta>
<journal-id journal-id-type="nlm-ta">PLoS ONE</journal-id>
<journal-id journal-id-type="publisher-id">plos</journal-id>
<journal-id journal-id-type="pmc">plosone</journal-id>
<journal-title-group>
<journal-title>PLOS ONE</journal-title>
</journal-title-group>
<issn pub-type="epub">1932-6203</issn>
<publisher>
<publisher-name>Public Library of Science</publisher-name>
<publisher-loc>San Francisco, CA USA</publisher-loc>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.1371/journal.pone.0162176</article-id>
<article-id pub-id-type="publisher-id">PONE-D-16-00940</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Research Article</subject>
</subj-group>
<subj-group subj-group-type="Discipline-v3"><subject>Biology and life sciences</subject><subj-group><subject>Neuroscience</subject><subj-group><subject>Cognitive science</subject><subj-group><subject>Cognitive psychology</subject><subj-group><subject>Learning</subject></subj-group></subj-group></subj-group></subj-group></subj-group><subj-group subj-group-type="Discipline-v3"><subject>Biology and life sciences</subject><subj-group><subject>Psychology</subject><subj-group><subject>Cognitive psychology</subject><subj-group><subject>Learning</subject></subj-group></subj-group></subj-group></subj-group><subj-group subj-group-type="Discipline-v3"><subject>Social sciences</subject><subj-group><subject>Psychology</subject><subj-group><subject>Cognitive psychology</subject><subj-group><subject>Learning</subject></subj-group></subj-group></subj-group></subj-group><subj-group subj-group-type="Discipline-v3"><subject>Biology and life sciences</subject><subj-group><subject>Neuroscience</subject><subj-group><subject>Learning and memory</subject><subj-group><subject>Learning</subject></subj-group></subj-group></subj-group></subj-group><subj-group subj-group-type="Discipline-v3"><subject>Social sciences</subject><subj-group><subject>Linguistics</subject><subj-group><subject>Language acquisition</subject></subj-group></subj-group></subj-group><subj-group subj-group-type="Discipline-v3"><subject>Biology and life sciences</subject><subj-group><subject>Organisms</subject><subj-group><subject>Animals</subject><subj-group><subject>Vertebrates</subject><subj-group><subject>Amniotes</subject><subj-group><subject>Mammals</subject><subj-group><subject>Dogs</subject></subj-group></subj-group></subj-group></subj-group></subj-group></subj-group></subj-group><subj-group subj-group-type="Discipline-v3"><subject>Social sciences</subject><subj-group><subject>Linguistics</subject><subj-group><subject>Lexicons</subject></subj-group></subj-group></subj-group><subj-group subj-group-type="Discipline-v3"><subject>Biology and life sciences</subject><subj-group><subject>Neuroscience</subject><subj-group><subject>Cognitive science</subject><subj-group><subject>Cognitive psychology</subject><subj-group><subject>Language</subject></subj-group></subj-group></subj-group></subj-group></subj-group><subj-group subj-group-type="Discipline-v3"><subject>Biology and life sciences</subject><subj-group><subject>Psychology</subject><subj-group><subject>Cognitive psychology</subject><subj-group><subject>Language</subject></subj-group></subj-group></subj-group></subj-group><subj-group subj-group-type="Discipline-v3"><subject>Social sciences</subject><subj-group><subject>Psychology</subject><subj-group><subject>Cognitive psychology</subject><subj-group><subject>Language</subject></subj-group></subj-group></subj-group></subj-group><subj-group subj-group-type="Discipline-v3"><subject>Biology and life sciences</subject><subj-group><subject>Organisms</subject><subj-group><subject>Animals</subject><subj-group><subject>Vertebrates</subject><subj-group><subject>Amniotes</subject><subj-group><subject>Mammals</subject><subj-group><subject>Bats</subject></subj-group></subj-group></subj-group></subj-group></subj-group></subj-group></subj-group><subj-group subj-group-type="Discipline-v3"><subject>Biology and life sciences</subject><subj-group><subject>Organisms</subject><subj-group><subject>Animals</subject><subj-group><subject>Vertebrates</subject><subj-group><subject>Amniotes</subject><subj-group><subject>Mammals</subject><subj-group><subject>Cats</subject></subj-group></subj-group></subj-group></subj-group></subj-group></subj-group></subj-group><subj-group subj-group-type="Discipline-v3"><subject>Social sciences</subject><subj-group><subject>Linguistics</subject><subj-group><subject>Phonology</subject></subj-group></subj-group></subj-group></article-categories>
<title-group>
<article-title>What Homophones Say about Words</article-title>
<alt-title alt-title-type="running-head">What Homophones Say about Words</alt-title>
</title-group>
<contrib-group>
<contrib contrib-type="author" corresp="yes" xlink:type="simple">
<name name-style="western">
<surname>Dautriche</surname>
<given-names>Isabelle</given-names>
</name>
<xref ref-type="corresp" rid="cor001">*</xref>
<xref ref-type="aff" rid="aff001"/>
</contrib>
<contrib contrib-type="author" xlink:type="simple">
<name name-style="western">
<surname>Chemla</surname>
<given-names>Emmanuel</given-names>
</name>
<xref ref-type="aff" rid="aff001"/>
</contrib>
</contrib-group>
<aff id="aff001"><addr-line>Laboratoire de Sciences Cognitives et Psycholinguistique, (DEC-ENS/EHESS/CNRS), Paris, France</addr-line></aff>
<contrib-group>
<contrib contrib-type="editor" xlink:type="simple">
<name name-style="western">
<surname>Allen</surname>
<given-names>Philip</given-names>
</name>
<role>Editor</role>
<xref ref-type="aff" rid="edit1"/>
</contrib>
</contrib-group>
<aff id="edit1"><addr-line>University of Akron, UNITED STATES</addr-line></aff>
<author-notes>
<fn fn-type="conflict" id="coi001">
<p>The authors have declared that no competing interests exist.</p>
</fn>
<fn fn-type="con">
<p><list list-type="simple"><list-item><p><bold>Conceived and designed the experiments:</bold> ID EC.</p></list-item> <list-item><p><bold>Performed the experiments:</bold> ID.</p></list-item> <list-item><p><bold>Analyzed the data:</bold> ID.</p></list-item> <list-item><p><bold>Contributed reagents/materials/analysis tools:</bold> ID EC.</p></list-item> <list-item><p><bold>Wrote the paper:</bold> ID EC.</p></list-item></list>
</p>
</fn>
<corresp id="cor001">* E-mail: <email xlink:type="simple">isabelle.dautriche@gmail.com</email></corresp>
</author-notes>
<pub-date pub-type="epub">
<day>1</day>
<month>9</month>
<year>2016</year>
</pub-date>
<pub-date pub-type="collection">
<year>2016</year>
</pub-date>
<volume>11</volume>
<issue>9</issue>
<elocation-id>e0162176</elocation-id>
<history>
<date date-type="received">
<day>9</day>
<month>1</month>
<year>2016</year>
</date>
<date date-type="accepted">
<day>18</day>
<month>8</month>
<year>2016</year>
</date>
</history>
<permissions>
<copyright-year>2016</copyright-year>
<copyright-holder>Dautriche, Chemla</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/" xlink:type="simple">
<license-p>This is an open access article distributed under the terms of the <ext-link ext-link-type="uri" xlink:href="http://creativecommons.org/licenses/by/4.0/" xlink:type="simple">Creative Commons Attribution License</ext-link>, which permits unrestricted use, distribution, and reproduction in any medium, provided the original author and source are credited.</license-p>
</license>
</permissions>
<self-uri content-type="pdf" xlink:href="info:doi/10.1371/journal.pone.0162176"/>
<abstract>
<p>The number of potential meanings for a new word is astronomic. To make the word-learning problem tractable, one must restrict the hypothesis space. To do so, current word learning accounts often incorporate constraints about cognition or about the mature lexicon directly in the learning device. We are concerned with the convexity constraint, which holds that concepts (privileged sets of entities that we think of as “coherent”) do not have gaps (if A and B belong to a concept, so does any entity “between” A and B). To leverage from it a linguistic constraint, learning algorithms have percolated this constraint from concepts, to word forms: some algorithms rely on the possibility that word forms are associated with convex sets of objects. Yet this does have to be the case: homophones are word forms associated with two separate words and meanings. Two sets of experiments show that when evidence suggests that a novel label is associated with a disjoint (non-convex) set of objects, either a) because there is a gap in conceptual space between the learning exemplars for a given word or b) because of the intervention of other lexical items in that gap, adults prefer to postulate homophony, where a single word form is associated with two separate words and meanings, rather than inferring that the word could have a disjunctive, discontinuous meaning. These results about homophony must be integrated to current word learning algorithms. We conclude by arguing for a weaker specialization of word learning algorithms, which too often could miss important constraints by focusing on a restricted empirical basis (e.g., non-homophonous content words).</p>
</abstract>
<funding-group>
<award-group id="award001">
<funding-source>
<institution-wrap>
<institution-id institution-id-type="funder-id">http://dx.doi.org/10.13039/501100000781</institution-id>
<institution>European Research Council</institution>
</institution-wrap>
</funding-source>
<award-id>313610</award-id>
<principal-award-recipient>
<name name-style="western">
<surname>Chemla</surname>
<given-names>Emmanuel</given-names>
</name>
</principal-award-recipient>
</award-group>
<award-group id="award002">
<funding-source>
<institution-wrap>
<institution-id institution-id-type="funder-id">http://dx.doi.org/10.13039/501100001665</institution-id>
<institution>Agence Nationale de la Recherche</institution>
</institution-wrap>
</funding-source>
<award-id>10-IDEX-0001-02</award-id>
</award-group>
<award-group id="award003">
<funding-source>
<institution-wrap>
<institution-id institution-id-type="funder-id">http://dx.doi.org/10.13039/501100001665</institution-id>
<institution>Agence Nationale de la Recherche</institution>
</institution-wrap>
</funding-source>
<award-id>10-LABX-0087</award-id>
</award-group>
<award-group id="award004">
<funding-source>
<institution-wrap>
<institution-id institution-id-type="funder-id">http://dx.doi.org/10.13039/501100006021</institution-id>
<institution>Direction Générale de lArmement</institution>
</institution-wrap>
</funding-source>
<principal-award-recipient>
<name name-style="western">
<surname>Dautriche</surname>
<given-names>Isabelle</given-names>
</name>
</principal-award-recipient>
</award-group>
<funding-statement>ID was supported by a Graduate Fellowship from the Direction Générale de l’Armement (PhD program Frontières du Vivant). The research was supported by grants from the European Research Council under FP/2007-2013-ERC n°313610 to EC, and from the Agence Nationale de la Recherche ANR-10-IDEX-0001-02, ANR-10-LABX-0087 for the facilities provided by the Departement d'Etudes Cognitives. The funders had no role in study design, data collection and analysis, decision to publish, or preparation of the manuscript.</funding-statement>
</funding-group>
<counts>
<fig-count count="6"/>
<table-count count="0"/>
<page-count count="19"/>
</counts>
<custom-meta-group>
<custom-meta id="data-availability">
<meta-name>Data Availability</meta-name>
<meta-value>All files are available from the OSF database (url <ext-link ext-link-type="uri" xlink:href="https://osf.io/u473e/" xlink:type="simple">https://osf.io/u473e/</ext-link>).</meta-value>
</custom-meta>
</custom-meta-group>
</article-meta>
</front>
<body>
<sec id="sec001" sec-type="intro">
<title>Introduction</title>
<p>Learning the word “cat” implies associating the sequence of sounds /kaet/ to the set of all cats and only cats. Quite generally one description of the meaning of a content word is its “extension”, i.e. the set of all entities to which that word refers (an idea discussed in detail in the tradition of formal semantics at least since [<xref ref-type="bibr" rid="pone.0162176.ref001">1</xref>]). But language learners need to infer the extension of a word based on a set of exemplars that surely do not exhaust that extension. The underlying inference problem would be unsolvable without prior knowledge, most notably some that could constrain the hypothesis space, which is the set of potential meanings for words (e.g., [<xref ref-type="bibr" rid="pone.0162176.ref002">2</xref>]–[<xref ref-type="bibr" rid="pone.0162176.ref005">5</xref>]; and [<xref ref-type="bibr" rid="pone.0162176.ref006">6</xref>] for a formal proof).</p>
<p>One way in which learners may reduce their hypothesis space is by privileging some meanings over others. For instance, toddlers and preschoolers prefer to extend a novel word (e.g., assume “blicket” is first associated with a dog) to an object of the same kind (e.g., a cat) rather than to an object of a different kind (e.g., a bone) (e.g., [<xref ref-type="bibr" rid="pone.0162176.ref007">7</xref>], [<xref ref-type="bibr" rid="pone.0162176.ref008">8</xref>]; see also the “shape bias”, showing that infants extend a label on the basis of the shape, [<xref ref-type="bibr" rid="pone.0162176.ref009">9</xref>]).</p>
<p>This follows if learners assume that those concepts that have words associated with them, are <italic>convex</italic> (e.g., [<xref ref-type="bibr" rid="pone.0162176.ref010">10</xref>]). A concept is convex if its members form a group that share a common set of properties that holds them to be contiguous in conceptual space (e.g., [<xref ref-type="bibr" rid="pone.0162176.ref010">10</xref>], see also [<xref ref-type="bibr" rid="pone.0162176.ref011">11</xref>] for the idea of “conceptual coherence”). For instance, the category <sc>dog or bone</sc> is not a possible concept because it does not form a coherent class of objects. Thus, if concepts are expected to be convex and words label concepts, learners may more readily extend the extension of a word (e.g., “blicket” designing a dog) to neighboring objects in conceptual space (cat) rather than to more distant objects (bone).</p>
<p>Thus, current experimental results provide evidence that a convexity constraint guides learners' inferences about word meanings: if A and B can be labeled using the sound /bliket/, then all objects falling “in between” A and B in conceptual space can also be labeled with the sound /bliket/. But these experiments do not distinguish between words and word forms, hence it is unclear whether this constraint applies at the level of the words or at the level of the word forms.</p>
<p>Yet words and word forms can be dissociated. A homophone is a phonological form associated arbitrarily with <italic>several</italic> meanings (contrary to polysemy, see e.g., Bréal, 1904), which together form a discontinuous set in conceptual space. For instance, the English word form “bat” applies both to the convex concept of <sc>animal bats</sc> and to the convex concept of <sc>baseball bats</sc>, but, regardless of how the conceptual space is constructed, not all intervening objects sharing a common property of animal bats and baseball bats count as bats.</p>
<p>Although the domain of application of the convexity constraint, words or word forms, is rarely specified explicitly, all current models of word learning (e.g., [<xref ref-type="bibr" rid="pone.0162176.ref012">12</xref>], [<xref ref-type="bibr" rid="pone.0162176.ref013">13</xref>], [<xref ref-type="bibr" rid="pone.0162176.ref014">14</xref>]) practically implement a convexity constraint at the level of word forms. It is a virtue of these word learning models that they work at the level of word forms, because this is the visible layer of the input, learners hear word forms, not words. But note that as a consequence of this implementation, these accounts mechanically predict that when encountering a word form that applies to animal bats and baseball bats, English learners should conclude that <italic>bat</italic> applies to any intervening object, as would words that apply very broadly, such as “thing” or “stuff”. The very existence of homophony in human languages thus shows that learners do not adhere blindly to a convexity constraint on word forms. In sum, learners make inferences about the meaning of words, based on the occurrence of some word forms across different situations, We will show how learners rely on the (hidden) <italic>word</italic> level of representation to master a lexicon and how they may capitalize on a proper convexity constraint at this level to learn homophones.</p>
<p>Concretely, our point of departure will be work by Xu and Tenenbaum (2007) [<xref ref-type="bibr" rid="pone.0162176.ref012">12</xref>]. One advantage of their study is that it implements the convexity constraint on word forms in a predictive model, but it also provides the means to test it in a non-circular way. To do so, they first gathered similarity judgments between pairs of objects, and inferred a tree-structure over the whole set. This tree structure represents the taxonomy between the objects: different dogs are close together and form a subtree, mammals form a (bigger) subtree, etc. Such a hypothesis space reflects the taxonomic assumption [<xref ref-type="bibr" rid="pone.0162176.ref005">5</xref>] that requires words to label the nodes of a tree-structured hierarchy of natural concepts, in line with developmental data (e.g., [<xref ref-type="bibr" rid="pone.0162176.ref004">4</xref>], [<xref ref-type="bibr" rid="pone.0162176.ref005">5</xref>], [<xref ref-type="bibr" rid="pone.0162176.ref007">7</xref>], [<xref ref-type="bibr" rid="pone.0162176.ref008">8</xref>]. Crucially, Xu and Tenenbaum (2007) [<xref ref-type="bibr" rid="pone.0162176.ref012">12</xref>] used this structured conceptual space to test a model of word learning according to which the extension inferred for a given word label should be a set of objects with no gap in conceptual space and which minimally includes all exemplars. Thus, <italic>intervening</italic> objects, i.e., objects that are <italic>in between</italic> two learning exemplars, are defined as all objects in the smallest subtree that includes both exemplars (their convex hull). Accordingly, the authors demonstrate that, when exposed to a set of learning exemplars, participants extend the exemplars’ label to all intervening objects belonging to the smallest subtree that contains all these exemplars. For example, when presented with three Dalmatians as exemplars for a new word “fep”, adults readily extend “fep” to the set of all Dalmatians; would they be presented with a Dalmatian, a Labrador and a German-shepherd for the word “fep”, they would extend the label to the set of all dogs. In other words, participants pick the smallest generalization that satisfies the convexity constraint on word forms.</p>
<p>The present study explores the situations that lead language learners to postulate homophony for a new word using the word learning paradigm used by Xu and Tenenbaum (2007) [<xref ref-type="bibr" rid="pone.0162176.ref012">12</xref>]. In Experiment 1, we manipulate two factors that should invite learners to favor a homophone interpretation of a novel label:</p>
<list list-type="order">
<list-item><p>The <italic>size of the gap</italic>, in conceptual space, that separates different learning exemplars of a given word. To learn a homophone, language learners are exposed to a discrete set of learning exemplars. For instance, for the word <italic>bat</italic>, they would observe several animal-bats and several baseball bats. However if the underlying true concept were the broad category that encompasses animal-bats, baseball-bats and all intervening objects (e.g., “thing”), then presumably learners would not observe exemplars confined to two corners of this set. Rather, they would observe a <italic>set</italic> of learning exemplars randomly (uniformly) sampled from the broad category. Observing exemplars clustered at two distant positions in the hypothesis space, i.e., observing a large gap between the exemplars may boost the likelihood that the exemplars are sampled from two independent categories, favoring a homophone interpretation.</p></list-item>
<list-item><p><italic>The intervention of other lexical items in that gap</italic>. Evidence for homophony may also come from other words in the lexicon. There has been much evidence that words and their underlying concepts mutually constrain each other. For instance, language learners assume that words do not overlap in meaning (the “mutual exclusivity effect”; e.g., [<xref ref-type="bibr" rid="pone.0162176.ref015">15</xref>]). Having evidence that an additional label point towards an intervening region of the conceptual space (e.g., between animal-bat and baseball bats) may help learners discover more subtle configurations about how words map onto meanings.</p></list-item>
</list>
<p>Our results show that participants refrain from associating a label to a broad concept encompassing all the exemplars. Yet it does not entail that learners postulate homophony in these cases: Learners could have accepted that a word map onto a discontinuous concept (e.g., <sc>dog or bone)</sc> therefore violating concept convexity. We address this question more directly in Experiment 2. All in all, our results suggest that the effects documented in Experiment 1 are the footprints of homophony: Learners prefer to associate a single word form to several words and associated convex concepts, thus preserving concept convexity at the expense of word form convexity. This shows that current accounts of word learning face new challenges when incorporating homophony into the picture and that homophony can reveal (some of) the existing constraints learners deploy while learning words.</p>
</sec>
<sec id="sec002">
<title>Experiment 1: Gap in Conceptual Space and Overall Structure of the Lexicon</title>
<p>We used a word learning <italic>paradigm à</italic> la Xu and Tenenbaum (2007) [<xref ref-type="bibr" rid="pone.0162176.ref012">12</xref>]: participants were exposed to a new label through a couple of learning exemplars and asked whether the label should be extended to test items. We introduced a) a large gap in conceptual space between learning exemplars b) an intervening exemplar with a different label in that gap. We predicted that these two manipulations would lead to a breaking point after which participants would violate a convexity constraint on word forms, i.e., exclude items in the gap from the extension of the label.</p>
<sec id="sec003" sec-type="materials|methods">
<title>Method</title>
<sec id="sec004">
<title>Ethic Statement</title>
<p>All research was approved by the Comité d'Ethique de la Recherche en Santé (2013/46). Following the committee's recommendations, prior to accepting to participate in the online studies, participants were presented with the informed consent document and instructions stating that by clicking “Agree” they indicated their consent to participate in the study.</p>
</sec>
<sec id="sec005">
<title>Participants</title>
<p>One hundred and five adults were recruited through Amazon’s Mechanical Turk (45 females; M = 33 years; 102 native speakers of English) and compensated $0.50 for their participation. We excluded participants for lack of engagement in the task (criterion: participants who selected no test item in more than 50% of the “attractive” trials, in which at least 3 items should have been selected, see below; <italic>n</italic> = 0 in Experiment 1A, <italic>n =</italic> 16 in Experiment 1B) and participated in both versions of the experiment or in a previous pilot version (<italic>n</italic> = 3 and 5). This resulted in 41 participants in Experiment 1A and 40 participants in Experiment 1B. Data collection was stopped when each of the experiment had at least 40 participants. The number of participants was established before data collection began.</p>
</sec>
<sec id="sec006">
<title>Procedure and display</title>
<p>Participants were tested online. They were instructed that they would be exposed to words from an alien language and would have to select images that correspond to those words. In the instructions, participants were shown an example of a trial with pictures and a label that would not appear during the test. In each trial, participants first saw 3 or 4 learning exemplars, presented as the combination of a picture and a sentence. The first three learning exemplars (referred to as <sc>le</sc>1, <sc>le</sc>2 and <sc>le</sc>3 below) were presented in random order and labeled with a novel word, e.g., <italic>blicket</italic>, via a prompt of the form “This is a blicket” underneath each of them. The fourth learning exemplar (<sc>le</sc>X below), if present, was labeled with another novel word highlighted in red, as in e.g., “This is a bosa” and was always the right-most exemplar. Once participants pressed a button “Show”, they would see a set of 4 pictures below the learning exemplars and be asked to select from these test items which one(s) could be labeled with the first novel word: <italic>“Do you see any other blicket(s)</italic>?<italic>”</italic> (see <xref ref-type="fig" rid="pone.0162176.g001">Fig 1A</xref>). They responded by clicking to select none, one or multiple test items. When a picture was selected, its frame became green. Participants could unselect their choice by clicking on it again. Once a response was validated, the set of selected pictures was recorded and the test continued to the next trial.</p>
<fig id="pone.0162176.g001" position="float">
<object-id pub-id-type="doi">10.1371/journal.pone.0162176.g001</object-id>
<label>Fig 1</label>
<caption>
<title/>
<p><bold>A) Screenshots from Experiment 1A.</bold> Participants first see the 3 learning exemplars for the word “blicket” and one optional learning exemplar for the word “bosa”. After pressing the “show” button they then see the test pictures and are asked to find the other blickets. Once the pictures are selected (green frame), participants submit their answers by pressing the “done” button. <bold>B) Schema of the structure of a trial in conceptual space.</bold> The first row of pictures corresponds to the learning exemplars (<sc>le</sc>1, <sc>le</sc>2, <sc>le</sc>3, <sc>le</sc>X) and the second row to the test items. The intervening item <sc>le</sc>X appeared only in half of the test trials (hence the parentheses).</p>
</caption>
<graphic mimetype="image" position="float" xlink:href="info:doi/10.1371/journal.pone.0162176.g001" xlink:type="simple"/>
</fig>
</sec>
<sec id="sec007">
<title>Conditions</title>
<p>Each participant saw 12 test trials and 10 filler trials.</p>
<p><italic>Test trials</italic>. The structure of test trials is represented schematically in <xref ref-type="fig" rid="pone.0162176.g001">Fig 1A</xref>, the key factor is how the learning exemplars (<sc>le</sc>1, <sc>le</sc>2, <sc>le</sc>3 and optionally <sc>le</sc>X) were spread in conceptual space (here a tree-structure) and how the test items were distributed between them. As shown in <xref ref-type="fig" rid="pone.0162176.g001">Fig 1B</xref>, there were two gaps between the exemplars: one small gap between <sc>le</sc>1 and <sc>le</sc>2 and one much larger gap between <sc>le</sc>2 and <sc>le</sc>3. Test items were picked somewhere in the middle of the first small gap (<italic>middle-small-gap</italic>), of the large gap (<italic>middle-large-gap</italic>), in the large gap but close to the corresponding exemplars (<italic>border-large-gap</italic>) or out of all the exemplars altogether (<italic>out</italic>).</p>
<p>Six of the test trials, “Gap trials”, were designed solely to test the effect of the size of a gap between learning exemplars. They displayed three learning exemplars (<sc>le</sc>1, <sc>le</sc>2, <sc>le</sc>3) associated with a to-be-learned label. According to the convexity constraint on word forms, participants should select all test items in the minimal subtree containing all learning exemplars, but we expected that participants would be willing to violate this constraint and exclude <italic>middle-large-gap</italic> (or not as much as <italic>middle-small-gap</italic>).</p>
<p>Another 6 test trials, “Gap+Intervention trials”, had a fourth learning exemplar with a secondary label (the <sc>le</sc>X <italic>bosa</italic> exemplar in <xref ref-type="fig" rid="pone.0162176.g001">Fig 1</xref>). The convexity constraint on word forms applies to single lexical entries and is in principle blind to the rest of the lexicon, but we expected that participants would select the <italic>middle-large-gap</italic> test item less in these trials with an intervening label than in the test trials without this intervening label.</p>
<p><italic>Filler trials</italic>. One filler trial was presented first so that participants could familiarize themselves with the task (with no particular indication of it however). Nine other fillers were randomly interspersed between the test trials. 6 “attractive” fillers were designed such that participants would select at least 3 of the 4 test pictures (3 of these filler trials contained three learning exemplars, all with the same label as in the Gap test trials, and 3 others included a fourth learning exemplar with a secondary label as in the Intervention test trials). 3 “repulsive” fillers implemented the opposite bias: participants were expected to select one or no test picture.</p>
</sec>
<sec id="sec008">
<title>Materials</title>
<p>Our stimuli relied on a set of to-be-learned labels and taxonomically organized objects.</p>
<p><italic>Labels</italic>. 28 phonotactically legal non-words of English were used for both experiments and were not repeated across trials.</p>
<p><italic>Objects in conceptual space</italic>. We tested participants on two sets of objects organized into drastically different taxonomic hierarchies: natural objects, with a similarity measure based on phylogenetic trees (Experiment 1A) and artificial objects constructed in a parametric fashion, so that a similarity measure between these objects can be defined in a canonical way (Experiment 1B; <xref ref-type="fig" rid="pone.0162176.g002">Fig 2</xref>). Objects from this artificial taxonomy do not exist such that the actual lexicon of our participants cannot influence our experimental results.</p>
<fig id="pone.0162176.g002" position="float">
<object-id pub-id-type="doi">10.1371/journal.pone.0162176.g002</object-id>
<label>Fig 2</label>
<caption>
<title>Stimuli of Experiment 1B.</title>
<p>Examples of the artificial stimuli used in Experiment 1B, out of a set of 1024 possible unique combinations obtained from 5 parameters (core pattern, core pattern occurrences, size of the core pattern, number of radial lines, number of bumps in the radial lines) with 4 levels each.</p>
</caption>
<graphic mimetype="image" position="float" xlink:href="info:doi/10.1371/journal.pone.0162176.g002" xlink:type="simple"/>
</fig>
<p>One important difference with Xu and Tenenbaum’s [<xref ref-type="bibr" rid="pone.0162176.ref012">12</xref>] paradigm is that our conceptual space did not rely on subjective, experimentally-gathered similarity judgments, but rather on objective similarity measures: one based on the distance in the phylogenetic tree and the other based on the parameterization of the objects. Surely these measures are only a proxy for participants’ representation of the similarity relationships between the objects. Yet, any effect that can be detected from these imperfect objective measures will retrospectively validate that it is a good approximation of the underlying subjective measure. We describe the two sets of objects at the basis of Experiments 1A and 1B, their structure, and how our experimental conditions were obtained in each case in <xref ref-type="supplementary-material" rid="pone.0162176.s001">S1 Supplemental Material</xref>. The experimental material for both experiments is available at <ext-link ext-link-type="uri" xlink:href="https://osf.io/u473e/?view_only=33576a1ac18746b08d7e3fcc96e10e9a" xlink:type="simple">https://osf.io/u473e/?view_only=33576a1ac18746b08d7e3fcc96e10e9a</ext-link></p>
</sec>
<sec id="sec009">
<title>Presentation and trial generation</title>
<p>The order of the trials as well as the pairing between the labels and the set of learning exemplars was fully randomized and differed for each participant. All trials were generated automatically following the algorithmic constraints described in <xref ref-type="supplementary-material" rid="pone.0162176.s001">S1 Supplemental</xref> Material for each stimuli type.</p>
</sec>
<sec id="sec010">
<title>Statistical analysis</title>
<p>In a mixed logit regression [<xref ref-type="bibr" rid="pone.0162176.ref016">16</xref>], we modeled the selection of a test item (coded as 0 or 1) for each experiment (natural or artificial stimuli). Both models included two categorical predictors with their interaction: Test Item (<italic>middle-small-gap</italic>, <italic>border-large-gap</italic>, <italic>middle-large-gap</italic>, <italic>out</italic>) and Trial Type (Gap vs. Gap+Intervention) as well as a random intercept and random slopes for both Test Item and Trial Type for participants. We coded our predictors such that selection of <italic>middle-large-gap</italic> for Gap trials served as a baseline (unless otherwise mentioned) against which we compared a) responses to the other test items, b) the responses to <italic>middle-large-gap</italic> in Gap+Intervention trials.</p>
<p>All analyses were conducted using the lme4 package [<xref ref-type="bibr" rid="pone.0162176.ref017">17</xref>] of R.</p>
</sec>
</sec>
<sec id="sec011" sec-type="results">
<title>Results</title>
<p><xref ref-type="fig" rid="pone.0162176.g003">Fig 3</xref> reports the average proportion of selection of each test item by Trial Type (Gap vs. Gap+Intervention trials) and Experiment (1A or 1B).</p>
<fig id="pone.0162176.g003" position="float">
<object-id pub-id-type="doi">10.1371/journal.pone.0162176.g003</object-id>
<label>Fig 3</label>
<caption>
<title>Results of Experiment 1.</title>
<p>Proportion of choice of each test item averaged across Experiment 1A with natural objects (upper panel) and Experiment 1B with artificial objects (lower panel) for each trial type (Gap vs. Gap+Intervention trials). The x-axis follows (with some simplification) the structure in conceptual space: the position of the learning exemplars is indicated among the bars for the test items with the dashed lines. Error bars indicate standard errors of the mean.</p>
</caption>
<graphic mimetype="image" position="float" xlink:href="info:doi/10.1371/journal.pone.0162176.g003" xlink:type="simple"/>
</fig>
<p>For Gap trials (<xref ref-type="fig" rid="pone.0162176.g003">Fig 3Aa and 3Ab</xref>), we replicate the minimal category effect seen in previous results (i.e., [<xref ref-type="bibr" rid="pone.0162176.ref012">12</xref>]) showing that participants are more likely to select a test item belonging to the category which is minimally consistent with the exemplars (<italic>middle-small-gap</italic>, <italic>border-large-gap</italic>, <italic>middle-large-gap</italic>) than a test item outside of this category (<italic>out</italic>), both for Experiment 1A (β = -3.75, <italic>z</italic> = -12.20, <italic>p</italic> &lt; .001) and Experiment 1B (β = -5.17, <italic>z</italic> = -11.46, <italic>p</italic> &lt; .001; We Helmert-coded the predictor Test Item to compare the choice of <italic>out</italic> to the choice of the rest of the test items as a group). Crucially, the size of the gap between learning exemplars modulated participants’ responses. That is, participants selected <italic>middle-small-gap</italic> items more than <italic>middle-large-gap</italic> items both in Experiment 1A (<italic>M</italic><sub><italic>middle-large-gap</italic></sub> <italic>=</italic> 0.59, <italic>M</italic><sub><italic>middle-small-gap</italic></sub> <italic>=</italic> 0.98; β = 3.72, <italic>z</italic> = 7.53, <italic>p</italic> &lt; .001) and in Experiment 1B (<italic>M</italic><sub><italic>middle-large-gap</italic></sub> <italic>=</italic> 0.44, <italic>M</italic><sub><italic>middle-small-gap</italic></sub> <italic>=</italic> 0.94; β = 3.45, <italic>z</italic> = 10.38, <italic>p</italic> &lt; .001). Participants were sensitive to the distribution of the learning exemplars with natural stimuli but also with unfamiliar stimuli. This latter case shows that familiarity with the categories (e.g., mammals, carnivores, animals) and possible existing labels for them cannot fully explain the results.</p>
<p>For Intervention trials (<xref ref-type="fig" rid="pone.0162176.g003">Fig 3Ba and 3Bb</xref>), we first replicate the effect described above: participants were sensitive to the size of the gap between the exemplars, that is, they selected <italic>middle-small-gap more</italic> than <italic>middle-large-gap</italic> in Experiments 1A (<italic>M</italic><sub><italic>middle-large-gap</italic></sub> <italic>=</italic> 0.43, <italic>M</italic><sub><italic>middle-small-gap</italic></sub> <italic>=</italic> 0.97; β = 3.95, <italic>z</italic> = 9.75, <italic>p</italic> &lt; .001) and in Experiment 1B (<italic>M</italic><sub><italic>middle-large-gap</italic></sub> <italic>=</italic> 0.32; <italic>M</italic><sub><italic>middle-small-gap</italic></sub> <italic>=</italic> 0.82; β = 2.70, <italic>z</italic> = 10.32, <italic>p</italic> &lt; .001). Crucially, we expected that the presence of an intervening item would increase participants’ violation of a convexity constraint on word forms.</p>
<p>Indeed, in Experiment 1A, participants selected <italic>middle-large-gap</italic> less in Gap+Intervention trials than in Gap trials (β = -0.72, <italic>z</italic> = -3.80, <italic>p</italic> &lt; .001). Yet, the presence of an intervening lexical item did not affect the choice of any other test items (all <italic>ps</italic> &gt; 0.4) leading to an interaction effect: the difference between the selection rate of <italic>middle-small-gap</italic> and <italic>middle-large-gap</italic> was greater in Gap+Intervention trials than in Gap trials (β = 0.68, <italic>z</italic> = 2.51, <italic>p</italic> &lt; .01).</p>
<p>In Experiment 1B, participants similarly selected <italic>middle-large-gap</italic> less in Gap+Intervention trials than in Gap trials (β = -0.61, <italic>z</italic> = -2.40, <italic>p</italic> &lt; .05). But we should pause and note that the same was true for <italic>middle-small-gap</italic> items (β = -1.40, <italic>z</italic> = -3.96, <italic>p</italic> &lt; .001; here the intercept reflected selection of <italic>middle-small-gap</italic> in Gap trials). This was because the intervening exemplar <sc>le</sc>X was sometimes close to <italic>middle-small-gap</italic> (and even closer than it was to <italic>middle-large-gap</italic>), thus introducing an independent reason not to select <italic>middle-small-gap</italic> in these intervention trials.</p>
<p>Overall, we did observe that intervening labels block the extension of a word to the minimal category including all observed exemplars, even though this effect was polluted for artificial stimuli.</p>
</sec>
<sec id="sec012" sec-type="conclusions">
<title>Discussion</title>
<p>We highlighted two factors that disturb the association of a word form to the single category that minimally includes all its learning exemplars: a) the size of the gap between the exemplars; b) the presence of intervening lexical items. There may be three potential interpretations for these results:</p>
<list list-type="order">
<list-item><p>Participants associated a label to two meanings that <italic>each</italic> satisfies concept convexity. That is, participants postulated homophony, a non-immediate way to bind labels and concepts. Note however that we did not test whether participants provide evidence that subjects generalized from the more distant trained exemplar, <sc>le</sc>3, a point that will be addressed in the next experiment.</p></list-item>
<list-item><p>Participants associated a label with a set covering entities from several <italic>disjoint</italic> concepts (e.g., as in <sc>dog or bone</sc>), breaking thus concept convexity, either because meaning discontinuity is acceptable or because the specific experimental task that we propose led them to do so.</p></list-item>
<list-item><p>Participants did not associate the new word with a meaning at all. Instead, they simply went by similarity of the test items to the learning exemplars: they selected more the objects close to the exemplars (<italic>middle-small-gap</italic>) than to the objects further away from them (<italic>middle</italic>-<italic>large-gap</italic>). The role of the intervening label may be harder to account for in this view, but one may imagine some strategic effect such that if an object is close to some irrelevant object X, it will decrease the tendency to say that this object belongs to a set that was not said to contain X.</p></list-item>
</list>
<p>Experiment 2 was designed to distinguish between these three interpretations.</p>
</sec>
</sec>
<sec id="sec013">
<title>Experiment 2: Linguistic Manipulations</title>
<p>Homophones interact with linguistic constructions in a characteristic way. Zeugmas are the typical rhetorical device used to pun on the different senses of ambiguous words (e.g., [<xref ref-type="bibr" rid="pone.0162176.ref018">18</xref>], [<xref ref-type="bibr" rid="pone.0162176.ref019">19</xref>]) and have been extensively used as a test to distinguish words with an extension that covers a broad category from polysemous and ambiguous words (e.g., [<xref ref-type="bibr" rid="pone.0162176.ref018">18</xref>], [<xref ref-type="bibr" rid="pone.0162176.ref020">20</xref>]). Consider for instance “John and his driving license expired last Thursday” [<xref ref-type="bibr" rid="pone.0162176.ref018">18</xref>], where the verb “expire” has two distinct, but related, senses (i.e. “died” and “not valid anymore”). If the zeugmatic sentence is acceptable, it shows that the relevant word is polysemous or ambiguous (the two meanings are distinct) rather than vague (the boundary between meanings are indistinct).</p>
<p>Interestingly, zeugmas can be used to distinguish between a homophone, where a label applies to two convex concepts, and a word associated with a disjunctive meaning, where a label would apply to a disjoint concept. For instance, if “blicket” maps onto a disjunctive concept, such as <sc>dog or bone</sc>, it should be possible to use a plural sentence “these are two blickets” when pointing to a dog and a bone, while it would be zeugmatic to say “these are two bats”, pointing at one animal-bat and one baseball-bat. This is explained in a theory of homophones in which two words, with different meanings, share the same form: one cannot use a single phonological form to refer to both meanings at the same time. However, different <italic>tokens</italic> of the phonological form may pick out different meanings: it may therefore be more natural to say in a situation as above “This is a bat (pointing at the animal-bat), this is <underline><italic>also</italic></underline> a bat (pointing at the baseball-bat)”.</p>
<p>We will use these two constructions to test whether the effects we documented in the experiments above are the signatures of homophony. If participants postulated homophony, the <italic>plural</italic> zeugmatic construction, which is not compatible with homophony, should increase the tendency to form a single convex category encompassing all learning exemplars (as dictated by the convexity constraint over word forms in the absence of homophony), compared to the <italic>also</italic> construction. This would be evidence that participants did not postulate that a label could map onto a discontinuous concept and that our effects are not solely driven by similarity, since the similarity of the test items to the exemplars is held constant across the two linguistic constructions.</p>
<sec id="sec014" sec-type="materials|methods">
<title>Method</title>
<sec id="sec015">
<title>Participants</title>
<p>Ninety adults were recruited through Amazon Mechanical Turk (28 females; M = 30 years; 87 native speakers of English) and were compensated $0.50 for their participation. We excluded subjects who participated in both conditions of the experiment (<italic>n</italic> = 3). This resulted in 44 participants in the <italic>also</italic>-condition and 43 participants in the <italic>plural-</italic>condition. Data collection was stopped when each of the conditions had at least 40 participants. The number of participants was established before data collection began.</p>
</sec>
<sec id="sec016">
<title>Procedure and display</title>
<p>Similar to Experiment 1, except that each trial now included 4 learning exemplars and 6 test items.</p>
</sec>
<sec id="sec017">
<title>Conditions</title>
<p>Each participant saw 8 test trials and 16 filler trials.</p>
<p><italic>Test trials</italic>. As schematized in <xref ref-type="fig" rid="pone.0162176.g004">Fig 4</xref>, each test trial contained 4 learning exemplars (<sc>le</sc>1, <sc>le</sc>2 and <sc>le</sc>1’, <sc>le</sc>2’). We implemented symmetry in the distribution of learning exemplars such that there were two small gaps (between <sc>le</sc>1 and <sc>le</sc>2 and between <sc>le</sc>1’ and <sc>le</sc>2’) and one large gap (between the two pairs of exemplars). This distribution of exemplars in conceptual space may favor the construction of sharp boundaries over two disjoint categories (see also discussion of the “size principle” in the <xref ref-type="sec" rid="sec023">General Discussion</xref>). The position of the six test items is shown in <xref ref-type="fig" rid="pone.0162176.g004">Fig 4</xref>. Two test items were placed inside the small gaps (<italic>middle-small-gap</italic> and <italic>middle-small-gap</italic>’), two items just outside of the minimal subtrees S(<sc>le</sc>1,<sc>le</sc>2) and S(<sc>le</sc>1’,<sc>le</sc>2’) containing each pair of exemplars (<italic>border-large-gap</italic> and <italic>border-large-gap</italic>’), one item inside the large gap (<italic>middle</italic>-<italic>large-gap</italic>, either attached to S(<sc>le</sc>1,<sc>le</sc>2) or to S(<sc>le</sc>1’,<sc>le</sc>2’)) and one item outside of the minimal subtree containing all four learning exemplars (<italic>out</italic>).</p>
<fig id="pone.0162176.g004" position="float">
<object-id pub-id-type="doi">10.1371/journal.pone.0162176.g004</object-id>
<label>Fig 4</label>
<caption>
<title>Schema of the tree-structure of the items used in a trial for Experiment 2.</title>
<p>The boxed items correspond to the learning exemplars.</p>
</caption>
<graphic mimetype="image" position="float" xlink:href="info:doi/10.1371/journal.pone.0162176.g004" xlink:type="simple"/>
</fig>
<p>The 8 test trials were created according to the schema in <xref ref-type="fig" rid="pone.0162176.g004">Fig 4</xref>, but their mode of presentation differed across the two conditions. In the <italic>also</italic>-condition, the four learning exemplars were presented in pairs: the left pair was labeled with a given word (e.g., “These are two blickets”) and the right pair with the same word using <italic>also</italic> (e.g., “These are <underline><italic>also</italic></underline> two blickets”). In the <italic>plural</italic>-condition, the four learning exemplars were ordered in pairs as in the <italic>also</italic>-condition but the four exemplars were grouped together in a gray frame and labeled at once via a plural sentence (e.g., “These are four blickets”; see <xref ref-type="fig" rid="pone.0162176.g005">Fig 5</xref>).</p>
<fig id="pone.0162176.g005" position="float">
<object-id pub-id-type="doi">10.1371/journal.pone.0162176.g005</object-id>
<label>Fig 5</label>
<caption>
<title>Example of a test trial in Experiment 2.</title>
<p>Possible learning exemplars for a test trial as presented in 1) the plural-condition and 2) the also-condition.</p>
</caption>
<graphic mimetype="image" position="float" xlink:href="info:doi/10.1371/journal.pone.0162176.g005" xlink:type="simple"/>
</fig>
<p>We expected that participants would select the test items <italic>middle-large-gap</italic> and <italic>border-large-gap</italic> more in the plural-condition than in the also-condition, because homophony is less of an option while using the plural construction.</p>
<p><italic>Filler trials</italic>. 16 filler trials were interspersed, half of which were visually similar to the test trials of the plural-condition (<xref ref-type="fig" rid="pone.0162176.g005">Fig 5A</xref>) and the other half were visually similar to the test trials of also-condition (<xref ref-type="fig" rid="pone.0162176.g005">Fig 5B</xref>; but with a different label for the two pairs of objects and, of course, no <italic>also</italic> in the description).</p>
</sec>
<sec id="sec018">
<title>Material</title>
<p>We used the same set of objects as in Experiment 1A and the same labels.</p>
</sec>
<sec id="sec019">
<title>Presentation and trial generation</title>
<p>The experiment always started with 3 filler trials. All trials were generated pseudo-randomly following the constraints described in <xref ref-type="supplementary-material" rid="pone.0162176.s001">S1 Supplemental</xref> Material.</p>
</sec>
<sec id="sec020">
<title>Statistical analysis</title>
<p>As before, we modeled the selection of a test item in a mixed logit model including two categorical predictors with their interaction: Test Item (<italic>middle-small-gap</italic>, <italic>border-large-gap</italic>, <italic>middle-large-gap</italic>, <italic>out</italic>) and Linguistic Condition (Plural vs. Also) as well as a random intercept and a random slope for Test Item on participants. The selection of <italic>middle-large-gap</italic> in the plural-condition served as a baseline (unless otherwise mentioned).</p>
</sec>
</sec>
<sec id="sec021" sec-type="results">
<title>Results</title>
<p>The results are presented in <xref ref-type="fig" rid="pone.0162176.g006">Fig 6</xref>. The pairs (<sc>le</sc>1, <sc>le</sc>2) and (<sc>le</sc>1’, <sc>le</sc>2’) played symmetric roles, we accordingly collapsed responses for <italic>middle-small-gap</italic> and <italic>middle-small-gap’</italic> and responses for <italic>border-small-gap</italic> and <italic>border-small-gap’</italic> (practically ignoring the prime sign in the report).</p>
<fig id="pone.0162176.g006" position="float">
<object-id pub-id-type="doi">10.1371/journal.pone.0162176.g006</object-id>
<label>Fig 6</label>
<caption>
<title>Results of Experiment 2.</title>
<p>Proportion of selection of each test item averaged across A) The plural-condition using a linguistic construction discarding the possibility of homophony and B) the also-condition using a linguistic construction more suitable to homophony. Error bars indicate standard errors of the mean.</p>
</caption>
<graphic mimetype="image" position="float" xlink:href="info:doi/10.1371/journal.pone.0162176.g006" xlink:type="simple"/>
</fig>
<p>First, the results confirm the existence of a gap effect. Participants showed sensitivity to the sampling distribution of the exemplars, in that they selected more <italic>middle-small-gap</italic> than <italic>middle-large-gap</italic> in both the also-condition (β = 2.20, <italic>z</italic> = 11.26, <italic>p</italic> &lt; .001) and the plural-condition (β = 1.70, <italic>z</italic> = 8.59, <italic>p</italic> &lt; .001).</p>
<p>Interestingly, the minimum path (in terms of number of branches) between the learning exemplars <sc>le</sc>1 and <sc>le</sc>2 was smaller in Experiment 1A (mean for d(<sc>le</sc>1, <sc>le</sc>2) = 3.65), compared both to (<sc>le</sc>1, <sc>le</sc>2) and (<sc>le</sc>1’, <sc>le</sc>2’) in Experiment 2 (mean for d(<sc>le</sc>1,<sc>le</sc>2) = 7.16) (see <xref ref-type="supplementary-material" rid="pone.0162176.s001">S1 Supplemental</xref> material). Accordingly, we found a cross-experiment gap effect such that <italic>middle-small-gap</italic> was less selected in Experiment 2 (<italic>M =</italic> 0.75) than in Experiment 1A (<italic>M</italic> = 0.95).</p>
<p>Our critical expectation concerned the comparison between linguistic presentations. Test items in the gap between (<sc>le</sc>1, <sc>le</sc>2) and (<sc>le</sc>1’, <sc>le</sc>2’) were selected more often in the plural-condition than in the also-condition: this was true both for <italic>middle-large-gap</italic> (β = 0.73, <italic>z</italic> = 2.24, <italic>p</italic> &lt; .05) and <italic>border-large-gap</italic> (β = 0.64, <italic>z</italic> = 2.13, <italic>p</italic> &lt; .05), resulting in an interaction effect: the difference between the selection rate of <italic>middle-small-gap</italic> (serving as a baseline) and the combined selection rate of <italic>middle-large-gap</italic> and <italic>border-large-gap</italic> was greater in the plural-condition than in in the also-condition (β = 0.46, <italic>z</italic> = 2.02, <italic>p</italic> &lt; .05; We Helmert coded the predictor Test Item to compare the choice of <italic>middle-small-gap</italic> to the choice of the <italic>middle-large-gap</italic> and <italic>border-large-gap</italic> as a group).</p>
</sec>
<sec id="sec022" sec-type="conclusions">
<title>Discussion</title>
<p>When presented with a plural construction (e.g., “These are blickets”), participants were more likely to associate the word to a category that spans over all the exemplars than when they were presented with a construction compatible with homophony (e.g., “This is a blicket and this is <underline><italic>also</italic></underline> a blicket”). This effect suggests that the gap effect documented in Experiment 1 is the footprint of homophony mapping two words with the same phonological form onto two convex concepts and not of the association of a single word to a discontinuous category.</p>
<p>Certainly, in line with previous results (e.g.,[<xref ref-type="bibr" rid="pone.0162176.ref021">21</xref>], participants were guided in part by similarity: they extended a label more to an object close to the exemplars (<italic>middle-small-gap</italic>) than to an object further away (<italic>middle</italic>-<italic>large-gap</italic>) and they extended the label less to <italic>middle-small-gap</italic> in Experiment 2 than they did in Experiment 1 due to a greater distance between the learning exemplars in Experiment 2 (Note that this is also compatible with the size principle documented by Xu and Tenenbaum (2007) [<xref ref-type="bibr" rid="pone.0162176.ref012">12</xref>]: since the boundaries of the categories defined by the pairs of exemplars (<sc>le</sc>1, <sc>le</sc>2) and (<sc>le</sc>1’, <sc>le</sc>2’) were less sharp than in Experiment 1A, the correct level of generalization was more uncertain in Experiment 2 than in Experiment 1A.). Yet, such similarity effects cannot explain the main result of Experiment 2 since the linguistic manipulation is realized holding constant similarity relations among learning exemplars and test items. The amount of similarity-driven generalization in participants’ responses could be quantified (see [<xref ref-type="bibr" rid="pone.0162176.ref012">12</xref>] for a model comparison of rule- vs. similarity-based model), but it is sufficient for our purposes to note that it cannot account for the entirety of the present effects, which are driven by linguistic manipulations alone in Experiment 2. Note that while we demonstrate that participants are sensitive to the linguistic constructions in which the words enter, we cannot tell whether the <italic>also</italic> construction alters the results in one direction (towards homophony), or whether the <italic>plural</italic> construction pushes in the opposite direction (against homophony), or both.</p>
<p>One may also ask whether the plural/also effect is linked to the very specific linguistic constructions involved or whether it is merely driven by the visual, two-part presentation that co-varies with these constructions in our experiments. Importantly, a visual effect (i.e., a visual “zeugma”) would make the same point as a more specific linguistic construction effect: all that matters for our argument is that there is room for two tokens of the same phonological form, either because two tokens are indeed present, or simply because the presentation introduces different labeling events.</p>
</sec>
</sec>
<sec id="sec023">
<title>General Discussion</title>
<p>We documented two factors that reduce the tendency to map a phonological form onto a single, convex extension, an explicit or implicit assumption about learners in the most explicit implementations of word learning accounts: a) the size of the gap in conceptual space between learning exemplars; b) the presence of an intervening label for entities in that gap. These effects were modulated by linguistic manipulations coherent with the presence/absence of homophony. We submit that when encountering novel words in such situations, learners prefer to postulate homophony, whereby a word form does <italic>not</italic> adhere to a convexity constraint but uncovers two distinct words associated with two convex concepts, thus preserving concept convexity.</p>
<p>In the following we first discuss whether our results depend on the objective definition we adopted for conceptual space. Second, we show how the current study of homophony is relevant to current accounts of word learning broadly, and why other phenomena should be subjected to the same scrutiny. Third, we discuss our findings for a different, albeit similar, phenomenon in the lexicon: polysemy. Finally, our results are based on adult data only and we discuss their relevance for children in the process of learning their native language.</p>
<sec id="sec024">
<title>How to work with concepts</title>
<p>The study of homophony allowed us to examine the existence of a convexity constraint over concepts/words. We started by defining that a concept is convex if its members form a group that share a common set of properties that holds them to be contiguous in conceptual space. The notion of <italic>contiguity</italic> in conceptual space is meaningful as long as the conceptual space is equipped with a metrical structure such as the one we defined throughout this study. One worry however is that there may not be a stable metric between abstract entities across contexts [<xref ref-type="bibr" rid="pone.0162176.ref022">22</xref>], such that two entities can be made arbitrarily similar by changing the dimension under consideration. For instance, one may consider that a Ferrari and a VW Beatle are closer to one another than a Ferrari and a diamond, but this similarity relation may reverse if the context involves paying attention to the value of entities. As a result, one may wonder whether the notion of conceptual similarity may not be too fragile to sustain word learning. This is an important challenge that needs to be looked at carefully, as the current results are incorporated into word learning algorithms. We submit however that in the long run this variability could be controlled if we accept that some “conceptual dimensions” are privileged: they are more stable across contexts, by default, and infants are biased to pay more attention to them (e.g., [<xref ref-type="bibr" rid="pone.0162176.ref023">23</xref>]). For instance, animacy may be a property that is <italic>privileged</italic> in that sense, over say color, to categorize objects. It does not always have to be the case, but on average this will create the basis for a stable set of privileged features to (partly) provide a structure for conceptual space (see also [<xref ref-type="bibr" rid="pone.0162176.ref024">24</xref>] for the notion of <italic>ad hoc</italic> categories and [<xref ref-type="bibr" rid="pone.0162176.ref025">25</xref>], [<xref ref-type="bibr" rid="pone.0162176.ref026">26</xref>] for the idea of concept naturalness).</p>
<p>The next worry then is to decide how one can objectively assess what the actual, “privileged” metric in conceptual space is and derive testable predictions from there. We have seen several responses to this issue: For instance, Xu and Tenenbaum (2007) [<xref ref-type="bibr" rid="pone.0162176.ref012">12</xref>] gathered subjective judgments of similarities, independently from the categorization task. In our studies, we decided on a structure of conceptual space prior to using it for our test. Specifically, our notion of convexity relied on phylogenetic trees and on an arbitrary metric over a multi-dimensional space of visual features. The hope was that there would be a sufficiently good matching between these idealized conceptual spaces and what participants would actually take to be the relations between the relevant entities. Since participants had access to the entities only through visual representations, one may worry that we over-evaluated the chances that perceptual features could determine concepts. Perceptual features as such may not be the determinant of conceptual structure, since concepts may be defined by non-observable properties. Several developmental studies show that, indeed, children prefer to draw inferences based on category membership than inferences based on perceptual appearances (e.g., [<xref ref-type="bibr" rid="pone.0162176.ref027">27</xref>]–[<xref ref-type="bibr" rid="pone.0162176.ref029">29</xref>]). Nevertheless, we use perceptual similarity as a proxy to reflect conceptual structure and follow previous work in that respect (see [<xref ref-type="bibr" rid="pone.0162176.ref012">12</xref>], [<xref ref-type="bibr" rid="pone.0162176.ref030">30</xref>]–[<xref ref-type="bibr" rid="pone.0162176.ref033">33</xref>]). Interestingly, we note that young children may also use such a proxy in their earliest word meaning inferences [<xref ref-type="bibr" rid="pone.0162176.ref009">9</xref>], [<xref ref-type="bibr" rid="pone.0162176.ref034">34</xref>].</p>
<p>The current inquiry was based on the hope that conceptual space could be approximately circumscribed by objective or scientifically based properties (e.g., phylogenetic trees). There surely has to be <italic>some</italic> correlation between such an objectively based categorization and actual, subjective categorization (e.g., [<xref ref-type="bibr" rid="pone.0162176.ref035">35</xref>]). Most importantly, the fact that our results come out the right way suggests <italic>a posteriori</italic> that our simplifying hypotheses are acceptable to a sufficient degree: our results could not be obtained if our assumptions to approximate the underlying conceptual structure were inappropriate.</p>
</sec>
<sec id="sec025">
<title>Challenges for accounts of word learning</title>
<p>The above discussion only refers to concepts, not to words. A natural assumption is that one word would map onto one concept, but it does not have to be so. For instance, a word could map onto a set of concepts, as if there was a word meaning <sc>dog or bone</sc> (where <sc>dog</sc> and <sc>bone</sc> here are supposed to be disjoint concepts). Our study of homophones shows that this does not happen. Instead, when a word could potentially have such a disjunctive, discontinuous meaning, a homophone is created.</p>
<p>As we described in the introduction, convexity holds for words and not for word forms. This distinction was not and could not be investigated through previous experimental work since word learning studies focused on non-homophonous cases, for which word forms and words are confounded (e.g., [<xref ref-type="bibr" rid="pone.0162176.ref007">7</xref>], [<xref ref-type="bibr" rid="pone.0162176.ref008">8</xref>]). As a result, nothing prevented the computational implementation of these accounts to hardwire the efficient convexity constraint on word forms as a systemic property of the learning mechanism, but the current results force us to reflect on the status and consequence of this assumption. To date, word learning models capitalize on a convexity constrain on word forms, whereby word forms map onto a single meaning that ought to be convex in conceptual space, as this is the most natural version to implement for a learner who is only exposed to word forms. One possibility is that the convexity constrain on word forms is a mere heuristic, that would provide the right outcome in most word learning circumstances, where word forms and words are equivalent. Yet, the existence of homophony suggests that learners must be able to detect cues that would help them depart from their default over-simplification, and relax their word form convexity heuristic into a concept convexity constraint. Importantly, if the learner deploys such a heuristic, the system must be able to decide when to trigger and when to silence it and we documented the empirical cues under which this switch may happen. But this leaves open a fundamental question about when the system is capable of such a sophistication: should we assume that words and the abstract notion of word forms are both primitive of the learning system? This would be an interesting and striking innateness type of claim, which should be assessed carefully. We saw that by specializing on a subset of cases (neglecting homophony), current word learning accounts missed to address that discussion (and actually ended up banning homophony from the system entirely).</p>
<p>Let us illustrate with another example showing that the assumptions behind word learning algorithms should be scrutinized with care when they focus their attention on a subclass of phenomena. Xu and Tenenbaum (2007) [<xref ref-type="bibr" rid="pone.0162176.ref012">12</xref>] propose a rational use of co-occurrences of words and objects to learn the meaning of words. But this crucially applies to content words only. Function words, however, occur in all sorts of contexts and may co-occur with all possible objects, in principle. Hence, without further assumptions, the model predicts that words like “the” or “and” mean the same as “thing” or “stuff”, which also co-occur with any kind of object. Arguably, learners deploy a different strategy to learn content words and function words (see relatedly [<xref ref-type="bibr" rid="pone.0162176.ref036">36</xref>] for the use of a different strategy for learning numerical concepts). But how does the learning system separate hypothesis spaces and learning algorithms for content words and function words? Ideally, one would propose a device, say a sorting algorithm, which capitalizes on differences between these words to orient function words and content words to the right learning algorithms for principled reasons (see, e.g., [<xref ref-type="bibr" rid="pone.0162176.ref037">37</xref>] for <italic>empirical</italic> facts about the relative frequency of function and content words that could be the input of such a sorting mechanism). But until such a mechanism is exhibited, the separation of word learning into sub-algorithms which selectively apply to different classes of words comes at the cost of presupposing that learners innately expect the distinction between these classes (i.e. that languages contain both function words and content words, or that languages contain homophones and non-homophones to go back to our central case).</p>
<p>In sum, current word learning accounts break the learning problem into manageable pieces of the puzzle, studying object labels, ambiguous words, functions words or numerical concepts separately. A reconciliation of these pieces into a single solution may be technically easy; one could say that the system “expects” these differences. But it has rich consequences because in the absence of a more complete picture, it amounts to postulating that subtle and quite specific phenomena such as the distinction between function words and content words or the existence of homophony have an innate basis.</p>
</sec>
<sec id="sec026">
<title>Homophony and Polysemy</title>
<p>There is a rich literature distinguishing homophony from polysemy. In short, polysemy is taken to be a form of motivated homophony, by which a word has two related meanings, with possibly systematic variation (e.g., [<xref ref-type="bibr" rid="pone.0162176.ref038">38</xref>]–[<xref ref-type="bibr" rid="pone.0162176.ref041">41</xref>]). For example, an object may take the same label as the artist who made it (e.g., Picasso / a Picasso) and this process is productive (e.g., “The museum owns a <italic>Mandela</italic>” would lead one to infer that Nelson Mandela was a painter). The current literature assumes that such productivity is caused by the presence of generative lexical or conceptual structures that allow meanings to be generated from the context (e.g., [<xref ref-type="bibr" rid="pone.0162176.ref040">40</xref>], [<xref ref-type="bibr" rid="pone.0162176.ref041">41</xref>]). As a result, the representation of polysemous words, for which different meanings are generated from a single entry, is different from the representation of homophones, for which different meanings are stored separately (e.g., [<xref ref-type="bibr" rid="pone.0162176.ref042">42</xref>]–[<xref ref-type="bibr" rid="pone.0162176.ref044">44</xref>]).</p>
<p>We showed that a word is more likely to yield homophony if it labels different portions of the conceptual space. In its simplest version, this would entail that learners would postulate homophony when facing polysemy, as the meanings of a polysemous word also label distant locations in conceptual space (e.g., Picasso, a human man and its painting, an artifact). Yet we know that this is not the case: Children as young as 4 years of age are able to distinguish polysemy from homophony ([<xref ref-type="bibr" rid="pone.0162176.ref045">45</xref>], [<xref ref-type="bibr" rid="pone.0162176.ref046">46</xref>]). Thus, any account of word learning must be able to explain how learners distinguish between homophony and polysemy, surely, capitalizing on the systematic link that exists between the different meanings of a polysemous word, a systematic link which exists even if the meanings are separated in conceptual space.</p>
<p>Depending on how we assume that polysemy is represented, the challenge may take a different form. For the sake of concreteness, assume that polysemy is underlined by Meaning Shift operations, MS, which transform the meaning of a word in systematic ways, (e.g., in the Picasso/Mandela example above, the relevant MS could be akin to “a piece of art made by…”) (see however [<xref ref-type="bibr" rid="pone.0162176.ref047">47</xref>], [<xref ref-type="bibr" rid="pone.0162176.ref048">48</xref>] for the hypothesis that the interpretation of polysemic words depends on pragmatic, gricean reasoning rather than on rule-like processes that generate an extended meaning from an existing one). Hence, polysemy does not involve a different <italic>word</italic> learning process, but the discovery of yet another layer of representation, that of the MS operations, which applies to words. Hence, <italic>Picasso</italic> really is a simple word, the name of an artist, but it can be interpreted as MS(<italic>Picasso</italic>). In practice then, if a word seems to apply to a non-convex set of examples, a learner would have three choices: (i) opt out from ambiguity altogether and postulate a broad meaning for the word, (ii) postulate ambiguity: assume that the word form corresponds to two distinct words, (iii) postulate polysemy: choose one of the meaning as primitive and postulate an MS operation to account for the rest of the examples. We have shown conditions under which learners prefer (ii) over (i), the conditions under which learners prefer (iii) remain to be studied.</p>
</sec>
<sec id="sec027">
<title>Early language acquisition</title>
<p>Through the study of homophones, our studies uncover several factors that play an important role in revealing the existing constraints on how words associate with concepts in general. An important open question is whether these factors influence word learning during the earliest stages of word acquisition (see Dautriche, Chemla &amp; Christophe, <italic>in press</italic>, for initial investigations with 4 year olds). While studying adults may inform us about the general strategies involved in word learning [<xref ref-type="bibr" rid="pone.0162176.ref049">49</xref>], children have different cognitive resources and biases and may consequently use different strategies. We discuss four factors that could lead to the emergence of homophony in children:</p>
<list list-type="order">
<list-item><p><italic>Concept convexity</italic>. Adults refrain from associating a label to a broad concept when positive evidence is missing for a large gap within the concept. The observation that a label applies to a discontinuous extension triggers the formation of novel word representations that are compatible with concept convexity. Do children also expect words to refer to coherent and convex concepts and, if so, what representation do they adopt when the convexity constraint over word forms is not met? [<xref ref-type="bibr" rid="pone.0162176.ref050">50</xref>] offer a relevant study in which they presented 10-month-old infants with exemplars of a word forming a gap in conceptual space: the presence of a similar label was enough for children to extend the label to all intervening items in that gap. Yet, they only tested rather small gaps, which may very well be before the breaking point of the convexity constraint over word forms.</p></list-item>
<list-item><p><italic>Sampling effect</italic>. [<xref ref-type="bibr" rid="pone.0162176.ref012">12</xref>] document a “size principle” according to which the sharpness of a concept is a function of the number of learning exemplars, for both children and adults. We showed an effect of the <italic>distribution</italic> of the learning exemplars in conceptual space: observing exemplars clustered at two distant positions in the hypothesis space boosted the likelihood that the exemplars were sampled from two independent categories. Children are sensitive to the size principle; they may also show sensitivity to such a “distribution principle”, a possibility that we explored elsewhere [<xref ref-type="bibr" rid="pone.0162176.ref051">51</xref>].</p></list-item>
<list-item><p><italic>The structure of the semantic lexicon</italic>. When confronted with a new word, adults consider the existence of <italic>other</italic> (potentially unknown) words. Specifically, they generalize a word A less to a new object if this new object comes in the vicinity of a concept labeled by a word B. This demonstrates that learners have expectations about the structure of the semantic lexicon as a whole and priors about how words may share the conceptual space. This new kind of evidence against individual word-by-word learning is coherent with simpler, so-called “mutual exclusivity effects” [<xref ref-type="bibr" rid="pone.0162176.ref015">15</xref>], according to which a new word should not occupy the same conceptual space as a known word. Interestingly, this effect has to be modulated by other factors, since some words surely overlap in conceptual space (e.g., compare <italic>cat</italic> and <italic>animal</italic>). To our knowledge, priors over the whole lexicon are missing from current word learning computational models—and their implementation raises immediate challenges.</p></list-item>
<list-item><p><italic>Linguistic factors</italic>: Adults’ generalization was modulated by the linguistic construction in which words were presented. While we used linguistic constructions as a linguistic test for homophony, these constructions may also be used to discover homophony (noting that homophones never appear in plural constructions but may appear in some more appropriate constructions such as the <italic>also</italic> construction we documented). Whether children are able to pick up on this is an empirical question, both because they may not be sensitive to these linguistic factors (effectively this would otherwise be a case of linguistic bootstrapping of homophony) or because the relevant facts may be too sparse in their input, e.g., if homophones cover distant concepts, it is unlikely that these two concepts will be mentioned within the same learning situation.</p></list-item>
</list>
</sec>
<sec id="sec028">
<title>Summary</title>
<p>In this work, we showed that a word is more likely to yield homophony if: (a) it is learnt from exemplars leaving an important gap between them (in conceptual space), (b) this gap in conceptual space is occupied by other words. We submit that encountering novel words in such situations may trigger forms of word representations which comply with concept convexity. More generally, we argue that incorporating homophony and other challenging word learning phenomena into current word learning accounts, will provide a better understanding of learners’ implicit knowledge and assumptions about how word forms map onto meanings.</p>
</sec>
</sec>
<sec id="sec029">
<title>Supporting Information</title>
<supplementary-material id="pone.0162176.s001" mimetype="application/vnd.openxmlformats-officedocument.wordprocessingml.document" position="float" xlink:href="info:doi/10.1371/journal.pone.0162176.s001" xlink:type="simple">
<label>S1 Supplemental Material</label>
<caption>
<title/>
<p>(DOCX)</p>
</caption>
</supplementary-material>
</sec>
</body>
<back>
<ack>
<p>We thank Anne Christophe, Paul Egré, Jean-Rémy Hochmann, Alexander Martin, Philippe Schlenker, Benjamin Spector, Brent Strickland.</p>
</ack>
<ref-list>
<title>References</title>
<ref id="pone.0162176.ref001"><label>1</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Frege</surname> <given-names>G.</given-names></name>, “<article-title>Ueber Begriff und Gegenstand</article-title>,” <source><italic>Vierteljahr</italic>. <italic>Fuer Wiss</italic>. <italic>Philos.</italic></source>, <volume>vol. 16</volume>, pp. <fpage>192</fpage>–<lpage>205</lpage>, <year>1892</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref002"><label>2</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Bloom</surname> <given-names>P.</given-names></name>, “<article-title>Précis of How children learn the meanings of words</article-title>,” <source><italic>Behav</italic>. <italic>Brain Sci.</italic></source>, <volume>vol. 24</volume>, <issue>no. 06</issue>, pp. <fpage>1095</fpage>–<lpage>1103</lpage>, <year>2001</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref003"><label>3</label><mixed-citation publication-type="book" xlink:type="simple"><name name-style="western"><surname>Goodman</surname> <given-names>N.</given-names></name>, <source><italic>Fact</italic>, <italic>fiction</italic>, <italic>and forecast</italic></source>. <publisher-name>Harvard University Press</publisher-name>, <year>1955</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref004"><label>4</label><mixed-citation publication-type="book" xlink:type="simple"><name name-style="western"><surname>Keil</surname> <given-names>F. C.</given-names></name>, <source><italic>Concepts</italic>, <italic>kinds</italic>, <italic>and cognitive development</italic></source>. <publisher-name>MIT Press</publisher-name>, <year>1989</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref005"><label>5</label><mixed-citation publication-type="book" xlink:type="simple"><name name-style="western"><surname>Markman</surname> <given-names>E. M.</given-names></name>, <source><italic>Categorization and naming in children</italic>: <italic>Problems of induction</italic></source>. <publisher-name>Mit Press</publisher-name>, <year>1989</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref006"><label>6</label><mixed-citation publication-type="book" xlink:type="simple"><name name-style="western"><surname>Mitchell</surname> <given-names>T. M.</given-names></name>, <source><italic>The need for biases in learning generalizations</italic></source>. <publisher-name>Department of Computer Science, Laboratory for Computer Science Research, Rutgers Univ.</publisher-name>, <year>1980</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref007"><label>7</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Markman</surname> <given-names>E. M.</given-names></name> and <name name-style="western"><surname>Hutchinson</surname> <given-names>J. E.</given-names></name>, “<article-title>Children’s sensitivity to constraints on word meaning: Taxonomic versus thematic relations</article-title>,” <source><italic>Cognit</italic>. <italic>Psychol.</italic></source>, <volume>vol. 16</volume>, <issue>no. 1</issue>, pp. <fpage>1</fpage>–<lpage>27</lpage>, <year>1984</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref008"><label>8</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Waxman</surname> <given-names>S.</given-names></name> and <name name-style="western"><surname>Gelman</surname> <given-names>R.</given-names></name>, “<article-title>Preschoolers’ use of superordinate relations in classification and language</article-title>,” <source><italic>Cogn</italic>. <italic>Dev.</italic></source>, <volume>vol. 1</volume>, <issue>no. 2</issue>, pp. <fpage>139</fpage>–<lpage>156</lpage>, <year>1986</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref009"><label>9</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Landau</surname> <given-names>B.</given-names></name>, <name name-style="western"><surname>Smith</surname> <given-names>L. B.</given-names></name>, and <name name-style="western"><surname>Jones</surname> <given-names>S. S.</given-names></name>, “<article-title>The importance of shape in early lexical learning</article-title>,” <source><italic>Cogn</italic>. <italic>Dev.</italic></source>, <volume>vol. 3</volume>, <issue>no. 3</issue>, pp. <fpage>299</fpage>–<lpage>321</lpage>, <year>1988</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref010"><label>10</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Gardenfors</surname> <given-names>P.</given-names></name>, “<article-title>Conceptual spaces as a framework for knowledge representation</article-title>,” <source><italic>Mind Matter</italic></source>, <volume>vol. 2</volume>, <issue>no. 2</issue>, pp. <fpage>9</fpage>–<lpage>27</lpage>, <year>2004</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref011"><label>11</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Murphy</surname> <given-names>G. L.</given-names></name> and <name name-style="western"><surname>Medin</surname> <given-names>D. L.</given-names></name>, “<article-title>The role of theories in conceptual coherence</article-title>.,” <source><italic>Psychol</italic>. <italic>Rev.</italic></source>, <volume>vol. 92</volume>, <issue>no. 3</issue>, p. <fpage>289</fpage>, <year>1985</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref012"><label>12</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Xu</surname> <given-names>F.</given-names></name> and <name name-style="western"><surname>Tenenbaum</surname> <given-names>J. B.</given-names></name>, “<article-title>Word learning as Bayesian inference</article-title>.,” <source><italic>Psychol</italic>. <italic>Rev.</italic></source>, <volume>vol. 114</volume>, <issue>no. 2</issue>, pp. <fpage>245</fpage>–<lpage>272</lpage>, <year>2007</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref013"><label>13</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Siskind</surname> <given-names>J. M.</given-names></name>, “<article-title>A computational study of cross-situational techniques for learning word-to-meaning mappings</article-title>,” <source><italic>Cognition</italic></source>, <volume>vol. 61</volume>, <issue>no. 1</issue>, pp. <fpage>39</fpage>–<lpage>91</lpage>, <year>1996</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref014"><label>14</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Regier</surname> <given-names>T.</given-names></name>, “<article-title>The emergence of words: Attentional learning in form and meaning</article-title>,” <source><italic>Cogn</italic>. <italic>Sci.</italic></source>, <volume>vol. 29</volume>, <issue>no. 6</issue>, pp. <fpage>819</fpage>–<lpage>865</lpage>, <year>2005</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref015"><label>15</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Markman</surname> <given-names>E. M.</given-names></name> and <name name-style="western"><surname>Wachtel</surname> <given-names>G. F.</given-names></name>, “<article-title>Children’s use of mutual exclusivity to constrain the meanings of words</article-title>,” <source><italic>Cognit</italic>. <italic>Psychol.</italic></source>, <volume>vol. 20</volume>, <issue>no. 2</issue>, pp. <fpage>121</fpage>–<lpage>157</lpage>, <year>1988</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref016"><label>16</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Jaeger</surname> <given-names>T. F.</given-names></name>, “<article-title>Categorical data analysis: Away from ANOVAs (transformation or not) and towards logit mixed models</article-title>,” <source><italic>J</italic>. <italic>Mem</italic>. <italic>Lang.</italic></source>, <volume>vol. 59</volume>, <issue>no. 4</issue>, pp. <fpage>434</fpage>–<lpage>446</lpage>, <month>Nov.</month> <year>2008</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref017"><label>17</label><mixed-citation publication-type="other" xlink:type="simple">Bates D. and Sarkar D., Ime4 library. Accessed, 2004.</mixed-citation></ref>
<ref id="pone.0162176.ref018"><label>18</label><mixed-citation publication-type="book" xlink:type="simple"><name name-style="western"><surname>Cruse</surname> <given-names>D. A.</given-names></name>, <source><italic>Lexical semantics</italic></source>. <publisher-name>Cambridge University Press</publisher-name>, <year>1986</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref019"><label>19</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Zwicky</surname> <given-names>A.</given-names></name> and <name name-style="western"><surname>Sadock</surname> <given-names>J.</given-names></name>, “<article-title>Ambiguity tests and how to fail them</article-title>,” <source><italic>Syntax Semant.</italic></source>, <volume>vol. 4</volume>, <issue>no. 1</issue>, pp. <fpage>1</fpage>–<lpage>36</lpage>, <year>1975</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref020"><label>20</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Geeraerts</surname> <given-names>D.</given-names></name>, “<article-title>Vagueness’s puzzles, polysemy’s vagaries</article-title>,” <source><italic>Cogn</italic>. <italic>Linguist</italic>. <italic>Incl</italic>. <italic>Cogn</italic>. <italic>Linguist</italic>. <italic>Bibliogr.</italic></source>, <volume>vol. 4</volume>, <issue>no. 3</issue>, pp. <fpage>223</fpage>–<lpage>272</lpage>, <year>1993</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref021"><label>21</label><mixed-citation publication-type="other" xlink:type="simple">Goldstone R. L., “Arguments for the Insufficiency of Similarity for Grounding Categorization,” 1994.</mixed-citation></ref>
<ref id="pone.0162176.ref022"><label>22</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Tversky</surname> <given-names>A.</given-names></name>, “<article-title>Similarity features</article-title>,” <source><italic>Psychol</italic>. <italic>Rev.</italic></source>, <volume>vol. 84</volume>, pp. <fpage>327</fpage>–<lpage>352</lpage>, <year>1977</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref023"><label>23</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Poulin-Dubois</surname> <given-names>D.</given-names></name>, <name name-style="western"><surname>Lepage</surname> <given-names>A.</given-names></name>, and <name name-style="western"><surname>Ferland</surname> <given-names>D.</given-names></name>, “<article-title>Infants’ concept of animacy</article-title>,” <source><italic>Cogn</italic>. <italic>Dev.</italic></source>, <volume>vol. 11</volume>, <issue>no. 1</issue>, pp. <fpage>19</fpage>–<lpage>36</lpage>, <year>1996</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref024"><label>24</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Barsalou</surname> <given-names>L. W.</given-names></name>, “<article-title>Ad hoc categories</article-title>,” <source><italic>Mem</italic>. <italic>Cognit.</italic></source>, <volume>vol. 11</volume>, <issue>no. 3</issue>, pp. <fpage>211</fpage>–<lpage>227</lpage>, <year>1983</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref025"><label>25</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Keil</surname> <given-names>F. C.</given-names></name>, “<article-title>Constraints on knowledge and cognitive development</article-title>.,” <source><italic>Psychol</italic>. <italic>Rev.</italic></source>, <volume>vol. 88</volume>, <issue>no. 3</issue>, pp. <fpage>197</fpage>–<lpage>227</lpage>, <year>1981</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref026"><label>26</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Osherson</surname> <given-names>D. N.</given-names></name>, “<article-title>Three conditions on conceptual naturalness</article-title>,” <source><italic>Cognition</italic></source>, <volume>vol. 6</volume>, <issue>no. 4</issue>, pp. <fpage>263</fpage>–<lpage>289</lpage>, <year>1978</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref027"><label>27</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Gelman</surname> <given-names>S. A.</given-names></name> and <name name-style="western"><surname>Coley</surname> <given-names>J. D.</given-names></name>, “<article-title>The importance of knowing a dodo is a bird: Categories and inferences in 2-year-old children</article-title>.,” <source><italic>Dev</italic>. <italic>Psychol.</italic></source>, <volume>vol. 26</volume>, <issue>no. 5</issue>, p. <fpage>796</fpage>, <year>1990</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref028"><label>28</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Gelman</surname> <given-names>S. A.</given-names></name> and <name name-style="western"><surname>Markman</surname> <given-names>E. M.</given-names></name>, “<article-title>Young children’s inductions from natural kinds: The role of categories and appearances</article-title>,” <source><italic>Child Dev.</italic></source>, pp. <fpage>1532</fpage>–<lpage>1541</lpage>, <year>1987</year>. <object-id pub-id-type="pmid">3691200</object-id></mixed-citation></ref>
<ref id="pone.0162176.ref029"><label>29</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Graham</surname> <given-names>S. A.</given-names></name>, <name name-style="western"><surname>Kilbreath</surname> <given-names>C. S.</given-names></name>, and <name name-style="western"><surname>Welder</surname> <given-names>A. N.</given-names></name>, “<article-title>Thirteen-Month-Olds Rely on Shared Labels and Shape Similarity for Inductive Inferences</article-title>,” <source><italic>Child Dev.</italic></source>, <volume>vol. 75</volume>, <issue>no. 2</issue>, pp. <fpage>409</fpage>–<lpage>427</lpage>, <year>2004</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref030"><label>30</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Medin</surname> <given-names>D. L.</given-names></name> and <name name-style="western"><surname>Schaffer</surname> <given-names>M. M.</given-names></name>, “<article-title>Context theory of classification learning</article-title>.,” <source><italic>Psychol</italic>. <italic>Rev.</italic></source>, <volume>vol. 85</volume>, <issue>no. 3</issue>, p. <fpage>207</fpage>, <year>1978</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref031"><label>31</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Nosofsky</surname> <given-names>R. M.</given-names></name>, “<article-title>Attention, similarity, and the identification–categorization relationship</article-title>.,” <source><italic>J</italic>. <italic>Exp</italic>. <italic>Psychol</italic>. <italic>Gen.</italic></source>, <volume>vol. 115</volume>, <issue>no. 1</issue>, p. <fpage>39</fpage>, <year>1986</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref032"><label>32</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Shepard</surname> <given-names>R. N.</given-names></name>, “<article-title>Attention and the metric structure of the stimulus space</article-title>,” <source><italic>J</italic>. <italic>Math</italic>. <italic>Psychol.</italic></source>, <volume>vol. 1</volume>, <issue>no. 1</issue>, pp. <fpage>54</fpage>–<lpage>87</lpage>, <year>1964</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref033"><label>33</label><mixed-citation publication-type="book" xlink:type="simple"><name name-style="western"><surname>Smith</surname> <given-names>E. E.</given-names></name> and <name name-style="western"><surname>Medin</surname> <given-names>D. L.</given-names></name>, <source><italic>Categories and concepts</italic></source>. <publisher-name>Harvard University Press</publisher-name> <publisher-loc>Cambridge, MA</publisher-loc>, <year>1981</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref034"><label>34</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Graham</surname> <given-names>S. A.</given-names></name> and <name name-style="western"><surname>Poulin-Dubois</surname> <given-names>D.</given-names></name>, “<article-title>Infants’ reliance on shape to generalize novel labels to animate and inanimate objects</article-title>,” <source><italic>J</italic>. <italic>Child Lang.</italic></source>, <volume>vol. 26</volume>, <issue>no. 02</issue>, pp. <fpage>295</fpage>–<lpage>320</lpage>, <year>1999</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref035"><label>35</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Atran</surname> <given-names>S.</given-names></name>, “<article-title>Folk biology and the anthropology of science: Cognitive universals and cultural particulars</article-title>,” <source><italic>Behav</italic>. <italic>Brain Sci.</italic></source>, <volume>vol. 21</volume>, <issue>no. 04</issue>, pp. <fpage>547</fpage>–<lpage>569</lpage>, <year>1998</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref036"><label>36</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Piantadosi</surname> <given-names>S. T.</given-names></name>, <name name-style="western"><surname>Tenenbaum</surname> <given-names>J. B.</given-names></name>, and <name name-style="western"><surname>Goodman</surname> <given-names>N. D.</given-names></name>, “<article-title>Bootstrapping in a language of thought: A formal model of numerical concept learning</article-title>,” <source><italic>Cognition</italic></source>, <volume>vol. 123</volume>, <issue>no. 2</issue>, pp. <fpage>199</fpage>–<lpage>217</lpage>, <month>May</month> <year>2012</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref037"><label>37</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Hochmann</surname> <given-names>J.-R.</given-names></name>, “<article-title>Word frequency, function words and the second gavagai problem</article-title>,” <source><italic>Cognition</italic></source>, <volume>vol. 128</volume>, <issue>no. 1</issue>, pp. <fpage>13</fpage>–<lpage>25</lpage>, <month>Jul.</month> <year>2013</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref038"><label>38</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Apresjan</surname> <given-names>J. D.</given-names></name>, “<article-title>Regular polysemy</article-title>,” <source><italic>Linguistics</italic></source>, <volume>vol. 12</volume>, <issue>no. 142</issue>, pp. <fpage>5</fpage>–<lpage>32</lpage>, <year>1974</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref039"><label>39</label><mixed-citation publication-type="book" xlink:type="simple"><name name-style="western"><surname>Bréal</surname> <given-names>M.</given-names></name>, <source><italic>Essai de sémantique</italic>:<italic>(science des significations)</italic></source>. <publisher-name>Hachette</publisher-name>, <year>1904</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref040"><label>40</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Copestake</surname> <given-names>A.</given-names></name> and <name name-style="western"><surname>Briscoe</surname> <given-names>T.</given-names></name>, “<article-title>Semi-productive polysemy and sense extension</article-title>,” <source><italic>J</italic>. <italic>Semant.</italic></source>, <volume>vol. 12</volume>, <issue>no. 1</issue>, pp. <fpage>15</fpage>–<lpage>67</lpage>, <year>1995</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref041"><label>41</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Pustejovsky</surname> <given-names>J.</given-names></name>, “<article-title>The generative lexicon</article-title>,” <source><italic>Comput</italic>. <italic>Linguist.</italic></source>, <volume>vol. 17</volume>, <issue>no. 4</issue>, pp. <fpage>409</fpage>–<lpage>441</lpage>, <year>1991</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref042"><label>42</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Caramazza</surname> <given-names>A.</given-names></name> and <name name-style="western"><surname>Grober</surname> <given-names>E.</given-names></name>, “<article-title>Polysemy and the structure of the subjective lexicon</article-title>,” <source><italic>Georget</italic>. <italic>Univ</italic>. <italic>Roundtable Lang</italic>. <italic>Linguist</italic>. <italic>Semant</italic>. <italic>Theory Appl.</italic></source>, pp. <fpage>181</fpage>–<lpage>206</lpage>, <year>1976</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref043"><label>43</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Rabagliati</surname> <given-names>H.</given-names></name> and <name name-style="western"><surname>Snedeker</surname> <given-names>J.</given-names></name>, “<article-title>The Truth About Chickens and Bats: Ambiguity Avoidance Distinguishes Types of Polysemy</article-title>,” <source><italic>Psychol</italic>. <italic>Sci.</italic></source>, <month>May</month> <year>2013</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref044"><label>44</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Seidenberg</surname> <given-names>M. S.</given-names></name>, <name name-style="western"><surname>Tanenhaus</surname> <given-names>M. K.</given-names></name>, <name name-style="western"><surname>Leiman</surname> <given-names>J. M.</given-names></name>, and <name name-style="western"><surname>Bienkowski</surname> <given-names>M.</given-names></name>, “<article-title>Automatic access of the meanings of ambiguous words in context: Some limitations of knowledge-based processing</article-title>,” <source><italic>Cognit</italic>. <italic>Psychol.</italic></source>, <volume>vol. 14</volume>, <issue>no. 4</issue>, pp. <fpage>489</fpage>–<lpage>537</lpage>, <year>1982</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref045"><label>45</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Srinivasan</surname> <given-names>M.</given-names></name> and <name name-style="western"><surname>Snedeker</surname> <given-names>J.</given-names></name>, “<article-title>Judging a book by its cover and its contents: The representation of polysemous and homophonous meanings in four-year-old children</article-title>,” <source><italic>Cognit</italic>. <italic>Psychol.</italic></source>, <volume>vol. 62</volume>, <issue>no. 4</issue>, pp. <fpage>245</fpage>–<lpage>272</lpage>, <month>Jun.</month> <year>2011</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref046"><label>46</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Srinivasan</surname> <given-names>M.</given-names></name> and <name name-style="western"><surname>Snedeker</surname> <given-names>J.</given-names></name>, “<article-title>Polysemy and the Taxonomic Constraint: Children’s Representation of Words that Label Multiple Kinds</article-title>,” <source><italic>Lang</italic>. <italic>Learn</italic>. <italic>Dev.</italic></source>, <volume>vol. 10</volume>, <issue>no. 2</issue>, pp. <fpage>97</fpage>–<lpage>128</lpage>, <year>2014</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref047"><label>47</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Nunberg</surname> <given-names>G.</given-names></name>, “<article-title>The non-uniqueness of semantic solutions: Polysemy</article-title>,” <source><italic>Linguist</italic>. <italic>Philos.</italic></source>, <volume>vol. 3</volume>, <issue>no. 2</issue>, pp. <fpage>143</fpage>–<lpage>184</lpage>, <year>1979</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref048"><label>48</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Nunberg</surname> <given-names>G.</given-names></name>, “<article-title>Transfers of meaning</article-title>,” <source><italic>J</italic>. <italic>Semant.</italic></source>, <volume>vol. 12</volume>, <issue>no. 2</issue>, pp. <fpage>109</fpage>–<lpage>132</lpage>, <year>1995</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref049"><label>49</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Markson</surname> <given-names>L.</given-names></name> and <name name-style="western"><surname>Bloom</surname> <given-names>P.</given-names></name>, “<article-title>Evidence against a dedicated system for word learning in children</article-title>,” <source><italic>Nature</italic></source>, <volume>vol. 385</volume>, <issue>no. 6619</issue>, pp. <fpage>813</fpage>–<lpage>815</lpage>, <year>1997</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref050"><label>50</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Plunkett</surname> <given-names>K.</given-names></name>, <name name-style="western"><surname>Hu</surname> <given-names>J.-F.</given-names></name>, and <name name-style="western"><surname>Cohen</surname> <given-names>L. B.</given-names></name>, “<article-title>Labels can override perceptual categories in early infancy</article-title>,” <source><italic>Cognition</italic></source>, <volume>vol. 106</volume>, <issue>no. 2</issue>, pp. <fpage>665</fpage>–<lpage>681</lpage>, <month>Feb.</month> <year>2008</year>.</mixed-citation></ref>
<ref id="pone.0162176.ref051"><label>51</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Dautriche</surname> <given-names>I.</given-names></name>, <name name-style="western"><surname>Chemla</surname> <given-names>E.</given-names></name>, and <name name-style="western"><surname>Christophe</surname> <given-names>A.</given-names></name>, “<article-title>Word Learning: Homophony and the Distribution of Learning Exemplars</article-title>,” <source><italic>Lang</italic>. <italic>Learn</italic>. <italic>Dev.</italic></source>, <volume>vol. 12</volume>, <issue>no. 3</issue>, pp. <fpage>231</fpage>–<lpage>251</lpage>, <year>2016</year>.</mixed-citation></ref>
</ref-list>
</back>
</article>