<?xml version="1.0" encoding="utf-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD JATS (Z39.96) Journal Publishing DTD v1.1d3 20150301//EN" "http://jats.nlm.nih.gov/publishing/1.1d3/JATS-journalpublishing1.dtd">
<article article-type="research-article" dtd-version="1.1d3" xml:lang="en" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink">
<front>
<journal-meta>
<journal-id journal-id-type="nlm-ta">PLoS ONE</journal-id>
<journal-id journal-id-type="publisher-id">plos</journal-id>
<journal-id journal-id-type="pmc">plosone</journal-id>
<journal-title-group>
<journal-title>PLOS ONE</journal-title>
</journal-title-group>
<issn pub-type="epub">1932-6203</issn>
<publisher>
<publisher-name>Public Library of Science</publisher-name>
<publisher-loc>San Francisco, CA USA</publisher-loc>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="publisher-id">PONE-D-16-33203</article-id>
<article-id pub-id-type="doi">10.1371/journal.pone.0172959</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Research Article</subject>
</subj-group>
<subj-group subj-group-type="Discipline-v3"><subject>Biology and life sciences</subject><subj-group><subject>Microbiology</subject><subj-group><subject>Medical microbiology</subject><subj-group><subject>Microbial pathogens</subject><subj-group><subject>Viral pathogens</subject><subj-group><subject>Immunodeficiency viruses</subject><subj-group><subject>HIV</subject></subj-group></subj-group></subj-group></subj-group></subj-group></subj-group></subj-group><subj-group subj-group-type="Discipline-v3"><subject>Medicine and health sciences</subject><subj-group><subject>Pathology and laboratory medicine</subject><subj-group><subject>Pathogens</subject><subj-group><subject>Microbial pathogens</subject><subj-group><subject>Viral pathogens</subject><subj-group><subject>Immunodeficiency viruses</subject><subj-group><subject>HIV</subject></subj-group></subj-group></subj-group></subj-group></subj-group></subj-group></subj-group><subj-group subj-group-type="Discipline-v3"><subject>Biology and life sciences</subject><subj-group><subject>Organisms</subject><subj-group><subject>Viruses</subject><subj-group><subject>Viral pathogens</subject><subj-group><subject>Immunodeficiency viruses</subject><subj-group><subject>HIV</subject></subj-group></subj-group></subj-group></subj-group></subj-group></subj-group><subj-group subj-group-type="Discipline-v3"><subject>Biology and life sciences</subject><subj-group><subject>Organisms</subject><subj-group><subject>Viruses</subject><subj-group><subject>Immunodeficiency viruses</subject><subj-group><subject>HIV</subject></subj-group></subj-group></subj-group></subj-group></subj-group><subj-group subj-group-type="Discipline-v3"><subject>Biology and life sciences</subject><subj-group><subject>Organisms</subject><subj-group><subject>Viruses</subject><subj-group><subject>RNA viruses</subject><subj-group><subject>Retroviruses</subject><subj-group><subject>Lentivirus</subject><subj-group><subject>HIV</subject></subj-group></subj-group></subj-group></subj-group></subj-group></subj-group></subj-group><subj-group subj-group-type="Discipline-v3"><subject>Biology and life sciences</subject><subj-group><subject>Microbiology</subject><subj-group><subject>Medical microbiology</subject><subj-group><subject>Microbial pathogens</subject><subj-group><subject>Viral pathogens</subject><subj-group><subject>Retroviruses</subject><subj-group><subject>Lentivirus</subject><subj-group><subject>HIV</subject></subj-group></subj-group></subj-group></subj-group></subj-group></subj-group></subj-group></subj-group><subj-group subj-group-type="Discipline-v3"><subject>Medicine and health sciences</subject><subj-group><subject>Pathology and laboratory medicine</subject><subj-group><subject>Pathogens</subject><subj-group><subject>Microbial pathogens</subject><subj-group><subject>Viral pathogens</subject><subj-group><subject>Retroviruses</subject><subj-group><subject>Lentivirus</subject><subj-group><subject>HIV</subject></subj-group></subj-group></subj-group></subj-group></subj-group></subj-group></subj-group></subj-group><subj-group subj-group-type="Discipline-v3"><subject>Biology and life sciences</subject><subj-group><subject>Organisms</subject><subj-group><subject>Viruses</subject><subj-group><subject>Viral pathogens</subject><subj-group><subject>Retroviruses</subject><subj-group><subject>Lentivirus</subject><subj-group><subject>HIV</subject></subj-group></subj-group></subj-group></subj-group></subj-group></subj-group></subj-group><subj-group subj-group-type="Discipline-v3"><subject>Medicine and health sciences</subject><subj-group><subject>Epidemiology</subject><subj-group><subject>HIV epidemiology</subject></subj-group></subj-group></subj-group><subj-group subj-group-type="Discipline-v3"><subject>Medicine and health sciences</subject><subj-group><subject>Infectious diseases</subject><subj-group><subject>Viral diseases</subject><subj-group><subject>HIV infections</subject></subj-group></subj-group></subj-group></subj-group><subj-group subj-group-type="Discipline-v3"><subject>People and places</subject><subj-group><subject>Geographical locations</subject><subj-group><subject>Africa</subject><subj-group><subject>Mozambique</subject></subj-group></subj-group></subj-group></subj-group><subj-group subj-group-type="Discipline-v3"><subject>Medicine and health sciences</subject><subj-group><subject>Public and occupational health</subject><subj-group><subject>Preventive medicine</subject><subj-group><subject>HIV prevention</subject></subj-group></subj-group></subj-group></subj-group><subj-group subj-group-type="Discipline-v3"><subject>Physical sciences</subject><subj-group><subject>Mathematics</subject><subj-group><subject>Probability theory</subject><subj-group><subject>Random variables</subject><subj-group><subject>Covariance</subject></subj-group></subj-group></subj-group></subj-group></subj-group><subj-group subj-group-type="Discipline-v3"><subject>Physical sciences</subject><subj-group><subject>Mathematics</subject><subj-group><subject>Statistics (mathematics)</subject><subj-group><subject>Statistical models</subject></subj-group></subj-group></subj-group></subj-group><subj-group subj-group-type="Discipline-v3"><subject>Medicine and health sciences</subject><subj-group><subject>Infectious diseases</subject><subj-group><subject>Viral diseases</subject><subj-group><subject>AIDS</subject></subj-group></subj-group></subj-group></subj-group></article-categories>
<title-group>
<article-title>A flexible method to model HIV serodiscordance among couples in Mozambique</article-title>
<alt-title alt-title-type="running-head">A flexible method to model HIV serodiscordance among couples in Mozambique</alt-title>
</title-group>
<contrib-group>
<contrib contrib-type="author" corresp="yes" xlink:type="simple">
<name name-style="western">
<surname>Juga</surname> <given-names>Adelino J. C.</given-names></name>
<xref ref-type="aff" rid="aff001"><sup>1</sup></xref>
<xref ref-type="aff" rid="aff002"><sup>2</sup></xref>
<xref ref-type="corresp" rid="cor001">*</xref>
</contrib>
<contrib contrib-type="author" xlink:type="simple">
<name name-style="western">
<surname>Hens</surname> <given-names>Niel</given-names></name>
<xref ref-type="aff" rid="aff002"><sup>2</sup></xref>
<xref ref-type="aff" rid="aff003"><sup>3</sup></xref>
<xref ref-type="aff" rid="aff004"><sup>4</sup></xref>
</contrib>
<contrib contrib-type="author" xlink:type="simple">
<name name-style="western">
<surname>Osman</surname> <given-names>Nafissa</given-names></name>
<xref ref-type="aff" rid="aff005"><sup>5</sup></xref>
<xref ref-type="aff" rid="aff006"><sup>6</sup></xref>
</contrib>
<contrib contrib-type="author" xlink:type="simple">
<name name-style="western">
<surname>Aerts</surname> <given-names>Marc</given-names></name>
<xref ref-type="aff" rid="aff002"><sup>2</sup></xref>
</contrib>
</contrib-group>
<aff id="aff001">
<label>1</label>
<addr-line>Department of Mathematics and Informatics, Faculty of Sciences, Eduardo Mondlane University, Maputo, Mozambique</addr-line>
</aff>
<aff id="aff002">
<label>2</label>
<addr-line>I-BioStat, Hasselt University, Diepenbeek, Belgium</addr-line>
</aff>
<aff id="aff003">
<label>3</label>
<addr-line>Centre for Health Economic Research and Modelling Infectious Diseases, Vaccine and Infectious Disease Institute (VAXINFECTIO), University of Antwerp, Antwerp, Belgium</addr-line>
</aff>
<aff id="aff004">
<label>4</label>
<addr-line>Epidemiology and Social Medicine (ESOC), Faculty of Medicine and Health Sciences, University of Antwerp, Antwerp, Belgium</addr-line>
</aff>
<aff id="aff005">
<label>5</label>
<addr-line>Department of Obstetrics and Gynaecology, Maputo Central Hospital, Maputo, Mozambique</addr-line>
</aff>
<aff id="aff006">
<label>6</label>
<addr-line>Faculty of Medicine, Eduardo Mondlane University, Maputo, Mozambique</addr-line>
</aff>
<contrib-group>
<contrib contrib-type="editor" xlink:type="simple">
<name name-style="western">
<surname>De Socio</surname> <given-names>Giuseppe Vittorio</given-names></name>
<role>Editor</role>
<xref ref-type="aff" rid="edit1"/>
</contrib>
</contrib-group>
<aff id="edit1">
<addr-line>Azienda Ospedaliera Universitaria di Perugia, ITALY</addr-line>
</aff>
<author-notes>
<fn fn-type="conflict" id="coi001">
<p>The authors have declared that no competing interests exist.</p>
</fn>
<fn fn-type="con">
<p>
<list list-type="simple">
<list-item>
<p><bold>Conceptualization:</bold> AJCJ MA.</p>
</list-item>
<list-item>
<p><bold>Data curation:</bold> AJCJ MA NH.</p>
</list-item>
<list-item>
<p><bold>Formal analysis:</bold> AJCJ.</p>
</list-item>
<list-item>
<p><bold>Funding acquisition:</bold> MA NO.</p>
</list-item>
<list-item>
<p><bold>Investigation:</bold> MA NH.</p>
</list-item>
<list-item>
<p><bold>Methodology:</bold> AJCJ MA.</p>
</list-item>
<list-item>
<p><bold>Project administration:</bold> NO.</p>
</list-item>
<list-item>
<p><bold>Resources:</bold> MA NO.</p>
</list-item>
<list-item>
<p><bold>Software:</bold> AJCJ.</p>
</list-item>
<list-item>
<p><bold>Supervision:</bold> MA NH.</p>
</list-item>
<list-item>
<p><bold>Validation:</bold> MA NH.</p>
</list-item>
<list-item>
<p><bold>Visualization:</bold> MA NH.</p>
</list-item>
<list-item>
<p><bold>Writing – original draft:</bold> AJCJ.</p>
</list-item>
<list-item>
<p><bold>Writing – review &amp; editing:</bold> AJCJ MA.</p>
</list-item>
</list>
</p>
</fn>
<corresp id="cor001">* E-mail: <email xlink:type="simple">adelino.juga@gmail.com</email></corresp>
</author-notes>
<pub-date pub-type="collection">
<year>2017</year>
</pub-date>
<pub-date pub-type="epub">
<day>2</day>
<month>3</month>
<year>2017</year>
</pub-date>
<volume>12</volume>
<issue>3</issue>
<elocation-id>e0172959</elocation-id>
<history>
<date date-type="received">
<day>19</day>
<month>8</month>
<year>2016</year>
</date>
<date date-type="accepted">
<day>13</day>
<month>2</month>
<year>2017</year>
</date>
</history>
<permissions>
<license xlink:href="https://creativecommons.org/publicdomain/zero/1.0/" xlink:type="simple">
<license-p>This is an open access article, free of all copyright, and may be freely reproduced, distributed, transmitted, modified, built upon, or otherwise used by anyone for any lawful purpose. The work is made available under the <ext-link ext-link-type="uri" xlink:href="https://creativecommons.org/publicdomain/zero/1.0/" xlink:type="simple">Creative Commons CC0</ext-link> public domain dedication.</license-p>
</license>
</permissions>
<self-uri content-type="pdf" xlink:href="info:doi/10.1371/journal.pone.0172959"/>
<abstract>
<p>Whereas the number of people newly infected by HIV is continuing to decline globally, the epidemic continues to expand in many parts of the world. As the HIV/AIDS epidemic has matured in many countries, it is believed that the proportion of new infections occurring within couples has risen. Across countries, including Mozambique, a sizeable proportion of couples with HIV infection are discordant. A serodiscordant couple is a couple in which one partner has tested positive for HIV and the other has not. To describe the HIV serodiscordance among couples, a variety of association measures can be used. In this paper, we propose the serodiscordance measure <monospace>(SDM)</monospace> as a new alternative measure. Focus is on the specification of flexible marginal and random effects models for multivariate correlated binary data together with a full-likelihood estimation method, to adequately and directly describe the measure of interest. Fitting joint models allows examining the effects of different risk factors and other covariates on the probability to be HIV positive for each member within a couple, and estimating common effects for both probabilities more efficiently, while accounting for the association between their infection status. Moreover, the interpretation of the proposed association parameter SDM is more direct and relevant and effects of covariates can be studied as well. Results show that the HIV prevalence for the province where a couple was located as well as the union number for the woman within a couple are factors associated with HIV serodiscordance. These findings are important for the Mozambican public health policy makers to design national prevention plans, which include policies to stimulate regular HIV testing for couples as well as adolescents and young adults, prior to getting married or living together as a couple.</p>
</abstract>
<funding-group>
<award-group id="award001">
<funding-source>
<institution>Flemish Interuniversity Council (VLIR-UOS)</institution>
</funding-source>
<principal-award-recipient>
<name name-style="western">
<surname>Juga</surname> <given-names>Adelino Jose Chingore</given-names></name>
</principal-award-recipient>
</award-group>
<funding-statement>Flemish Interuniversity Council (VLIR-UOS) in collaboration with Eduardo Mondlane University (UEM) through the DESAFIO Program provided financial support.</funding-statement>
</funding-group>
<counts>
<fig-count count="0"/>
<table-count count="3"/>
<page-count count="14"/>
</counts>
<custom-meta-group>
<custom-meta id="data-availability">
<meta-name>Data Availability</meta-name>
<meta-value>The data of this paper are available in DHS program website <ext-link ext-link-type="uri" xlink:href="http://dhsprogram.com/what-we-do/survey/survey-display-322.cfm" xlink:type="simple">http://dhsprogram.com/what-we-do/survey/survey-display-322.cfm</ext-link>.</meta-value>
</custom-meta>
</custom-meta-group>
</article-meta>
</front>
<body>
<sec id="sec001" sec-type="intro">
<title>Introduction</title>
<p>In recent years, there has been increasing interest in the spread of HIV within stable sexual partnerships [<xref ref-type="bibr" rid="pone.0172959.ref001">1</xref>], since the HIV epidemic has matured in many countries and it is believed that the proportion of new infections occurring within couples has risen [<xref ref-type="bibr" rid="pone.0172959.ref002">2</xref>]. Evidence has shown that across countries, a sizeable proportion of couples with any HIV infection are discordant [<xref ref-type="bibr" rid="pone.0172959.ref002">2</xref>]. A serodiscordant couple is a couple in which one partner has tested positive for HIV and the other has not.</p>
<p>Evidence suggests that women have the greatest risk of contracting HIV within a marital relationship. Fewer attempts have been made to understand a man’s risk. But regions with generalised epidemics, high rates of serodiscordance among heterosexual couples in which the woman is HIV positive suggest that a man’s risk of marital HIV acquisition could also be substantial [<xref ref-type="bibr" rid="pone.0172959.ref003">3</xref>].</p>
<p>In Mozambique, in among 1.6 million people living with HIV/AIDS in 2009, around one in every ten couples were discordant. In 5% of all couples, both members were HIV positive (concordant positive), while in 10% of all couples, one member was HIV positive whilst the other was HIV negative (female or male discordant). The estimates of 2011 showed that the proportion of male discordant couples was similar to the proportion of female discordant couples [<xref ref-type="bibr" rid="pone.0172959.ref001">1</xref>]. By female discordant we refer to a couple where the woman was HIV positive while the man was HIV negative. A couple where the man was HIV positive whilst the woman was HIV negative is referred to as a male discordant couple.</p>
<p>Despite the empirical evidence pointing to their programmatic importance, serodiscordant couples are often overlooked or, at best, only vaguely addressed in many national prevention plans. This omission may stem not only from sensitivity surrounding HIV within couples but also from misperceptions about the extent of serodiscordance and failure to understand that it is possible to prevent transmission within a stable union once one partner has become infected [<xref ref-type="bibr" rid="pone.0172959.ref004">4</xref>].</p>
<p>Statistical methods and models to investigate risk factors associated with HIV serodiscordance among couples can be very helpful. Simple methods are based on the use of descriptive statistics as well as bivariate associations. Extensions of such descriptive analyses have been covered by (standard) logistic regression models. Kaiser <italic>et al</italic>. [<xref ref-type="bibr" rid="pone.0172959.ref005">5</xref>] developed two different models to assess factors associated with couple status, one for HIV discordance as the outcome and a second one with HIV concordance as the binary outcome. In the first model, concordant HIV-infected couples were excluded from the denominator and in the second model concordant uninfected couples were excluded from the denominator. Based on the 2007 Kenya AIDS Indicator Survey, they found that the following factors were independently associated with HIV-discordance: young age in women, increasing number of lifetime sexual partners in women, HSV-2 (herpes simplex virus type 2) infection in either or both partners, and lack of male circumcision. Independent factors for HIV-concordance included HSV-2 infection in both partners and lack of male circumcision. Fishel <italic>et al</italic>. [<xref ref-type="bibr" rid="pone.0172959.ref001">1</xref>] based their findings on an analysis including two logistic regression models. The first model compared concordant positive couples with all discordant couples, regardless of whether the discordant couple is male discordant or female discordant. The second multinomial logistic regression compares concordant positive couples separately to male discordant and female discordant couples. They concluded that HIV discordance of couples in Mozambique does not appear to have a strong association with whether or not the couple is polygynous, the amount of time passed since the couple last had sexual intercourse with each other, or risk factors for non-sexual transmission of HIV. Although there are some differences in the HIV status of couples by characteristics of interest, the differences do not outline a profile that makes discordant couples easy to distinguish from the general population.</p>
<p>Our approach is different in two ways. Instead of using two separate logistic models, we propose to use genuine bivariate models for correlated binary data, allowing one to fit separate logistic regression models in one single analysis while taking into account the correlation between both binary responses (HIV status of man and woman within a couple). Moreover we propose to reparametrize the model such that serodiscordance itself is a model parameter, allowing to investigate the factors affecting the serodiscordance in a direct way. The particular measure of serodiscordance proposed in this paper is inspired by the conditional synchrony measure (<monospace>CSM</monospace>) as introduced by [<xref ref-type="bibr" rid="pone.0172959.ref006">6</xref>] to measure synchrony in neuronal firing.</p>
<p>There exist different families of models for correlated binary data: so-called marginal, population averaged models and conditional models. Full likelihood models such as the bivariate Dale model [<xref ref-type="bibr" rid="pone.0172959.ref007">7</xref>] and alternative models estimated by the method of moments (Generalized Estimating Equations, GEE, [<xref ref-type="bibr" rid="pone.0172959.ref008">8</xref>]) are of the first type. Generalized linear mixed models [<xref ref-type="bibr" rid="pone.0172959.ref009">9</xref>], such as the logistic-normal model, are based on the use of random effects and constitute a second type of models. Which model family or combination of types is to be preferred depends on the design of the study and on the research question(s) of interest.</p>
<p>In this paper, we will propose the combination of a reparametrized version of the bivariate Dale model together with multivariate random effects with different covariance structures to take into account the design of the study. Focus is on the specification and estimation of appropriate covariate models for all model components including the serodiscordance parameter. The models were designed to investigate in particular the relationship/association between the HIV status of woman and man within a couple and the risk factors associated with HIV serodiscordance among couples.</p>
<p>The paper is organized as follows. A first section on materials and methods provides information on the INSIDA survey with some details about the survey design and the variables used in this study. The statistical models are introduced together with a new measure of serodiscordance. The use of survey weights and the model building strategy is briefly discussed. The second section summarizes the results of the application of the proposed methodology on the INSIDA data and the paper ends with conclusions and a final discussion.</p>
</sec>
<sec id="sec002" sec-type="materials|methods">
<title>Materials and methods</title>
<sec id="sec003">
<title>Survey design</title>
<p>We used data from the 2009 National Survey of Prevalence, Risk Behavioural and Information about HIV and AIDS (INSIDA, [<xref ref-type="bibr" rid="pone.0172959.ref010">10</xref>]). This survey is a cross-sectional two-stage survey, carried out by the National Institute of Health in collaboration with the National Bureau of Statistics of Mozambique. It was the first survey designed to collect comprehensive data on the prevalence of HIV infection, knowledge, attitude, behaviour risk factors and access to information on HIV and AIDS in the Mozambican population. Of particular interest in this survey was the HIV status of cohabiting couples. Access to the data can be requested by logging in at <ext-link ext-link-type="uri" xlink:href="http://dhsprogram.com/data/dataset/Mozambique_Standard-AIS_2009.cfm?flag=0" xlink:type="simple">http://dhsprogram.com/data/dataset/Mozambique_Standard-AIS_2009.cfm?flag=0</ext-link> and by submitting a research project proposal.</p>
<p>Stratification and cluster sampling methods were applied to ensure that for each province inference was possible with nearly the same precision. Moreover, two-stage sampling was also used to access individuals within households. Enumeration Areas (EAs), households and individuals were Primary Sampling Units (PSU), Secondary Sampling Units (SSU) and Tertiary Sampling Units (TSU) respectively. Under stratification, 11 provinces were considered as stratums. Moreover, 270 EAs were selected in all provinces from a total of 45000 EAs defined according to the cartography of general census, 2007. Out of selected 270 EAs, 122 were urban and 148 rural. Then a fixed number of households were systematically selected within each EA in the second stage. In this phase, 22 households from urban EAs and 24 from rural EAs were selected [<xref ref-type="bibr" rid="pone.0172959.ref001">1</xref>]. Men and women aged between 15–64 years were eligible to participate in an individual interview and to provide a blood sample for the HIV test. For ethical reasons, it was not asked whether the respondents knew their HIV status [<xref ref-type="bibr" rid="pone.0172959.ref001">1</xref>]. To identify couples, for each woman that was married or lived together with her male partner, an attempt was made to match her with her husband/partner using his household line number. A confirmation was then made by checking the man’s interview information that was named by the woman. If a man was polygamous (married or living with more than one wife/woman) he may appear in the database multiple times, once for each wife [<xref ref-type="bibr" rid="pone.0172959.ref001">1</xref>]. 334 men appeared to be polygamous. Out of 6190 households, 10(0.16%) of them had two couples and only 1 household had 3 couples.</p>
</sec>
<sec id="sec004">
<title>Description of variables</title>
<p><xref ref-type="table" rid="pone.0172959.t001">Table 1</xref> shows a brief description of all variables used in the final model. None of them has missing values. The HIV status (infected yes/no) for each member within a couple constitute the (binary) response variables of interest. The following covariates appear in the final model with fixed effects. The variables union number for man/woman refer to whether the respondent has been married or lived with a woman/man once or more than once. Wealth index refers to the economic status of the couple while condom used by man refers to whether the male respondent of the couple used a condom the last time he had sexual intercourse with the other partner. The HIV prevalence for province was categorized into three categories using cutpoints of 5 and 15% [<xref ref-type="bibr" rid="pone.0172959.ref001">1</xref>]. Finally, heterogeneity across the randomly selected EAs was modelled using random effects.</p>
<table-wrap id="pone.0172959.t001" position="float">
<object-id pub-id-type="doi">10.1371/journal.pone.0172959.t001</object-id>
<label>Table 1</label>
<caption>
<title>INSIDA survey: basic description of variables used in the final model.</title>
</caption>
<alternatives>
<graphic id="pone.0172959.t001g" mimetype="image" position="float" xlink:href="info:doi/10.1371/journal.pone.0172959.t001" xlink:type="simple"/>
<table border="0" frame="box" rules="all">
<colgroup>
<col align="left" valign="middle"/>
<col align="left" valign="middle"/>
<col align="left" valign="middle"/>
</colgroup>
<thead>
<tr>
<th align="left">Variables</th>
<th align="left">Type</th>
<th align="left">Levels</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" colspan="3"><bold>At the individual level</bold></td>
</tr>
<tr>
<td align="left" rowspan="2">HIV status for woman</td>
<td align="left" rowspan="2">Binary</td>
<td align="left">0: HIV Negative (refence category)</td>
</tr>
<tr>
<td align="left">1: HIV Positive</td>
</tr>
<tr>
<td align="left" rowspan="2">HIV status for man</td>
<td align="left" rowspan="2">Binary</td>
<td align="left">0: HIV Negative (reference category)</td>
</tr>
<tr>
<td align="left">1: HIV Positive</td>
</tr>
<tr>
<td align="left" rowspan="2">Union number for woman</td>
<td align="left" rowspan="2">Binary</td>
<td align="left">1: once (reference category)</td>
</tr>
<tr>
<td align="left">2: more than once</td>
</tr>
<tr>
<td align="left" rowspan="2">Union number for man</td>
<td align="left" rowspan="2">Binary</td>
<td align="left">1: once (reference category)</td>
</tr>
<tr>
<td align="left">2: more than once</td>
</tr>
<tr>
<td align="left" colspan="3"><bold>At the level of the couple</bold></td>
</tr>
<tr>
<td align="left" rowspan="2">Condom used by man</td>
<td align="left" rowspan="2">Binary</td>
<td align="left">1: Used(reference category)</td>
</tr>
<tr>
<td align="left">2: Not used</td>
</tr>
<tr>
<td align="left" rowspan="3">Wealth index</td>
<td align="left" rowspan="3">Categorical</td>
<td align="left">1: Poorer</td>
</tr>
<tr>
<td align="left">2: Middle</td>
</tr>
<tr>
<td align="left">3: Richer (reference category)</td>
</tr>
<tr>
<td align="left" colspan="3"><bold>At the level of the province</bold></td>
</tr>
<tr>
<td align="left" rowspan="3">HIV prevalence of province</td>
<td align="left" rowspan="3">Categorical</td>
<td align="left">1: &lt; 5%(reference category)</td>
</tr>
<tr>
<td align="left">2: 5%–15%</td>
</tr>
<tr>
<td align="left">3: &gt; 15%</td>
</tr>
</tbody>
</table>
</alternatives>
</table-wrap>
</sec>
<sec id="sec005">
<title>Joint models for bivariate binary outcomes</title>
<p>In this section the statistical model is introduced and different measures of association are discussed. A new measure of association, the serodiscordance measure, is introduced and embedded in the statistical model. Next, the model is extended with random effects accounting for heterogeneity across enumeration areas. Finally more information is provided about the use of sample weights accounting for the survey design. Finally the model building strategy is briefly described.</p>
<sec id="sec006">
<title>Joint marginal model</title>
<p>Let <italic>y</italic> = (<italic>y</italic><sub>1</sub>, <italic>y</italic><sub>2</sub>) denote the HIV status of a couple, woman and man respectively, and <italic>x</italic> a vector of covariates (including the constant 1 associated with the intercept). Some of the covariates are couple-specific (such as wealth index), whereas others are individual specific (union number). But any variable is considered to have a potential effect on the joint distribution of <italic>y</italic> = (<italic>y</italic><sub>1</sub>, <italic>y</italic><sub>2</sub>). So, e.g. condom use by man could have an effect on the probability that the HIV status of the man is positive, on that of the woman as well and possibly on the association between both statuses (while correcting for all other covariates in the model).</p>
<p>The model arises from the decomposition of the joint probabilities (given a covariate pattern <italic>x</italic>)
<disp-formula id="pone.0172959.e001"><alternatives><graphic id="pone.0172959.e001g" mimetype="image" position="anchor" xlink:href="info:doi/10.1371/journal.pone.0172959.e001" xlink:type="simple"/><mml:math display="block" id="M1"><mml:mrow><mml:msub><mml:mi>π</mml:mi> <mml:mrow><mml:msub><mml:mi>j</mml:mi> <mml:mn>1</mml:mn></mml:msub> <mml:msub><mml:mi>j</mml:mi> <mml:mn>2</mml:mn></mml:msub></mml:mrow></mml:msub> <mml:mo>=</mml:mo> <mml:mi>P</mml:mi> <mml:mrow><mml:mo>(</mml:mo> <mml:msub><mml:mi>y</mml:mi> <mml:mn>1</mml:mn></mml:msub> <mml:mo>=</mml:mo> <mml:msub><mml:mi>j</mml:mi> <mml:mn>1</mml:mn></mml:msub> <mml:mo>,</mml:mo> <mml:msub><mml:mi>y</mml:mi> <mml:mn>2</mml:mn></mml:msub> <mml:mo>=</mml:mo> <mml:msub><mml:mi>j</mml:mi> <mml:mn>2</mml:mn></mml:msub> <mml:mo>|</mml:mo> <mml:mi>x</mml:mi> <mml:mo>)</mml:mo></mml:mrow> <mml:mo>,</mml:mo> <mml:mspace width="0.277778em"/><mml:mspace width="0.277778em"/><mml:msub><mml:mi>j</mml:mi> <mml:mn>1</mml:mn></mml:msub> <mml:mo>,</mml:mo> <mml:msub><mml:mi>j</mml:mi> <mml:mn>2</mml:mn></mml:msub> <mml:mo>=</mml:mo> <mml:mn>0</mml:mn> <mml:mo>,</mml:mo> <mml:mn>1</mml:mn> <mml:mo>,</mml:mo></mml:mrow></mml:math></alternatives> <label>(1)</label></disp-formula>
into the marginal probabilities of each couple member’s status to be positive, <italic>π</italic><sub><italic>F</italic></sub> = <italic>P</italic>(<italic>y</italic><sub>1</sub> = 1) = <italic>π</italic><sub>10</sub> + <italic>π</italic><sub>11</sub> and <italic>π</italic><sub><italic>M</italic></sub> = <italic>P</italic>(<italic>y</italic><sub>2</sub> = 1) = <italic>π</italic><sub>01</sub> + <italic>π</italic><sub>11</sub>, together with an association parameter <italic>ϕ</italic>. The effect of covariates on <italic>π</italic><sub><italic>F</italic></sub>, <italic>π</italic><sub><italic>M</italic></sub> and <italic>ϕ</italic> can be modelled as follows (extending the logistic regression model):
<disp-formula id="pone.0172959.e002"><alternatives><graphic id="pone.0172959.e002g" mimetype="image" position="anchor" xlink:href="info:doi/10.1371/journal.pone.0172959.e002" xlink:type="simple"/><mml:math display="block" id="M2"><mml:mtable><mml:mtr><mml:mtd><mml:mrow><mml:msub><mml:mi>h</mml:mi> <mml:mn>1</mml:mn></mml:msub> <mml:mrow><mml:mo>(</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mi>F</mml:mi></mml:msub> <mml:mo>)</mml:mo></mml:mrow> <mml:mo>=</mml:mo> <mml:msubsup><mml:mi>β</mml:mi> <mml:mn>1</mml:mn> <mml:mi>T</mml:mi></mml:msubsup> <mml:mi>x</mml:mi> <mml:mo>,</mml:mo></mml:mrow></mml:mtd></mml:mtr> <mml:mtr><mml:mtd><mml:mrow><mml:msub><mml:mi>h</mml:mi> <mml:mn>2</mml:mn></mml:msub> <mml:mrow><mml:mo>(</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mi>M</mml:mi></mml:msub> <mml:mo>)</mml:mo></mml:mrow> <mml:mo>=</mml:mo> <mml:msubsup><mml:mi>β</mml:mi> <mml:mn>2</mml:mn> <mml:mi>T</mml:mi></mml:msubsup> <mml:mi>x</mml:mi> <mml:mo>,</mml:mo></mml:mrow></mml:mtd></mml:mtr> <mml:mtr><mml:mtd><mml:mrow><mml:msub><mml:mi>h</mml:mi> <mml:mn>3</mml:mn></mml:msub> <mml:mrow><mml:mo>(</mml:mo> <mml:mi>ϕ</mml:mi> <mml:mo>)</mml:mo></mml:mrow> <mml:mo>=</mml:mo> <mml:msubsup><mml:mi>β</mml:mi> <mml:mn>3</mml:mn> <mml:mi>T</mml:mi></mml:msubsup> <mml:mi>x</mml:mi> <mml:mo>,</mml:mo></mml:mrow></mml:mtd></mml:mtr></mml:mtable></mml:math></alternatives> <label>(2)</label></disp-formula>
where <italic>h</italic><sub>1</sub>, <italic>h</italic><sub>2</sub> and <italic>h</italic><sub>3</sub> are appropriate link functions. Typical choices for <italic>h</italic><sub>1</sub> and <italic>h</italic><sub>2</sub> are the logit, probit, cloglog link (see e.g. [<xref ref-type="bibr" rid="pone.0172959.ref011">11</xref>]). Here we focus on the logit link as it allows an appealing interpretation of the effects of the covariates in terms of odds ratios. The choice of the link function <italic>h</italic><sub>3</sub> is related to the particular choice of the association parameter: the correlation coefficient <inline-formula id="pone.0172959.e003"><alternatives><graphic id="pone.0172959.e003g" mimetype="image" position="anchor" xlink:href="info:doi/10.1371/journal.pone.0172959.e003" xlink:type="simple"/><mml:math display="inline" id="M3"><mml:mrow><mml:mi>ρ</mml:mi> <mml:mo>=</mml:mo> <mml:mrow><mml:mo>(</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mn>11</mml:mn></mml:msub> <mml:mo>-</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mi>F</mml:mi></mml:msub> <mml:msub><mml:mi>π</mml:mi> <mml:mi>M</mml:mi></mml:msub> <mml:mo>)</mml:mo></mml:mrow> <mml:mo>/</mml:mo> <mml:msqrt><mml:mrow><mml:msub><mml:mi>π</mml:mi> <mml:mi>F</mml:mi></mml:msub> <mml:mrow><mml:mo>(</mml:mo> <mml:mn>1</mml:mn> <mml:mo>-</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mi>F</mml:mi></mml:msub> <mml:mo>)</mml:mo></mml:mrow> <mml:msub><mml:mi>π</mml:mi> <mml:mi>M</mml:mi></mml:msub> <mml:mrow><mml:mo>(</mml:mo> <mml:mn>1</mml:mn> <mml:mo>-</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mi>M</mml:mi></mml:msub> <mml:mo>)</mml:mo></mml:mrow></mml:mrow></mml:msqrt></mml:mrow></mml:math></alternatives></inline-formula>, the odds ratio OR = (<italic>π</italic><sub>00</sub><italic>π</italic><sub>11</sub>)/(<italic>π</italic><sub>10</sub><italic>π</italic><sub>01</sub>) or just <italic>π</italic><sub>11</sub>, or any other choice that more clearly represents the parameter of interest, being the serodiscordance in our setting.</p>
<p>The model components <xref ref-type="disp-formula" rid="pone.0172959.e002">Eq (2)</xref> can be embedded in different frameworks of estimation and inference. The standard application of Generalized Estimating Equations (GEE, moment estimation, see e.g. [<xref ref-type="bibr" rid="pone.0172959.ref011">11</xref>]) focuses on the two model components for <italic>π</italic><sub><italic>F</italic></sub> and <italic>π</italic><sub><italic>M</italic></sub> while treating the “intra-couple” correlation as a nuisance parameter, not allowing to incorporate any model <inline-formula id="pone.0172959.e004"><alternatives><graphic id="pone.0172959.e004g" mimetype="image" position="anchor" xlink:href="info:doi/10.1371/journal.pone.0172959.e004" xlink:type="simple"/><mml:math display="inline" id="M4"><mml:mrow><mml:msub><mml:mi>h</mml:mi> <mml:mn>3</mml:mn></mml:msub> <mml:mrow><mml:mo>(</mml:mo> <mml:mi>ϕ</mml:mi> <mml:mo>)</mml:mo></mml:mrow> <mml:mo>=</mml:mo> <mml:msubsup><mml:mi>β</mml:mi> <mml:mn>3</mml:mn> <mml:mi>T</mml:mi></mml:msubsup> <mml:mi>x</mml:mi></mml:mrow></mml:math></alternatives></inline-formula>. Extensions such as Alternating Logistic Regression(ALR) proposed by [<xref ref-type="bibr" rid="pone.0172959.ref012">12</xref>] extend standard GEE with the possibility to include a model <inline-formula id="pone.0172959.e005"><alternatives><graphic id="pone.0172959.e005g" mimetype="image" position="anchor" xlink:href="info:doi/10.1371/journal.pone.0172959.e005" xlink:type="simple"/><mml:math display="inline" id="M5"><mml:mrow><mml:msub><mml:mi>h</mml:mi> <mml:mn>3</mml:mn></mml:msub> <mml:mrow><mml:mo>(</mml:mo> <mml:mi>ϕ</mml:mi> <mml:mo>)</mml:mo></mml:mrow> <mml:mo>=</mml:mo> <mml:msubsup><mml:mi>β</mml:mi> <mml:mn>3</mml:mn> <mml:mi>T</mml:mi></mml:msubsup> <mml:mi>x</mml:mi></mml:mrow></mml:math></alternatives></inline-formula> and uses the OR rather than the correlation coefficient. But existing software is limiting the type of such models one can fit. Here we opt however for full Maximum Likelihood (ML), as this paradigm allows modelling in principle any association parameter by any kind of covariate model and is able to combine fixed effects for the covariates with random effects to represent other sources of hierarchical heterogeneity in a conceptually straightforward way (such as the EAs in our application).</p>
<p>The joint probabilities <italic>π</italic><sub><italic>j</italic><sub>1</sub><italic>j</italic><sub>2</sub></sub> used to construct the multinomial (log-)likelihood are reparametrized in terms of <italic>π</italic><sub><italic>F</italic></sub>, <italic>π</italic><sub><italic>M</italic></sub> and <italic>π</italic><sub>11</sub> as follows
<disp-formula id="pone.0172959.e006"><alternatives><graphic id="pone.0172959.e006g" mimetype="image" position="anchor" xlink:href="info:doi/10.1371/journal.pone.0172959.e006" xlink:type="simple"/><mml:math display="block" id="M6"><mml:mrow><mml:msub><mml:mi>π</mml:mi> <mml:mrow><mml:msub><mml:mi>j</mml:mi> <mml:mn>1</mml:mn></mml:msub> <mml:msub><mml:mi>j</mml:mi> <mml:mn>2</mml:mn></mml:msub></mml:mrow></mml:msub> <mml:mo>=</mml:mo> <mml:mfenced close="" open="{" separators=""><mml:mtable><mml:mtr><mml:mtd columnalign="left"><mml:msub><mml:mi>π</mml:mi> <mml:mn>11</mml:mn></mml:msub></mml:mtd> <mml:mtd columnalign="right"><mml:mtext>if</mml:mtext></mml:mtd> <mml:mtd columnalign="right"><mml:mrow><mml:msub><mml:mi>j</mml:mi> <mml:mn>1</mml:mn></mml:msub> <mml:mo>=</mml:mo> <mml:mn>1</mml:mn></mml:mrow></mml:mtd> <mml:mtd columnalign="right"><mml:mtext>and</mml:mtext></mml:mtd> <mml:mtd columnalign="right"><mml:mrow><mml:msub><mml:mi>j</mml:mi> <mml:mn>2</mml:mn></mml:msub> <mml:mo>=</mml:mo> <mml:mn>1</mml:mn> <mml:mo>,</mml:mo></mml:mrow></mml:mtd></mml:mtr> <mml:mtr><mml:mtd columnalign="left"><mml:mrow><mml:msub><mml:mi>π</mml:mi> <mml:mi>F</mml:mi></mml:msub> <mml:mo>-</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mn>11</mml:mn></mml:msub></mml:mrow></mml:mtd> <mml:mtd columnalign="right"><mml:mtext>if</mml:mtext></mml:mtd> <mml:mtd columnalign="right"><mml:mrow><mml:msub><mml:mi>j</mml:mi> <mml:mn>1</mml:mn></mml:msub> <mml:mo>=</mml:mo> <mml:mn>1</mml:mn></mml:mrow></mml:mtd> <mml:mtd columnalign="right"><mml:mtext>and</mml:mtext></mml:mtd> <mml:mtd columnalign="right"><mml:mrow><mml:msub><mml:mi>j</mml:mi> <mml:mn>2</mml:mn></mml:msub> <mml:mo>=</mml:mo> <mml:mn>0</mml:mn> <mml:mo>,</mml:mo></mml:mrow></mml:mtd></mml:mtr> <mml:mtr><mml:mtd columnalign="left"><mml:mrow><mml:msub><mml:mi>π</mml:mi> <mml:mi>M</mml:mi></mml:msub> <mml:mo>-</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mn>11</mml:mn></mml:msub></mml:mrow></mml:mtd> <mml:mtd columnalign="right"><mml:mtext>if</mml:mtext></mml:mtd> <mml:mtd columnalign="right"><mml:mrow><mml:msub><mml:mi>j</mml:mi> <mml:mn>1</mml:mn></mml:msub> <mml:mo>=</mml:mo> <mml:mn>0</mml:mn></mml:mrow></mml:mtd> <mml:mtd columnalign="right"><mml:mtext>and</mml:mtext></mml:mtd> <mml:mtd columnalign="right"><mml:mrow><mml:msub><mml:mi>j</mml:mi> <mml:mn>2</mml:mn></mml:msub> <mml:mo>=</mml:mo> <mml:mn>1</mml:mn> <mml:mo>,</mml:mo></mml:mrow></mml:mtd></mml:mtr> <mml:mtr><mml:mtd columnalign="left"><mml:mrow><mml:mn>1</mml:mn> <mml:mo>-</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mi>F</mml:mi></mml:msub> <mml:mo>-</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mi>M</mml:mi></mml:msub> <mml:mo>+</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mn>11</mml:mn></mml:msub></mml:mrow></mml:mtd> <mml:mtd columnalign="right"><mml:mtext>if</mml:mtext></mml:mtd> <mml:mtd columnalign="right"><mml:mrow><mml:msub><mml:mi>j</mml:mi> <mml:mn>1</mml:mn></mml:msub> <mml:mo>=</mml:mo> <mml:mn>0</mml:mn></mml:mrow></mml:mtd> <mml:mtd columnalign="right"><mml:mtext>and</mml:mtext></mml:mtd> <mml:mtd columnalign="right"><mml:mrow><mml:msub><mml:mi>j</mml:mi> <mml:mn>2</mml:mn></mml:msub> <mml:mo>=</mml:mo> <mml:mn>0</mml:mn> <mml:mo>.</mml:mo></mml:mrow></mml:mtd></mml:mtr></mml:mtable></mml:mfenced></mml:mrow></mml:math></alternatives> <label>(3)</label></disp-formula>
The joint probability <italic>π</italic><sub>11</sub> that both partners of the same couple are HIV positive can be replaced, as the association parameter, by the correlation coefficient <italic>ρ</italic>, or by the odds ratio, the latter being the most natural choice for binary data. The multinomial likelihood with probabilities <italic>π</italic><sub><italic>j</italic><sub>1</sub><italic>j</italic><sub>2</sub></sub>, via identities Eqs (<xref ref-type="disp-formula" rid="pone.0172959.e002">2</xref>) and (<xref ref-type="disp-formula" rid="pone.0172959.e006">3</xref>) written in terms of the regression parameters (<italic>β</italic><sub>1</sub>, <italic>β</italic><sub>2</sub>, <italic>β</italic><sub>3</sub>), is the basis for ML estimation and inference.</p>
<p>Within this ML framework, <xref ref-type="disp-formula" rid="pone.0172959.e002">model (2)</xref> with logit links for <italic>π</italic><sub><italic>F</italic></sub> and <italic>π</italic><sub><italic>M</italic></sub>, and the log link for <italic>ϕ</italic> = <monospace>OR</monospace> corresponds to the Bivariate Dale Model (BDM) [<xref ref-type="bibr" rid="pone.0172959.ref007">7</xref>]. The reparameterisation formulas for <italic>π</italic><sub>11</sub> and <monospace>OR</monospace> are given by
<disp-formula id="pone.0172959.e007"><alternatives><graphic id="pone.0172959.e007g" mimetype="image" position="anchor" xlink:href="info:doi/10.1371/journal.pone.0172959.e007" xlink:type="simple"/><mml:math display="block" id="M7"><mml:mrow><mml:mtext mathvariant="sans-serif">OR</mml:mtext> <mml:mo>=</mml:mo> <mml:mfrac><mml:mrow><mml:msub><mml:mi>π</mml:mi> <mml:mn>11</mml:mn></mml:msub> <mml:mrow><mml:mo>(</mml:mo> <mml:mn>1</mml:mn> <mml:mo>-</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mi>F</mml:mi></mml:msub> <mml:mo>-</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mi>M</mml:mi></mml:msub> <mml:mo>+</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mn>11</mml:mn></mml:msub> <mml:mo>)</mml:mo></mml:mrow></mml:mrow> <mml:mrow><mml:mrow><mml:mo>(</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mi>F</mml:mi></mml:msub> <mml:mo>-</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mn>11</mml:mn></mml:msub> <mml:mo>)</mml:mo></mml:mrow> <mml:mrow><mml:mo>(</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mi>M</mml:mi></mml:msub> <mml:mo>-</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mn>11</mml:mn></mml:msub> <mml:mo>)</mml:mo></mml:mrow></mml:mrow></mml:mfrac></mml:mrow></mml:math></alternatives> <label>(4)</label></disp-formula>
and, if <monospace>OR</monospace> ≠ 1
<disp-formula id="pone.0172959.e008"><alternatives><graphic id="pone.0172959.e008g" mimetype="image" position="anchor" xlink:href="info:doi/10.1371/journal.pone.0172959.e008" xlink:type="simple"/><mml:math display="block" id="M8"><mml:mtable><mml:mtr><mml:mtd><mml:mstyle displaystyle="true" scriptlevel="0"><mml:mrow><mml:msub><mml:mi>π</mml:mi> <mml:mn>11</mml:mn></mml:msub> <mml:mo>=</mml:mo> <mml:mfrac><mml:mrow><mml:mn>1</mml:mn> <mml:mo>+</mml:mo> <mml:mrow><mml:mo>(</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mi>F</mml:mi></mml:msub> <mml:mo>+</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mi>M</mml:mi></mml:msub> <mml:mo>)</mml:mo></mml:mrow> <mml:mrow><mml:mo>(</mml:mo> <mml:mtext mathvariant="sans-serif">OR</mml:mtext> <mml:mo>-</mml:mo> <mml:mn>1</mml:mn> <mml:mo>)</mml:mo></mml:mrow> <mml:mo>-</mml:mo> <mml:msup><mml:mrow><mml:mo>{</mml:mo> <mml:mrow><mml:mo>[</mml:mo> <mml:mn>1</mml:mn> <mml:mo>+</mml:mo> <mml:mrow><mml:mo>(</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mi>F</mml:mi></mml:msub> <mml:mo>+</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mi>M</mml:mi></mml:msub> <mml:mo>)</mml:mo></mml:mrow> <mml:msup><mml:mrow><mml:mo>(</mml:mo> <mml:mtext mathvariant="sans-serif">OR</mml:mtext> <mml:mo>-</mml:mo> <mml:mn>1</mml:mn> <mml:mo>)</mml:mo></mml:mrow> <mml:mn>2</mml:mn></mml:msup> <mml:mo>]</mml:mo></mml:mrow> <mml:mo>+</mml:mo> <mml:mn>4</mml:mn> <mml:mtext mathvariant="sans-serif">OR</mml:mtext> <mml:mrow><mml:mo>(</mml:mo> <mml:mn>1</mml:mn> <mml:mo>-</mml:mo> <mml:mtext mathvariant="sans-serif">OR</mml:mtext> <mml:mo>)</mml:mo></mml:mrow> <mml:msub><mml:mi>π</mml:mi> <mml:mi>F</mml:mi></mml:msub> <mml:msub><mml:mi>π</mml:mi> <mml:mi>M</mml:mi></mml:msub> <mml:mo>}</mml:mo></mml:mrow> <mml:mfrac><mml:mn>1</mml:mn> <mml:mn>2</mml:mn></mml:mfrac></mml:msup></mml:mrow> <mml:mrow><mml:mo>[</mml:mo> <mml:mn>2</mml:mn> <mml:mo>(</mml:mo> <mml:mtext mathvariant="sans-serif">OR</mml:mtext> <mml:mo>-</mml:mo> <mml:mn>1</mml:mn> <mml:mo>)</mml:mo> <mml:mo>]</mml:mo></mml:mrow></mml:mfrac></mml:mrow></mml:mstyle></mml:mtd></mml:mtr></mml:mtable></mml:math></alternatives> <label>(5)</label></disp-formula>
and when <monospace>OR</monospace> = 1, <italic>π</italic><sub>11</sub> = <italic>π</italic><sub><italic>F</italic></sub><italic>π</italic><sub><italic>M</italic></sub>. In general the OR is preferred for its interpretation and its orthogonality to marginal probabilities <italic>π</italic><sub><italic>F</italic></sub> and <italic>π</italic><sub><italic>M</italic></sub> (in the sense that the corresponding elements in the expected covariance matrix are zero). Although the <monospace>OR</monospace> is an attractive measure describing the association between binary variables, it treats concordant positive and concordant negative couples fully equally, inflating the <monospace>OR</monospace> as a measure for (con/dis)cordance. Furthermore, the magnitude of association between the two outcome may be misunderstood, leading to unrealistic estimates [<xref ref-type="bibr" rid="pone.0172959.ref013">13</xref>]. Indeed, as the majority of couples is concordant negative, the OR might be high even if there is only a small number of concordant positive couples. For example, for the INSIDA data, the overall sample estimates are <inline-formula id="pone.0172959.e009"><alternatives><graphic id="pone.0172959.e009g" mimetype="image" position="anchor" xlink:href="info:doi/10.1371/journal.pone.0172959.e009" xlink:type="simple"/><mml:math display="inline" id="M9"><mml:mrow><mml:msub><mml:mover accent="true"><mml:mi>π</mml:mi> <mml:mo>^</mml:mo></mml:mover> <mml:mn>00</mml:mn></mml:msub> <mml:mo>=</mml:mo> <mml:mn>0</mml:mn> <mml:mo>.</mml:mo> <mml:mn>84</mml:mn></mml:mrow></mml:math></alternatives></inline-formula> (95% CI [0.82,0.85]) and <inline-formula id="pone.0172959.e010"><alternatives><graphic id="pone.0172959.e010g" mimetype="image" position="anchor" xlink:href="info:doi/10.1371/journal.pone.0172959.e010" xlink:type="simple"/><mml:math display="inline" id="M10"><mml:mrow><mml:mover accent="true"><mml:mtext>OR</mml:mtext> <mml:mo>^</mml:mo></mml:mover> <mml:mo>=</mml:mo> <mml:mn>14</mml:mn> <mml:mo>.</mml:mo> <mml:mn>75</mml:mn></mml:mrow></mml:math></alternatives></inline-formula> (95% CI [10.16,19.34]), so both very high, whereas the estimate for concordant positive couples is small: <inline-formula id="pone.0172959.e011"><alternatives><graphic id="pone.0172959.e011g" mimetype="image" position="anchor" xlink:href="info:doi/10.1371/journal.pone.0172959.e011" xlink:type="simple"/><mml:math display="inline" id="M11"><mml:mrow><mml:msub><mml:mover accent="true"><mml:mi>π</mml:mi> <mml:mo>^</mml:mo></mml:mover> <mml:mn>11</mml:mn></mml:msub> <mml:mo>=</mml:mo> <mml:mn>0</mml:mn> <mml:mo>.</mml:mo> <mml:mn>05</mml:mn></mml:mrow></mml:math></alternatives></inline-formula> (95% CI [0.04,0.06]). For the purpose of our particular research question, interest goes to couples being discordant (or concordant), given that at least one of both is positive. This type of association measure is introduced and discussed in the next subsection.</p>
</sec>
<sec id="sec007">
<title>A new serodiscordance measure</title>
<p>In a completely different context of measuring synchrony in neuronal firing, [<xref ref-type="bibr" rid="pone.0172959.ref006">6</xref>] proposed a new measure of synchrony, the conditional synchrony measure CSM, which is the probability of two neurons firing together, given that at least one of the two is active. Faes <italic>et al</italic>. [<xref ref-type="bibr" rid="pone.0172959.ref006">6</xref>] state that, although the odds ratio is an attractive association measure with nice mathematical properties (such as the absence of range restrictions, regardless of the marginal probabilities), it is less suitable to quantify synchrony due to its symmetry, treating 0–0 matches of equal importance as 1–1 matches. The CSM is defined as
<disp-formula id="pone.0172959.e012"><alternatives><graphic id="pone.0172959.e012g" mimetype="image" position="anchor" xlink:href="info:doi/10.1371/journal.pone.0172959.e012" xlink:type="simple"/><mml:math display="block" id="M12"><mml:mrow><mml:mtext mathvariant="sans-serif">CSM</mml:mtext> <mml:mo>=</mml:mo> <mml:mfrac><mml:msub><mml:mi>π</mml:mi> <mml:mn>11</mml:mn></mml:msub> <mml:mrow><mml:msub><mml:mi>π</mml:mi> <mml:mi>F</mml:mi></mml:msub> <mml:mo>+</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mi>M</mml:mi></mml:msub> <mml:mo>-</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mn>11</mml:mn></mml:msub></mml:mrow></mml:mfrac> <mml:mo>,</mml:mo></mml:mrow></mml:math></alternatives></disp-formula>
and translated to our application, it is the conditional probability that both, man and woman, are HIV positive, given that at least one of them, man or woman, is HIV positive.</p>
<p>As we are rather interested in discordance, we define the HIV serodiscordance measure SDM as the conditional probability that the couple is HIV discordant, given that at least one of them, man or woman, is HIV positive, or
<disp-formula id="pone.0172959.e013"><alternatives><graphic id="pone.0172959.e013g" mimetype="image" position="anchor" xlink:href="info:doi/10.1371/journal.pone.0172959.e013" xlink:type="simple"/><mml:math display="block" id="M13"><mml:mrow><mml:mtext mathvariant="sans-serif">SDM</mml:mtext> <mml:mo>=</mml:mo> <mml:mn>1</mml:mn> <mml:mo>-</mml:mo> <mml:mtext mathvariant="sans-serif">CSM</mml:mtext> <mml:mo>=</mml:mo> <mml:mfrac><mml:mrow><mml:msub><mml:mi>π</mml:mi> <mml:mi>F</mml:mi></mml:msub> <mml:mo>+</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mi>M</mml:mi></mml:msub> <mml:mo>-</mml:mo> <mml:mn>2</mml:mn> <mml:msub><mml:mi>π</mml:mi> <mml:mn>11</mml:mn></mml:msub></mml:mrow> <mml:mrow><mml:msub><mml:mi>π</mml:mi> <mml:mi>F</mml:mi></mml:msub> <mml:mo>+</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mi>M</mml:mi></mml:msub> <mml:mo>-</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mn>11</mml:mn></mml:msub></mml:mrow></mml:mfrac> <mml:mo>.</mml:mo></mml:mrow></mml:math></alternatives> <label>(6)</label></disp-formula>
Being a (conditional) probability, SDM takes values between 0 and 1, with higher values indicating more HIV discordance within the couple. The global sample estimate for the INSIDA data, <inline-formula id="pone.0172959.e014"><alternatives><graphic id="pone.0172959.e014g" mimetype="image" position="anchor" xlink:href="info:doi/10.1371/journal.pone.0172959.e014" xlink:type="simple"/><mml:math display="inline" id="M14"><mml:mrow><mml:mover accent="true"><mml:mtext mathvariant="sans-serif">SDM</mml:mtext> <mml:mo>^</mml:mo></mml:mover> <mml:mo>=</mml:mo> <mml:mn>0</mml:mn> <mml:mo>.</mml:mo> <mml:mn>6747</mml:mn></mml:mrow></mml:math></alternatives></inline-formula> (95% CI [0.63,0.72]), indicates a quite high amount of discordance. It is the objective to examine how the SDM depends on different factors, while accounting for the effects of (possibly) partly common, partly different factors on the marginal probability <italic>π</italic><sub><italic>F</italic></sub> and <italic>π</italic><sub><italic>M</italic></sub>. A logit link function is used for <italic>h</italic><sub>3</sub> in the definition of the model components <xref ref-type="disp-formula" rid="pone.0172959.e002">Eq (2)</xref>, allowing to interpret effects of factors as odds ratios contrasting the probabilities to be discordant to that of being concordant, given that at least one member of the couple is HIV positive. The joint probability <italic>π</italic><sub>11</sub>, used in the loglikelihood expressions based on <xref ref-type="disp-formula" rid="pone.0172959.e006">Eq (3)</xref>, can be expressed in terms of the SDM, <italic>π</italic><sub><italic>F</italic></sub> and <italic>π</italic><sub><italic>M</italic></sub> as follows
<disp-formula id="pone.0172959.e015"><alternatives><graphic id="pone.0172959.e015g" mimetype="image" position="anchor" xlink:href="info:doi/10.1371/journal.pone.0172959.e015" xlink:type="simple"/><mml:math display="block" id="M15"><mml:mtable><mml:mtr><mml:mtd><mml:mstyle displaystyle="true" scriptlevel="0"><mml:mrow><mml:msub><mml:mi>π</mml:mi> <mml:mn>11</mml:mn></mml:msub> <mml:mo>=</mml:mo> <mml:mfrac><mml:mrow><mml:mn>1</mml:mn> <mml:mo>-</mml:mo> <mml:mtext mathvariant="sans-serif">SDM</mml:mtext></mml:mrow> <mml:mrow><mml:mn>2</mml:mn> <mml:mo>-</mml:mo> <mml:mtext mathvariant="sans-serif">SDM</mml:mtext></mml:mrow></mml:mfrac> <mml:mrow><mml:mo>[</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mi>F</mml:mi></mml:msub> <mml:mo>+</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mi>M</mml:mi></mml:msub> <mml:mo>]</mml:mo></mml:mrow> <mml:mo>.</mml:mo></mml:mrow></mml:mstyle></mml:mtd></mml:mtr></mml:mtable></mml:math></alternatives> <label>(7)</label></disp-formula></p>
<p>In the next section the fixed effects <xref ref-type="disp-formula" rid="pone.0172959.e002">model (2)</xref> using association parameter <italic>ϕ</italic> = <monospace>SDM</monospace> and logit links for all three parameter models, is extended with random EA effects to accommodate heterogeneity across EAs.</p>
</sec>
<sec id="sec008">
<title>Extension with random effects to capture heterogeneity at the EA level</title>
<p>Let <italic>y</italic><sub><italic>ij</italic></sub> = (<italic>y</italic><sub><italic>ij</italic>1</sub>, <italic>y</italic><sub><italic>ij</italic>2</sub>) denote the HIV status of a couple <italic>j</italic> = 1, …, <italic>n</italic><sub><italic>i</italic></sub> in enumeration area <italic>i</italic> with <italic>n</italic><sub><italic>i</italic></sub> sampled couples, <italic>i</italic> = 1, …, <italic>N</italic>. For the INSIDA data, the total number of enumeration areas is <italic>N</italic> = 270 and <inline-formula id="pone.0172959.e016"><alternatives><graphic id="pone.0172959.e016g" mimetype="image" position="anchor" xlink:href="info:doi/10.1371/journal.pone.0172959.e016" xlink:type="simple"/><mml:math display="inline" id="M16"><mml:mrow><mml:msubsup><mml:mo>∑</mml:mo> <mml:mrow><mml:mi>i</mml:mi> <mml:mo>=</mml:mo> <mml:mn>1</mml:mn></mml:mrow> <mml:mi>N</mml:mi></mml:msubsup> <mml:msub><mml:mi>n</mml:mi> <mml:mi>i</mml:mi></mml:msub> <mml:mo>=</mml:mo> <mml:mn>2159</mml:mn></mml:mrow></mml:math></alternatives></inline-formula> is the total number of couples. As <italic>π</italic><sub><italic>F</italic></sub>, <italic>π</italic><sub><italic>M</italic></sub>, and SDM are expected to be heterogeneous across EAs, <xref ref-type="disp-formula" rid="pone.0172959.e002">model (2)</xref> has to be extended with an EA effect. <xref ref-type="disp-formula" rid="pone.0172959.e002">Model (2)</xref> is therefore extended with EA random effects as follows, adding subscripts <italic>i</italic> and <italic>j</italic> and expressing that the covariates <italic>x</italic><sub>1,<italic>ij</italic></sub>, <italic>x</italic><sub>2,<italic>ij</italic></sub> and <italic>x</italic><sub>3,<italic>ij</italic></sub> can be possibly different subvectors of the full covariate/factor vector <italic>x</italic><sub><italic>ij</italic></sub>,
<disp-formula id="pone.0172959.e017"><alternatives><graphic id="pone.0172959.e017g" mimetype="image" position="anchor" xlink:href="info:doi/10.1371/journal.pone.0172959.e017" xlink:type="simple"/><mml:math display="block" id="M17"><mml:mtable><mml:mtr><mml:mtd><mml:mrow><mml:mtext>logit</mml:mtext> <mml:mrow><mml:mo>(</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mrow><mml:mi>F</mml:mi> <mml:mo>,</mml:mo> <mml:mi>i</mml:mi> <mml:mi>j</mml:mi></mml:mrow></mml:msub> <mml:mo>)</mml:mo></mml:mrow> <mml:mo>=</mml:mo> <mml:msubsup><mml:mi>β</mml:mi> <mml:mn>1</mml:mn> <mml:mi>T</mml:mi></mml:msubsup> <mml:msub><mml:mi>x</mml:mi> <mml:mrow><mml:mn>1</mml:mn> <mml:mo>,</mml:mo> <mml:mi>i</mml:mi> <mml:mi>j</mml:mi></mml:mrow></mml:msub> <mml:mo>+</mml:mo> <mml:msub><mml:mi>b</mml:mi> <mml:mrow><mml:mi>F</mml:mi> <mml:mo>,</mml:mo> <mml:mi>i</mml:mi></mml:mrow></mml:msub> <mml:mo>,</mml:mo></mml:mrow></mml:mtd></mml:mtr> <mml:mtr><mml:mtd><mml:mrow><mml:mtext>logit</mml:mtext> <mml:mrow><mml:mo>(</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mrow><mml:mi>M</mml:mi> <mml:mo>,</mml:mo> <mml:mi>i</mml:mi> <mml:mi>j</mml:mi></mml:mrow></mml:msub> <mml:mo>)</mml:mo></mml:mrow> <mml:mo>=</mml:mo> <mml:msubsup><mml:mi>β</mml:mi> <mml:mn>2</mml:mn> <mml:mi>T</mml:mi></mml:msubsup> <mml:msub><mml:mi>x</mml:mi> <mml:mrow><mml:mn>2</mml:mn> <mml:mo>,</mml:mo> <mml:mi>i</mml:mi> <mml:mi>j</mml:mi></mml:mrow></mml:msub> <mml:mo>+</mml:mo> <mml:msub><mml:mi>b</mml:mi> <mml:mrow><mml:mi>M</mml:mi> <mml:mo>,</mml:mo> <mml:mi>i</mml:mi></mml:mrow></mml:msub> <mml:mo>,</mml:mo></mml:mrow></mml:mtd></mml:mtr> <mml:mtr><mml:mtd><mml:mrow><mml:mtext>logit</mml:mtext> <mml:mrow><mml:mo>(</mml:mo> <mml:msub><mml:mtext mathvariant="sans-serif">SDM</mml:mtext> <mml:mrow><mml:mi>i</mml:mi> <mml:mi>j</mml:mi></mml:mrow></mml:msub> <mml:mo>)</mml:mo></mml:mrow> <mml:mo>=</mml:mo> <mml:msubsup><mml:mi>β</mml:mi> <mml:mn>3</mml:mn> <mml:mi>T</mml:mi></mml:msubsup> <mml:msub><mml:mi>x</mml:mi> <mml:mrow><mml:mn>3</mml:mn> <mml:mo>,</mml:mo> <mml:mi>i</mml:mi> <mml:mi>j</mml:mi></mml:mrow></mml:msub> <mml:mo>+</mml:mo> <mml:msub><mml:mi>b</mml:mi> <mml:mrow><mml:mi>D</mml:mi> <mml:mo>,</mml:mo> <mml:mi>i</mml:mi></mml:mrow></mml:msub> <mml:mo>,</mml:mo></mml:mrow></mml:mtd></mml:mtr></mml:mtable></mml:math></alternatives> <label>(8)</label></disp-formula>
where
<disp-formula id="pone.0172959.e018"><alternatives><graphic id="pone.0172959.e018g" mimetype="image" position="anchor" xlink:href="info:doi/10.1371/journal.pone.0172959.e018" xlink:type="simple"/><mml:math display="block" id="M18"><mml:mrow><mml:mrow><mml:mo>(</mml:mo> <mml:msub><mml:mi>b</mml:mi> <mml:mrow><mml:mi>F</mml:mi> <mml:mo>,</mml:mo> <mml:mi>i</mml:mi></mml:mrow></mml:msub> <mml:mo>,</mml:mo> <mml:msub><mml:mi>b</mml:mi> <mml:mrow><mml:mi>M</mml:mi> <mml:mo>,</mml:mo> <mml:mi>i</mml:mi></mml:mrow></mml:msub> <mml:mo>,</mml:mo> <mml:msub><mml:mi>b</mml:mi> <mml:mrow><mml:mi>D</mml:mi> <mml:mo>,</mml:mo> <mml:mi>i</mml:mi></mml:mrow></mml:msub> <mml:mo>)</mml:mo></mml:mrow> <mml:mo>∼</mml:mo> <mml:msub><mml:mi>N</mml:mi> <mml:mn>3</mml:mn></mml:msub> <mml:mrow><mml:mo>(</mml:mo> <mml:mn>0</mml:mn> <mml:mo>,</mml:mo> <mml:mo>Σ</mml:mo> <mml:mo>)</mml:mo></mml:mrow> <mml:mo>,</mml:mo></mml:mrow></mml:math></alternatives> <label>(9)</label></disp-formula>
are distributed as a trivariate normal distribtion with mean zero-vector and covariance matrix Σ. Note that <italic>x</italic><sub>ℓ,<italic>ij</italic></sub> (ℓ = 1, 2, 3) are vectors of covariates, some of these covariates are specific for the male/female individual, some specific for the couple, and some are at the level of province.</p>
<p>The variance components <inline-formula id="pone.0172959.e019"><alternatives><graphic id="pone.0172959.e019g" mimetype="image" position="anchor" xlink:href="info:doi/10.1371/journal.pone.0172959.e019" xlink:type="simple"/><mml:math display="inline" id="M19"><mml:mrow><mml:msubsup><mml:mi>σ</mml:mi> <mml:mi>F</mml:mi> <mml:mn>2</mml:mn></mml:msubsup> <mml:mo>,</mml:mo> <mml:msubsup><mml:mi>σ</mml:mi> <mml:mi>M</mml:mi> <mml:mn>2</mml:mn></mml:msubsup></mml:mrow></mml:math></alternatives></inline-formula> and <inline-formula id="pone.0172959.e020"><alternatives><graphic id="pone.0172959.e020g" mimetype="image" position="anchor" xlink:href="info:doi/10.1371/journal.pone.0172959.e020" xlink:type="simple"/><mml:math display="inline" id="M20"><mml:msubsup><mml:mi>σ</mml:mi> <mml:mi>D</mml:mi> <mml:mn>2</mml:mn></mml:msubsup></mml:math></alternatives></inline-formula> quantify the degree of heterogeneity across the AEs for each of the three parameters. Different choices considered for the covariance matrix Σ = <italic>VRV</italic> with
<disp-formula id="pone.0172959.e021"><alternatives><graphic id="pone.0172959.e021g" mimetype="image" position="anchor" xlink:href="info:doi/10.1371/journal.pone.0172959.e021" xlink:type="simple"/><mml:math display="block" id="M21"><mml:mrow><mml:mi>V</mml:mi> <mml:mo>=</mml:mo> <mml:mfenced close=")" open="(" separators=""><mml:mtable><mml:mtr><mml:mtd><mml:msub><mml:mi>σ</mml:mi> <mml:mi>F</mml:mi></mml:msub></mml:mtd> <mml:mtd><mml:mn>0</mml:mn></mml:mtd> <mml:mtd><mml:mn>0</mml:mn></mml:mtd></mml:mtr> <mml:mtr><mml:mtd><mml:mn>0</mml:mn></mml:mtd> <mml:mtd><mml:msub><mml:mi>σ</mml:mi> <mml:mi>M</mml:mi></mml:msub></mml:mtd> <mml:mtd><mml:mn>0</mml:mn></mml:mtd></mml:mtr> <mml:mtr><mml:mtd><mml:mn>0</mml:mn></mml:mtd> <mml:mtd><mml:mn>0</mml:mn></mml:mtd> <mml:mtd><mml:msub><mml:mi>σ</mml:mi> <mml:mi>D</mml:mi></mml:msub></mml:mtd></mml:mtr></mml:mtable></mml:mfenced> <mml:mo>,</mml:mo></mml:mrow></mml:math></alternatives></disp-formula>
are based on different choices for the correlation matrix <italic>R</italic>.</p>
<p>(Partial) Correlated Random Effects</p>
<p>A first option is a fully correlated random effects structure with
<disp-formula id="pone.0172959.e022"><alternatives><graphic id="pone.0172959.e022g" mimetype="image" position="anchor" xlink:href="info:doi/10.1371/journal.pone.0172959.e022" xlink:type="simple"/><mml:math display="block" id="M22"><mml:mrow><mml:msub><mml:mi>R</mml:mi> <mml:mtext>COR</mml:mtext></mml:msub> <mml:mo>=</mml:mo> <mml:mfenced close=")" open="(" separators=""><mml:mtable><mml:mtr><mml:mtd><mml:mn>1</mml:mn></mml:mtd> <mml:mtd><mml:msub><mml:mi>ρ</mml:mi> <mml:mrow><mml:mi>F</mml:mi> <mml:mi>M</mml:mi></mml:mrow></mml:msub></mml:mtd> <mml:mtd><mml:msub><mml:mi>ρ</mml:mi> <mml:mrow><mml:mi>F</mml:mi> <mml:mi>D</mml:mi></mml:mrow></mml:msub></mml:mtd></mml:mtr> <mml:mtr><mml:mtd><mml:msub><mml:mi>ρ</mml:mi> <mml:mrow><mml:mi>F</mml:mi> <mml:mi>M</mml:mi></mml:mrow></mml:msub></mml:mtd> <mml:mtd><mml:mn>1</mml:mn></mml:mtd> <mml:mtd><mml:msub><mml:mi>ρ</mml:mi> <mml:mrow><mml:mi>M</mml:mi> <mml:mi>D</mml:mi></mml:mrow></mml:msub></mml:mtd></mml:mtr> <mml:mtr><mml:mtd><mml:msub><mml:mi>ρ</mml:mi> <mml:mrow><mml:mi>F</mml:mi> <mml:mi>D</mml:mi></mml:mrow></mml:msub></mml:mtd> <mml:mtd><mml:msub><mml:mi>ρ</mml:mi> <mml:mrow><mml:mi>M</mml:mi> <mml:mi>D</mml:mi></mml:mrow></mml:msub></mml:mtd> <mml:mtd><mml:mn>1</mml:mn></mml:mtd></mml:mtr></mml:mtable></mml:mfenced> <mml:mo>.</mml:mo></mml:mrow></mml:math></alternatives></disp-formula>
A reduced partial-correlated version assumes only one correlation component to be non-zero: <italic>ρ</italic><sub><italic>FM</italic></sub> ≠ 0 and <italic>ρ</italic><sub><italic>FD</italic></sub> = <italic>ρ</italic><sub><italic>MD</italic></sub> = 0, so HIV status and SDM random effects are not correlated.</p>
<p>(Partial) Shared Random Effects</p>
<p>Another simplified structure corresponds to <italic>ρ</italic><sub><italic>FD</italic></sub> = <italic>ρ</italic><sub><italic>MD</italic></sub> = <italic>ρ</italic><sub><italic>FM</italic></sub> = 1 in <italic>R</italic><sub>COR</sub> leading to <italic>R</italic><sub>COR</sub> = <italic>J</italic><sub>3</sub>, the all-ones matrix. This implies a full-shared random effect, with <italic>b</italic><sub><italic>M</italic>,<italic>i</italic></sub> = (<italic>σ</italic><sub><italic>M</italic></sub>/<italic>σ</italic><sub><italic>F</italic></sub>)<italic>b</italic><sub><italic>F</italic>,<italic>i</italic></sub> and <italic>b</italic><sub><italic>D</italic>,<italic>i</italic></sub> = (<italic>σ</italic><sub><italic>D</italic></sub>/<italic>σ</italic><sub><italic>F</italic></sub>)<italic>b</italic><sub><italic>F</italic>,<italic>i</italic></sub>. This shared version is simplified further by putting <italic>σ</italic><sub><italic>D</italic></sub> = 0, leaving only (two) shared random effects in the models for <italic>π</italic><sub><italic>F</italic>,<italic>ij</italic></sub> and <italic>π</italic><sub><italic>M</italic>,<italic>ij</italic></sub> and assuming SDM<sub><italic>ij</italic></sub> to be constant across AEs (after correcting for covariates). This model will be referred to as the partial-shared random effects model.</p>
<p>Partial Equal Random Effects</p>
<p>A final reduction concerns setting <italic>σ</italic> = <italic>σ</italic><sub><italic>M</italic></sub> = <italic>σ</italic><sub><italic>F</italic></sub> in the partial-shared model, such that <italic>b</italic><sub><italic>M</italic>,<italic>i</italic></sub> = <italic>b</italic><sub><italic>F</italic>,<italic>i</italic></sub>, resulting in the partial-equal random effects model.</p>
<p>Independent Random Effects</p>
<p>As a final model, consider <italic>ρ</italic><sub><italic>FD</italic></sub> = <italic>ρ</italic><sub><italic>MD</italic></sub> = <italic>ρ</italic><sub><italic>FM</italic></sub> = 0 resulting in <italic>R</italic><sub>COR</sub> = <italic>I</italic><sub>3</sub>, the identity matrix, and corresponding to independent random effects.</p>
</sec>
<sec id="sec009">
<title>Accounting for survey design using weights</title>
<p>Most of the sample designs for household surveys such as INSIDA are complex and involve stratification, multistage sampling, and unequal sampling rates. Such survey designs are often more appropriate as coverage of the entire population of interest is better achieved and are often practically more efficient for interviewing subjects [<xref ref-type="bibr" rid="pone.0172959.ref014">14</xref>]. But it is necessary to account for the particular survey design in the statistical analyses using appropriate weights. For details on the calculation of the weights as used in our analyses, we refer to [<xref ref-type="bibr" rid="pone.0172959.ref001">1</xref>].</p>
</sec>
<sec id="sec010">
<title>Model building</title>
<p>For the selection of the covariates (shown in <xref ref-type="table" rid="pone.0172959.t001">Table 1</xref>), the following model building strategy was followed. Ideally one would select the covariates in all three model components <xref ref-type="disp-formula" rid="pone.0172959.e017">Eq (8)</xref> starting from the most complicated random effects type of <xref ref-type="disp-formula" rid="pone.0172959.e018">model (9)</xref>, but that appeared to be computationally not feasible. Therefore, as a pragmatic but reasonable alternative, in our view, stepwise regression was first applied to the two univariate logistic regression models for <italic>π</italic><sub><italic>F</italic></sub> and <italic>π</italic><sub><italic>M</italic></sub> separately. Later, these covariates were used in the joint marginal as well as in its extension to the AE-specific random effects models, hereby extending the logistic model for SDM in a stepwise forward manner. In a last step it was investigated whether the intercepts and some or all of the slopes in the two univariate logistic regression models for <italic>π</italic><sub><italic>F</italic></sub> and <italic>π</italic><sub><italic>M</italic></sub> could be taken in common. The final joint model was selected using the Akaike Information Criterion (smaller values of <monospace>AIC</monospace> indicate a better fit, [<xref ref-type="bibr" rid="pone.0172959.ref015">15</xref>]).</p>
<p>Data analysis was performed in SAS version 9.3 using PROC NLMIXED (code in <xref ref-type="sec" rid="sec013">Appendix 1</xref>).</p>
</sec>
</sec>
</sec>
<sec id="sec011" sec-type="results">
<title>Results</title>
<p>As the selected set of covariates for the two univariate logistic regression models for <italic>π</italic><sub><italic>F</italic></sub> and <italic>π</italic><sub><italic>M</italic></sub> only differs in one covariate for each model component, it was decided to keep the set equal for both, allowing a more complete comparison and interpretation of the estimated effects on the probability to be HIV positive for the female and the male partner within the same couple. The covariates involved are: the HIV prevalence of the province, the union number of the woman and that of the man, the condom use by the man, and the wealth index. The stepwise procedure to identify the factors influencing the serodiscordance measure SDM led to significant effects of the HIV prevalence of the province, and the union number of the woman.</p>
<p>The model components <xref ref-type="disp-formula" rid="pone.0172959.e017">Eq (8)</xref> with these selected sets of covariates were then fitted simultaneously with different (multivariate) random effects structures, capturing the heterogeneity across the EAs. <xref ref-type="table" rid="pone.0172959.t002">Table 2</xref> shows the values of -2×log-likelihood, the number of parameters, the AIC and BIC goodness of fit measures of all fitted models: the joint marginal model without random AE effects and a subset of all random effects models, with different and with partly common slopes for the models for <italic>π</italic><sub><italic>F</italic></sub> and <italic>π</italic><sub><italic>M</italic></sub>. The last two columns shows the models’ ranks according to AIC and BIC. The independent random effects and the full or partial random effects models did not convergence. So the results and the discussion is limited to the full-, partial-shared and partial equal random effects models.</p>
<table-wrap id="pone.0172959.t002" position="float">
<object-id pub-id-type="doi">10.1371/journal.pone.0172959.t002</object-id>
<label>Table 2</label>
<caption>
<title>Comparison of marginal model, and full-shared, partial-shared and partial-equal random effects models, all without or with common intercept and common slope for HIV prevalence and wealth index for the models for <italic>π</italic><sub><italic>F</italic></sub> and <italic>π</italic><sub><italic>M</italic></sub>.</title>
<p>The column ‘-2ll’ shows the values of -2×log-likelihood; the column ‘#Par’ shows the number of parameters and the columns ‘Rank’ refers to the ranking of the models according to the AIC and BIC criterion.</p>
</caption>
<alternatives>
<graphic id="pone.0172959.t002g" mimetype="image" position="float" xlink:href="info:doi/10.1371/journal.pone.0172959.t002" xlink:type="simple"/>
<table border="0" frame="box" rules="all">
<colgroup>
<col align="left" valign="middle"/>
<col align="left" valign="middle"/>
<col align="left" valign="middle"/>
<col align="left" valign="middle"/>
<col align="left" valign="middle"/>
<col align="left" valign="middle"/>
<col align="left" valign="middle"/>
</colgroup>
<thead>
<tr>
<th align="left">Model</th>
<th align="center">
<monospace>-2ll</monospace>
</th>
<th align="center"># <monospace>Par</monospace></th>
<th align="center">
<monospace>AIC</monospace>
</th>
<th align="center">
<monospace>BIC</monospace>
</th>
<th align="center">
<monospace>Rank<sub><italic>A</italic></sub></monospace>
</th>
<th align="center">
<monospace>Rank<sub><italic>B</italic></sub></monospace>
</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" style="background-color:#fbfbfb">MM (marginal model)</td>
<td align="char" char="." style="background-color:#fbfbfb">2584.8</td>
<td align="center" style="background-color:#fbfbfb">20</td>
<td align="char" char="." style="background-color:#fbfbfb">2624.8</td>
<td align="char" char="." style="background-color:#fbfbfb">2738.3</td>
<td align="center" style="background-color:#fbfbfb">8</td>
<td align="center" style="background-color:#fbfbfb">8</td>
</tr>
<tr>
<td align="left" style="background-color:#fbfbfb">CE-MM (common effects MM)</td>
<td align="char" char="." style="background-color:#fbfbfb">2590.4</td>
<td align="center" style="background-color:#fbfbfb">15</td>
<td align="char" char="." style="background-color:#fbfbfb">2620.4</td>
<td align="char" char="." style="background-color:#fbfbfb">2705.5</td>
<td align="center" style="background-color:#fbfbfb">7</td>
<td align="center" style="background-color:#fbfbfb">7</td>
</tr>
<tr>
<td align="left">FS (full-shared RE model)</td>
<td align="char" char=".">2534.2</td>
<td align="center">23</td>
<td align="char" char=".">2580.2</td>
<td align="char" char=".">2662.8</td>
<td align="center">6</td>
<td align="center">6</td>
</tr>
<tr>
<td align="left">CE-FS (common effects FS)</td>
<td align="char" char=".">2538.0</td>
<td align="center">18</td>
<td align="char" char=".">2574.0</td>
<td align="char" char=".">2638.7</td>
<td align="center">4</td>
<td align="center">3</td>
</tr>
<tr>
<td align="left">PS (partial-shared RE model)</td>
<td align="char" char=".">2534.2</td>
<td align="center">22</td>
<td align="char" char=".">2578.2</td>
<td align="char" char=".">2657.2</td>
<td align="center">5</td>
<td align="center">5</td>
</tr>
<tr>
<td align="left">CE-PS (common effects PS)</td>
<td align="char" char=".">2538.1</td>
<td align="center">17</td>
<td align="char" char=".">2572.1</td>
<td align="char" char=".">2633.2</td>
<td align="center">2</td>
<td align="center">2</td>
</tr>
<tr>
<td align="left">PE (partial-equal RE model)</td>
<td align="char" char=".">2534.7</td>
<td align="center">21</td>
<td align="char" char=".">2576.7</td>
<td align="char" char=".">2652.1</td>
<td align="center">3</td>
<td align="center">4</td>
</tr>
<tr>
<td align="left">CE-PE (common effects PE)</td>
<td align="char" char=".">2539.8</td>
<td align="center">16</td>
<td align="char" char=".">2571.8</td>
<td align="char" char=".">2629.3</td>
<td align="center">1</td>
<td align="center">1</td>
</tr>
</tbody>
</table>
</alternatives>
</table-wrap>
<p>According to AIC and BIC the best model is the partial-equal random effects model with common effects (CE-PE) for the HIV prevalence of the province and the wealth index, rather closely followed by the partial-shared random effects model with the same common effects (CE-PS). Based on the likelihood ratio test (LRT), the null hypothesis of equal random effects <italic>σ</italic><sub><italic>M</italic></sub> = <italic>σ</italic><sub><italic>F</italic></sub> cannot be rejected at level 0.05 (p-val = 0.19). So, the CE-PE model is taken as final model.</p>
<p>The estimates of all 16 parameters of the CE-PE model and their standard error estimates are shown in <xref ref-type="table" rid="pone.0172959.t003">Table 3</xref>. In the following discussion 95% confidence intervals for the odds ratio (OR-CI) are also shown. First of all, the probability to be HIV positive increases with the HIV prevalence in the province, for both women and men in the same way, as to be expected. As compared to a province with less than 5% prevalence, the odds to be HIV positive increases with a factor of about 6.9 when living in a province with more than 15% prevalence (OR-CI [4.10,11.62]). The effect of wealth index is also identical for women and men: as compared to the “richer” reference category, the odds to be HIV positive is estimated to be about 32% lower for the “middle” category, and 46% lower for “poorer” category (OR-CI [0.48,0.96] and [0.38.0.75] respectively). In Mozambique it is quite usual that one (or both but especially the male) partner within a richer couple has multiple extramarital partners or even visits sex workers, increasing the risk for HIV infection [<xref ref-type="bibr" rid="pone.0172959.ref016">16</xref>], [<xref ref-type="bibr" rid="pone.0172959.ref017">17</xref>]. The effects of union number and condom use are different for women and men. Whereas condom use has no significant effect on the odds for men to be HIV positive, not using a condom is estimated to increase the odds for women to be HIV positive with about 90% (OR-CI [1.16,3.12]). That the male partner has been married or lived with a female partner more than once has no significant effect on the odds for the woman to be HIV positive, but it increases the odds for the man to be HIV positive with about 45% (OR-CI [1.09,1.93]). A female partner being married or having lived with a male partner more than once has a much higher odds to be positive (increase of about 121%, OR-CI [1.61,3.02]) and the same holds for her male partner, though with a more limited increase of about 56% (OR-CI [1.13,2.16]). For most combinations of the covariate values the estimated probability to be HIV positive is higher for the female partner, in line with what is known. For the female partner the fitted probability to be HIV positive varies between 1.55% and 48.50% depending on the covariate pattern; for the male partner between 1.55% and 31.86% which is less but still substantial. The above discussed OR estimates and confidence intervals are for a given EA, and the estimated probabilities for a “central” EA with specific EA-effect equal to 0.</p>
<table-wrap id="pone.0172959.t003" position="float">
<object-id pub-id-type="doi">10.1371/journal.pone.0172959.t003</object-id>
<label>Table 3</label>
<caption>
<title>Parameters estimates and standard error estimates for the CE-PE model.</title>
</caption>
<alternatives>
<graphic id="pone.0172959.t003g" mimetype="image" position="float" xlink:href="info:doi/10.1371/journal.pone.0172959.t003" xlink:type="simple"/>
<table border="0" frame="box" rules="all">
<colgroup>
<col align="left" valign="middle"/>
<col align="left" valign="middle"/>
<col align="left" valign="middle"/>
<col align="left" valign="middle"/>
</colgroup>
<thead>
<tr>
<th align="left">Effect</th>
<th align="center">HIV Woman</th>
<th align="center">HIV Man</th>
<th align="center">
<monospace>SDM</monospace>
</th>
</tr>
<tr>
<th align="left"/>
<th align="center">Estimates(SE)</th>
<th align="center">Estimates(SE)</th>
<th align="center">Estimates(SE)</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" style="background-color:#fbfbfb"><bold>Intercept</bold></td>
<td align="center" colspan="2" style="background-color:#fbfbfb">-3.53(0.254)<xref ref-type="table-fn" rid="t003fn001">*</xref></td>
<td align="center" style="background-color:#fbfbfb">1.67(0.440)<xref ref-type="table-fn" rid="t003fn001">*</xref></td>
</tr>
<tr>
<td align="left"><bold>HIV prevalence</bold></td>
<td align="center" colspan="2"/>
<td align="center"/>
</tr>
<tr>
<td align="left">5–15%</td>
<td align="center" colspan="2">1.10(0.248)<xref ref-type="table-fn" rid="t003fn001">*</xref></td>
<td align="center">-0.56(0.452)</td>
</tr>
<tr>
<td align="left">&gt; 15%</td>
<td align="center" colspan="2">1.93(0.264)<xref ref-type="table-fn" rid="t003fn001">*</xref></td>
<td align="center">-1.08(0.455)<xref ref-type="table-fn" rid="t003fn001">*</xref></td>
</tr>
<tr>
<td align="left"><bold>Union number woman</bold></td>
<td align="center" colspan="2"/>
<td align="center"/>
</tr>
<tr>
<td align="left">More than once</td>
<td align="center">0.79(0.189)<xref ref-type="table-fn" rid="t003fn001">*</xref></td>
<td align="center">0.44(0.165)<xref ref-type="table-fn" rid="t003fn001">*</xref></td>
<td align="center">-0.57(0.232)<xref ref-type="table-fn" rid="t003fn001">*</xref></td>
</tr>
<tr>
<td align="left"><bold>Union number man</bold></td>
<td align="center" colspan="2"/>
<td align="center"/>
</tr>
<tr>
<td align="left">More than once</td>
<td align="center">0.11(0.150)</td>
<td align="center">0.37(0.144)<xref ref-type="table-fn" rid="t003fn001">*</xref></td>
<td align="center">-</td>
</tr>
<tr>
<td align="left"><bold>Condom use man</bold></td>
<td align="center" colspan="2"/>
<td align="center"/>
</tr>
<tr>
<td align="left">Not used</td>
<td align="center">0.64(0.251)<xref ref-type="table-fn" rid="t003fn001">*</xref></td>
<td align="center">0.03(0.248)</td>
<td align="center">-</td>
</tr>
<tr>
<td align="left"><bold>Wealth index</bold></td>
<td align="center" colspan="2"/>
<td align="center"/>
</tr>
<tr>
<td align="left">Poorer</td>
<td align="center" colspan="2">-0.62(0.172)<xref ref-type="table-fn" rid="t003fn001">*</xref></td>
<td align="center">-</td>
</tr>
<tr>
<td align="left">Middle</td>
<td align="center" colspan="2">-0.39(0.177)<xref ref-type="table-fn" rid="t003fn001">*</xref></td>
<td align="center">-</td>
</tr>
<tr>
<td align="left"><bold>Variance component</bold></td>
<td align="center" colspan="2"/>
<td align="center"/>
</tr>
<tr>
<td align="left"><italic>σ</italic><sup>2</sup></td>
<td align="center" colspan="2">0.46(0.115)<xref ref-type="table-fn" rid="t003fn002">†</xref></td>
<td align="center"/>
</tr>
</tbody>
</table>
</alternatives>
<table-wrap-foot>
<fn id="t003fn001">
<p>* Significant at 5% level (Wald test)</p>
</fn>
<fn id="t003fn002">
<p><sup>†</sup> Significant at 5% level, using <inline-formula id="pone.0172959.e023"><alternatives><graphic id="pone.0172959.e023g" mimetype="image" position="anchor" xlink:href="info:doi/10.1371/journal.pone.0172959.e023" xlink:type="simple"/><mml:math display="inline" id="M23"><mml:msubsup><mml:mi>χ</mml:mi> <mml:mrow><mml:mn>0</mml:mn> <mml:mo>,</mml:mo> <mml:mn>1</mml:mn></mml:mrow> <mml:mn>2</mml:mn></mml:msubsup></mml:math></alternatives></inline-formula> mixture</p>
</fn>
</table-wrap-foot>
</table-wrap>
<p>The null hypothesis <italic>H</italic><sub>0</sub>: <italic>σ</italic> = 0 is rejected at 5% level, using a <inline-formula id="pone.0172959.e024"><alternatives><graphic id="pone.0172959.e024g" mimetype="image" position="anchor" xlink:href="info:doi/10.1371/journal.pone.0172959.e024" xlink:type="simple"/><mml:math display="inline" id="M24"><mml:msubsup><mml:mi>χ</mml:mi> <mml:mrow><mml:mn>0</mml:mn> <mml:mo>,</mml:mo> <mml:mn>1</mml:mn></mml:mrow> <mml:mn>2</mml:mn></mml:msubsup></mml:math></alternatives></inline-formula> mixture, indicating a significant EA effect on the probability to be HIV positive (common to both members of the couple in that EA) and implying highly varying probabilities to be HIV positive across EAs. For an EA with an extreme negative EA-effect <inline-formula id="pone.0172959.e025"><alternatives><graphic id="pone.0172959.e025g" mimetype="image" position="anchor" xlink:href="info:doi/10.1371/journal.pone.0172959.e025" xlink:type="simple"/><mml:math display="inline" id="M25"><mml:mrow><mml:mo>-</mml:mo> <mml:mn>2</mml:mn> <mml:mover accent="true"><mml:mi>σ</mml:mi> <mml:mo>^</mml:mo></mml:mover> <mml:mo>=</mml:mo> <mml:mo>-</mml:mo> <mml:mn>1</mml:mn> <mml:mo>.</mml:mo> <mml:mn>356</mml:mn></mml:mrow></mml:math></alternatives></inline-formula> the estimated probability to be HIV positive for the female partner varies between 0.4% and 19.50% (depending on the covariate pattern); for the male partner between 0.4% and 10.75%. On the other hand, for an EA with an extreme positive EA-effect <inline-formula id="pone.0172959.e026"><alternatives><graphic id="pone.0172959.e026g" mimetype="image" position="anchor" xlink:href="info:doi/10.1371/journal.pone.0172959.e026" xlink:type="simple"/><mml:math display="inline" id="M26"><mml:mrow><mml:mn>2</mml:mn> <mml:mover accent="true"><mml:mi>σ</mml:mi> <mml:mo>^</mml:mo></mml:mover> <mml:mo>=</mml:mo> <mml:mn>1</mml:mn> <mml:mo>.</mml:mo> <mml:mn>356</mml:mn></mml:mrow></mml:math></alternatives></inline-formula> the estimated probability to be HIV positive for the female partner varies between 5.77% and 78.52% (depending on the covariate pattern); for the male partner between 5.77% and 64.48%.</p>
<p>The odds for a couple to be serodiscordant decreases with the HIV prevalence of the province. As compared to a province with less than 5% prevalence, the odds for a couple to be serodiscordant decreases with 66% when living in a province with more than 15% prevalence (OR-CI [0.14,0.84]). In other words, given that at least one of the two partners is positive, the probability that both are positive increases with the HIV prevalence of the province where the couple is living. The odds of serodiscordance also decreases when the woman has been married or living with a man more than once, with an estimated decrease of about 43% (OR-CI [0.36,0.89]). Given that one of the two partners is positive, the probability that the other is not varies between 50.50% and 84.16%. Note that, according to the CE-PE model, the probability to be serodiscordant does not vary across EAs, in contrast to the probabilities for men and women to be HIV positive.</p>
</sec>
<sec id="sec012" sec-type="conclusions">
<title>Discussion and conclusions</title>
<p>In this paper, a new serodiscordance measure SDM is defined as the conditional probability that both partners differ in their HIV status, given that one of them is positive. Together with the probabilities to be HIV positive for man and woman, the SDM parameter is embedded in a bivariate statistical model with fixed and random effects for covariates and design variables on each of the three parameters. A (final) model with significant fixed effects for the HIV prevalence of the province, the union number of woman and man, condom use and wealth index, and a significant random effect for the EAs was fitted using data from the 2009 INSIDA survey. An extension of the current analysis with more recent survey data and extending the model to investigate any time trend is a very interesting and relevant topic for further research.</p>
<p>In Zambia, a retrospectively study of 65 couples to estimate the likely origin of HIV infection, found that at least one quarter of cases of HIV infection in recently married men were acquired from extramarital partnerships, and for both men and women, less than one half of cases of HIV infection were acquired from their spouse/husband [<xref ref-type="bibr" rid="pone.0172959.ref004">4</xref>]. In addition, they report that many infections in married men, even in those with HIV-infected wives, could be acquired from outside the marriage [<xref ref-type="bibr" rid="pone.0172959.ref004">4</xref>]. Furthermore, a study conduct in South Africa to investigate who was infecting whom among migrant and non-migrant within concordance as well as discordance couples, found that non-migrant men were 10 times more likely to be infected from outside their regular relationships than inside [<xref ref-type="bibr" rid="pone.0172959.ref018">18</xref>].</p>
<p>Hence, it seems also relevant for Mozambican public health authorities to get more insights in this matter and to launch national prevention plans promoting regular HIV testing for couples (even prior to getting married or living together) and to support couples to prevent transmission once one partner has become infected. A similar recommendation was given by Kaiser <italic>at al</italic>. in Kenya, where they recommended that prevention interventions should begin early in relationships and include mutual knowledge of HIV status [<xref ref-type="bibr" rid="pone.0172959.ref005">5</xref>].</p>
<p>An interesting methodological topic for further research is the modification and application of the model to same-sex couples. Two approaches seem to be readily applicable, although further research is needed. A first possibility is the use of another unique identifier for both dimensions of the bivariate distribution. For instance, instead of <italic>y</italic><sub>1</sub> and <italic>y</italic><sub>2</sub> referring to the HIV status of woman and man respectively, one could use <italic>y</italic><sub>1</sub> to be the HIV status of the youngest of the two same-sex partners and <italic>y</italic><sub>2</sub> the status of the oldest of the two partners. Exactly the same model could be applied but now the two marginal success probabilities would represent the probabilities for the youngest partner and the oldest partner to be HIV positive. Any other unique identifier could be used, if available for all paired observations and if allowing interpretations of interest. A second approach would fully reflect the exchangeability. The HIV status of both partners could be put in any order to be <italic>y</italic><sub>1</sub> and <italic>y</italic><sub>2</sub>. This would only make sense if both models for both probabilities would be taken exactly the same with all effects in common (not only partial as in our model), such that all data contribute to the estimation of that same probability <italic>π</italic> = <italic>P</italic>(<italic>y</italic><sub>1</sub> = 1|<italic>x</italic>) = <italic>P</italic>(<italic>y</italic><sub>2</sub> = 1|<italic>x</italic>). This would result in a model with only two components; the first two components of <xref ref-type="disp-formula" rid="pone.0172959.e002">Eq (2)</xref> would be identical (<italic>h</italic>(<italic>π</italic>) = <italic>β</italic><sup><italic>T</italic></sup> <italic>x</italic>). As the SDM is symmetric, its definition <xref ref-type="disp-formula" rid="pone.0172959.e013">Eq (6)</xref> can remain unchanged, but it reduces to (as both marginal probabilities are identical):
<disp-formula id="pone.0172959.e027"><alternatives><graphic id="pone.0172959.e027g" mimetype="image" position="anchor" xlink:href="info:doi/10.1371/journal.pone.0172959.e027" xlink:type="simple"/><mml:math display="block" id="M27"><mml:mrow><mml:mtext mathvariant="sans-serif">SDM</mml:mtext> <mml:mo>=</mml:mo> <mml:mfrac><mml:mrow><mml:mn>2</mml:mn> <mml:mo>(</mml:mo> <mml:mi>π</mml:mi> <mml:mo>-</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mn>11</mml:mn></mml:msub> <mml:mo>)</mml:mo></mml:mrow> <mml:mrow><mml:mn>2</mml:mn> <mml:mi>π</mml:mi> <mml:mo>-</mml:mo> <mml:msub><mml:mi>π</mml:mi> <mml:mn>11</mml:mn></mml:msub></mml:mrow></mml:mfrac> <mml:mo>.</mml:mo></mml:mrow></mml:math></alternatives></disp-formula>
Similar modifications are to be applied to the joint probabilities <xref ref-type="disp-formula" rid="pone.0172959.e006">Eq (3)</xref>. What the potential as well as limitations of such an approach are, needs still to be investigated.</p>
</sec>
<sec id="sec013">
<title>Appendix: SAS code for the final CE-PE model (see results in <xref ref-type="table" rid="pone.0172959.t003">Table 3</xref>)</title>
<p><monospace>proc nlmixed data = couples qpoints = 100;</monospace></p>
<p><monospace>/* <xref ref-type="disp-formula" rid="pone.0172959.e017">Eq (8)</xref>: probability model for P(y1 = 1) on logit scale (with random effect a) */</monospace></p>
<p><monospace>eta1 = b0_1 + b1_1*prev_d2 + b2_1*prev_d3 + b3_1*N_union_dW2 + b7_1*N_union_dM2 + b4_1*po + b5_1*mid + b8_1*Condom_dM1 + a;</monospace></p>
<p><monospace>/* probability model for p1_1 = P(y1 = 1) on probability scale */</monospace></p>
<p><monospace>p1_1 = 1 / (1 + exp(-eta1));</monospace></p>
<p><monospace>/* <xref ref-type="disp-formula" rid="pone.0172959.e017">Eq (8)</xref>: probability model for P(y2 = 1) on logit scale (with random effect a) */</monospace></p>
<p><monospace>eta2 = b0_1 + b1_1*prev_d2 + b2_1*prev_d3 + b3_2*N_union_dW2 + b7_2*N_Union_dM2 + b4_1*po + b5_1*mid + b8_2*Condom_dM1 + a;</monospace></p>
<p><monospace>/* probability model for p1_2 = P(y2 = 1) on probability scale */</monospace></p>
<p><monospace>p1_2 = 1 / (1 + exp(-eta2));</monospace></p>
<p><monospace>/* <xref ref-type="disp-formula" rid="pone.0172959.e017">Eq (8)</xref>: probability model for SDM on logit scale (without random effect) */</monospace></p>
<p><monospace>eta3 = b0_3 + b1_3*prev_d2 + b2_3*prev_d3 + b3_3*N_union_dW2;</monospace></p>
<p><monospace>SDM = 1/(1 + exp(-eta3));</monospace></p>
<p><monospace>/* probability model for SDM on probability scale */</monospace></p>
<p><monospace>/* Eqs (<xref ref-type="disp-formula" rid="pone.0172959.e006">3</xref>)&amp;(<xref ref-type="disp-formula" rid="pone.0172959.e015">7</xref>) joint probabilities as function of p1_1 = P(y1 = 1), p1_2 = P(y2 = 1), SDM */</monospace></p>
<p><monospace>p11 = ((1-SDM)*(p1_1 + p1_2))/(2-SDM);</monospace></p>
<p><monospace>p10 = p1_1 − p11;</monospace></p>
<p><monospace>p01 = p1_2 − p11;</monospace></p>
<p><monospace>p00 = 1 − p1_1 − p1_2 + p11;</monospace></p>
<p><monospace>/* weighted log-likelihood for multinomial distribution (weights of survey design) */</monospace></p>
<p><monospace>ll = HIV_W*HIV_M*log(p11) + HIV_W*(1-HIV_M)*log(p10) + (1-HIV_W)*HIV_M*log(p01) + (1-HIV_W)*(1-HIV_M)*log(p00);</monospace></p>
<p><monospace>ll = ll*weight;</monospace></p>
<p><monospace>model ll ~ general(ll);</monospace></p>
<p><monospace>random a ~ normal(0,sigma2_a) subject = EA;</monospace></p>
<p><monospace>run;</monospace></p>
</sec>
</body>
<back>
<ack>
<p>This study was only possible thanks to the financial support of the Flemish Interuniversity Council (VLIR-UOS) in collaboration with Eduardo Mondlane University (UEM) through the DESAFIO Program. The authors would also like to acknowledge the support given by the Mozambican Health Ministry (MISAU) and Demography Health Survey (DHS) Program for providing the INSIDA survey data. The authors would also like to thank the editors and the reviewers for their interesting and helpful comments and suggestions.</p>
</ack>
<ref-list>
<title>References</title>
<ref id="pone.0172959.ref001">
<label>1</label>
<mixed-citation publication-type="book" xlink:type="simple">
<name name-style="western"><surname>Fishel</surname> <given-names>JD</given-names></name>, <name name-style="western"><surname>Bradley</surname> <given-names>SEK</given-names></name>, <name name-style="western"><surname>Young</surname> <given-names>PW</given-names></name>, <name name-style="western"><surname>Mbofana</surname> <given-names>F</given-names></name>, <name name-style="western"><surname>Botão</surname> <given-names>C</given-names></name>. <chapter-title>HIV among Couples in Mozambique: HIV Status, Knowledge of Status, and Factors Associated with HIV Serodiscordance</chapter-title>. <source>Further Analysis of the 2009 Inquérito Nacional de Prevalência, Riscos Comportamentais e Informação sobre o HIV e SIDA em Moçambique 2009</source>. <publisher-loc>Calverton, Maryland, USA</publisher-loc>. <publisher-name>ICF International</publisher-name>. <year>2011</year></mixed-citation>
</ref>
<ref id="pone.0172959.ref002">
<label>2</label>
<mixed-citation publication-type="journal" xlink:type="simple">
<name name-style="western"><surname>Ewayo</surname> <given-names>O</given-names></name>, <name name-style="western"><surname>de Walque</surname> <given-names>D</given-names></name>, <name name-style="western"><surname>Ford</surname> <given-names>N</given-names></name>, <name name-style="western"><surname>Gakii</surname> <given-names>G</given-names></name>, <name name-style="western"><surname>Lester</surname> <given-names>RT</given-names></name>, <name name-style="western"><surname>Mills</surname> <given-names>EJ</given-names></name>. <article-title>HIV status in discordant couples in sub-Saharan Africa: a systematic review and meta-analysis</article-title>. <source>The Lancet Infectious Diseases</source>. <year>2010</year>.<volume>10</volume>:<fpage>770</fpage>–<lpage>777</lpage> <comment>doi: <ext-link ext-link-type="uri" xlink:href="http://dx.doi.org/10.1016/S1473-3099(10)70189-4" xlink:type="simple">10.1016/S1473-3099(10)70189-4</ext-link></comment></mixed-citation>
</ref>
<ref id="pone.0172959.ref003">
<label>3</label>
<mixed-citation publication-type="journal" xlink:type="simple">
<name name-style="western"><surname>Dunkle</surname> <given-names>KL</given-names></name>, <name name-style="western"><surname>Stephenson</surname> <given-names>R</given-names></name>, <name name-style="western"><surname>Karita</surname> <given-names>E</given-names></name>, <name name-style="western"><surname>Chomba</surname> <given-names>E</given-names></name>, <name name-style="western"><surname>Kayitenkore</surname> <given-names>K</given-names></name>, <name name-style="western"><surname>Vwalika</surname> <given-names>C</given-names></name>, <etal>et al</etal>. <article-title>New heterosexually transmitted HIV infections in married or cohabiting couples in urban Zambia and Rwanda: an analysis of survey and clinical data</article-title>. <source>The Lancet Infectious Diseases</source>. <year>2008</year>.<volume>371</volume>:<fpage>2183</fpage>–<lpage>91</lpage></mixed-citation>
</ref>
<ref id="pone.0172959.ref004">
<label>4</label>
<mixed-citation publication-type="journal" xlink:type="simple">
<name name-style="western"><surname>Glynn</surname> <given-names>JR</given-names></name>, <name name-style="western"><surname>Carael</surname> <given-names>M</given-names></name>, <name name-style="western"><surname>Buve</surname> <given-names>A</given-names></name>, <name name-style="western"><surname>Musonda</surname> <given-names>R M</given-names></name>, <name name-style="western"><surname>Kahindo</surname> <given-names>M</given-names></name>. <article-title>HIV risk in relation to marriage in areas with high prevalence of HIV infection</article-title>. <source>J Acquir Immune Defic Syndr</source>. <year>2003</year>.<volume>33</volume>:<fpage>526</fpage>–<lpage>535</lpage> <object-id pub-id-type="pmid">12869843</object-id></mixed-citation>
</ref>
<ref id="pone.0172959.ref005">
<label>5</label>
<mixed-citation publication-type="journal" xlink:type="simple">
<name name-style="western"><surname>Kaiser</surname> <given-names>R</given-names></name>, <name name-style="western"><surname>Bunnell</surname> <given-names>R</given-names></name>, <name name-style="western"><surname>Hightower</surname> <given-names>A</given-names></name>, <name name-style="western"><surname>Kim</surname> <given-names>A</given-names></name>, <name name-style="western"><surname>Cherutich</surname> <given-names>P</given-names></name>, <name name-style="western"><surname>Mwangi</surname> <given-names>M</given-names></name>, <etal>et al</etal>. <article-title>Factors Associated with HIV Infection in Married or Cohabitating Couples in Kenya: Results from a Nationally Representative Study, for the KAIS Study Group</article-title>. <source>PLoS ONE</source>. <year>2011</year>.<volume>6</volume>(<issue>3</issue>): <fpage>e17842</fpage>. <comment>doi: <ext-link ext-link-type="uri" xlink:href="http://dx.doi.org/10.1371/journal.pone.0017842" xlink:type="simple">10.1371/journal.pone.0017842</ext-link></comment> <object-id pub-id-type="pmid">21423615</object-id></mixed-citation>
</ref>
<ref id="pone.0172959.ref006">
<label>6</label>
<mixed-citation publication-type="journal" xlink:type="simple">
<name name-style="western"><surname>Faes</surname> <given-names>C</given-names></name>, <name name-style="western"><surname>Geys</surname> <given-names>H</given-names></name>, <name name-style="western"><surname>Molenberghs</surname> <given-names>G</given-names></name>, <name name-style="western"><surname>Aerts</surname> <given-names>M</given-names></name>, <name name-style="western"><surname>Cadarso-Suaresz</surname> <given-names>C</given-names></name>, <name name-style="western"><surname>Acuna</surname> <given-names>C</given-names></name>, <etal>et al</etal>. <article-title>A Flexible Method to Measure Synchrony in Neuronal Firing</article-title>. <source>Journal Of The American Statistical Association</source>. <year>2008</year>.<volume>103</volume>:<fpage>149</fpage>–<lpage>161</lpage>. <comment>doi: <ext-link ext-link-type="uri" xlink:href="http://dx.doi.org/10.1198/016214507000000419" xlink:type="simple">10.1198/016214507000000419</ext-link></comment></mixed-citation>
</ref>
<ref id="pone.0172959.ref007">
<label>7</label>
<mixed-citation publication-type="journal" xlink:type="simple">
<name name-style="western"><surname>Dale</surname> <given-names>JR</given-names></name>. <article-title>Global cross-ratio models for bivariate, discrete, ordered responses</article-title>. <source>Biometrics</source>. <year>1986</year>.<volume>42</volume>:<fpage>909</fpage>–<lpage>917</lpage> <comment>doi: <ext-link ext-link-type="uri" xlink:href="http://dx.doi.org/10.2307/2530704" xlink:type="simple">10.2307/2530704</ext-link></comment> <object-id pub-id-type="pmid">3814731</object-id></mixed-citation>
</ref>
<ref id="pone.0172959.ref008">
<label>8</label>
<mixed-citation publication-type="journal" xlink:type="simple">
<name name-style="western"><surname>Liang</surname> <given-names>KY</given-names></name>, <name name-style="western"><surname>Zeger</surname> <given-names>SL</given-names></name>. <article-title>Longitudinal data analysis using generalized linear models</article-title>. <source>Biometrika</source>. <year>1986</year>.<volume>73</volume>:<fpage>13</fpage>–<lpage>22</lpage>. <comment>doi: <ext-link ext-link-type="uri" xlink:href="http://dx.doi.org/10.1093/biomet/73.1.13" xlink:type="simple">10.1093/biomet/73.1.13</ext-link></comment></mixed-citation>
</ref>
<ref id="pone.0172959.ref009">
<label>9</label>
<mixed-citation publication-type="journal" xlink:type="simple">
<name name-style="western"><surname>Breslow</surname> <given-names>N</given-names></name>, <name name-style="western"><surname>Clayton</surname> <given-names>DG</given-names></name>. <article-title>Approximate inference in generalized linear mixed models</article-title>. <source>Journal American Statistical Association</source>. <year>1993</year>.<volume>88</volume>:<fpage>9</fpage>–<lpage>25</lpage>. <comment>doi: <ext-link ext-link-type="uri" xlink:href="http://dx.doi.org/10.2307/2290687" xlink:type="simple">10.2307/2290687</ext-link></comment></mixed-citation>
</ref>
<ref id="pone.0172959.ref010">
<label>10</label>
<mixed-citation publication-type="other" xlink:type="simple">Instituto Nacional de Saúde (INS), Instituto Nacional de Estatística (INE). Inquérito Nacional de Prevalência, Riscos Comportamentais e Informação sobre o HIV e SIDA em Moçambique de 2009. 2010:333.</mixed-citation>
</ref>
<ref id="pone.0172959.ref011">
<label>11</label>
<mixed-citation publication-type="book" xlink:type="simple">
<name name-style="western"><surname>Agresti</surname> <given-names>A</given-names></name>. <source>Categorical Data Analysis</source>. <edition>Third Edition</edition>. <publisher-loc>New York</publisher-loc>: <publisher-name>John Wiley and Sons Inc</publisher-name>; <year>2013</year></mixed-citation>
</ref>
<ref id="pone.0172959.ref012">
<label>12</label>
<mixed-citation publication-type="journal" xlink:type="simple">
<name name-style="western"><surname>Carey</surname> <given-names>V</given-names></name>, <name name-style="western"><surname>Zeger</surname> <given-names>SL</given-names></name>, <name name-style="western"><surname>Diggle</surname> <given-names>P</given-names></name>. <article-title>Modelling multivariate binary data with alternating logistic regressions</article-title>. <source>Biometrika</source>. <year>1993</year>.<volume>80</volume>:<fpage>517</fpage>–<lpage>526</lpage>. <comment>doi: <ext-link ext-link-type="uri" xlink:href="http://dx.doi.org/10.1093/biomet/80.3.517" xlink:type="simple">10.1093/biomet/80.3.517</ext-link></comment></mixed-citation>
</ref>
<ref id="pone.0172959.ref013">
<label>13</label>
<mixed-citation publication-type="journal" xlink:type="simple">
<name name-style="western"><surname>Lovasi</surname> <given-names>GS</given-names></name>, <name name-style="western"><surname>Underhill</surname> <given-names>LJ</given-names></name>, <name name-style="western"><surname>Jack</surname> <given-names>D</given-names></name>, <name name-style="western"><surname>Richards</surname> <given-names>C</given-names></name>, <name name-style="western"><surname>Weiss</surname> <given-names>C</given-names></name>, <name name-style="western"><surname>Rundle</surname> <given-names>A</given-names></name>. <article-title>At Odds: Concerns Raised by Using Odds Ratios for Continuous or Common Dichotomous Outcomes in Research on Physical Activity and Obesity</article-title>. <source>The open epidemiology journal</source>. <year>2012</year>;<volume>5</volume>:<fpage>13</fpage>–<lpage>17</lpage>. <comment>doi: <ext-link ext-link-type="uri" xlink:href="http://dx.doi.org/10.2174/1874297101205010013" xlink:type="simple">10.2174/1874297101205010013</ext-link></comment> <object-id pub-id-type="pmid">23002407</object-id></mixed-citation>
</ref>
<ref id="pone.0172959.ref014">
<label>14</label>
<mixed-citation publication-type="journal" xlink:type="simple">
<name name-style="western"><surname>Thomas</surname> <given-names>SL</given-names></name>, <name name-style="western"><surname>Heck</surname> <given-names>RH</given-names></name>. <article-title>Analysis of Large-Scale Secondary Data in Higher Education Research: Potential Perils Associated With Complex Sampling Designs</article-title>. <source>Research in Higher Education</source>. <year>2001</year>.<volume>42</volume>:<fpage>517</fpage>–<lpage>540</lpage>.</mixed-citation>
</ref>
<ref id="pone.0172959.ref015">
<label>15</label>
<mixed-citation publication-type="other" xlink:type="simple">Akaike H. Information theory and an extension of the maximum likelihood princi- ple. In Petrov BN and Csaki F. editors. 2nd International Symposium on Information Theory: 267–281. Akademiai K, Budapest (Reproduced in Breakthroughs in Statistics, Volume 1 (eds. S. Kotz and N. L. Johnson), Springer Verlag, New York (1992)). 1973</mixed-citation>
</ref>
<ref id="pone.0172959.ref016">
<label>16</label>
<mixed-citation publication-type="journal" xlink:type="simple">
<name name-style="western"><surname>Dokubo</surname> <given-names>EK</given-names></name>, <name name-style="western"><surname>Shiraishi</surname> <given-names>RW</given-names></name>, <name name-style="western"><surname>Young</surname> <given-names>PW</given-names></name>, <name name-style="western"><surname>Neal</surname> <given-names>JJ</given-names></name>, <name name-style="western"><surname>Aberle-Grasse</surname> <given-names>J</given-names></name>, <etal>et al</etal>. <article-title>Awareness of HIV Status, Prevention Knowledge and Condom Use among People Living with HIV in Mozambique</article-title>. <source>PLoS ONE</source> <year>2014</year>.<volume>9</volume>(<issue>9</issue>): <fpage>e106760</fpage>. <comment>doi: <ext-link ext-link-type="uri" xlink:href="http://dx.doi.org/10.1371/journal.pone.0106760" xlink:type="simple">10.1371/journal.pone.0106760</ext-link></comment> <object-id pub-id-type="pmid">25222010</object-id></mixed-citation>
</ref>
<ref id="pone.0172959.ref017">
<label>17</label>
<mixed-citation publication-type="journal" xlink:type="simple">
<name name-style="western"><surname>Vera Cruz</surname> <given-names>G</given-names></name>, and <name name-style="western"><surname>Maússe</surname> <given-names>L</given-names></name>. <article-title>Multiple and Concurrent Sexual Partnerships among Mozambican Women from High Socio-Economic Status and with High Education Degrees: Involvement Motives</article-title>. <source>Psychology</source>, <year>2014</year>.<volume>5</volume>, <fpage>1260</fpage>–<lpage>1267</lpage>. <comment>doi: <ext-link ext-link-type="uri" xlink:href="http://dx.doi.org/10.4236/psych.2014.510138" xlink:type="simple">10.4236/psych.2014.510138</ext-link></comment></mixed-citation>
</ref>
<ref id="pone.0172959.ref018">
<label>18</label>
<mixed-citation publication-type="journal" xlink:type="simple">
<name name-style="western"><surname>Lurie</surname> <given-names>MN</given-names></name>, <name name-style="western"><surname>Williams</surname> <given-names>BG</given-names></name>, <name name-style="western"><surname>Zuma</surname> <given-names>K</given-names></name>, <name name-style="western"><surname>Mkaya-Mwamburi</surname> <given-names>D</given-names></name>, <name name-style="western"><surname>Garnett</surname> <given-names>GP</given-names></name>, <name name-style="western"><surname>Sweat</surname> <given-names>MD</given-names></name>, <etal>et al</etal>. <article-title>Who infects whom? HIV-1 concordance and discordance among migrant and non-migrant couples in South Africa</article-title>. <source>AIDS</source>. <year>2003</year>.<volume>17</volume>:<fpage>2245</fpage>–<lpage>2252</lpage>. <object-id pub-id-type="pmid">14523282</object-id></mixed-citation>
</ref>
</ref-list>
</back>
</article>