<?xml version="1.0" encoding="UTF-8"?><!DOCTYPE article SYSTEM "http://jats.nlm.nih.gov/archiving/1.2/JATS-archivearticle1.dtd"><article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" dtd-version="1.2" article-type="research-article" xml:lang="en"><?properties open_access?><front><journal-meta><journal-id journal-id-type="publisher-id">41467</journal-id><journal-id journal-id-type="doi">10.1038/41467.2041-1723</journal-id><journal-title-group><journal-title>Nature Communications</journal-title><abbrev-journal-title abbrev-type="publisher">Nat Commun</abbrev-journal-title></journal-title-group><issn pub-type="epub">2041-1723</issn><publisher><publisher-name>Nature Publishing Group UK</publisher-name><publisher-loc>London</publisher-loc></publisher></journal-meta><article-meta><article-id pub-id-type="publisher-id">s41467-021-21593-7</article-id><article-id pub-id-type="manuscript">21593</article-id><article-id pub-id-type="doi">10.1038/s41467-021-21593-7</article-id><article-categories><subj-group subj-group-type="heading"><subject>Article</subject></subj-group><subj-group subj-group-type="SubjectPath"><subject>/631/208/191/2018</subject></subj-group><subj-group subj-group-type="SubjectPath"><subject>/631/208/200</subject></subj-group><subj-group subj-group-type="SubjectPath"><subject>/631/208/205/2138</subject></subj-group><subj-group subj-group-type="SubjectPath"><subject>/631/208/212/2019</subject></subj-group><subj-group subj-group-type="TechniquePath"><subject>/45/91</subject></subj-group><subj-group subj-group-type="TechniquePath"><subject>/45</subject></subj-group><subj-group subj-group-type="NatureArticleTypeID"><subject>article</subject></subj-group></article-categories><title-group><article-title xml:lang="en">A molecular quantitative trait locus map for osteoarthritis</article-title></title-group><contrib-group><contrib contrib-type="author" id="Au1"><contrib-id contrib-id-type="orcid">http://orcid.org/0000-0002-0585-2312</contrib-id><name><surname>Steinberg</surname><given-names>Julia</given-names></name><xref ref-type="aff" rid="Aff1">1</xref><xref ref-type="aff" rid="Aff2">2</xref><xref ref-type="aff" rid="Aff3">3</xref><xref ref-type="aff" rid="Aff4">4</xref></contrib><contrib contrib-type="author" id="Au2"><name><surname>Southam</surname><given-names>Lorraine</given-names></name><xref ref-type="aff" rid="Aff1">1</xref><xref ref-type="aff" rid="Aff3">3</xref></contrib><contrib contrib-type="author" id="Au3"><name><surname>Roumeliotis</surname><given-names>Theodoros I.</given-names></name><xref ref-type="aff" rid="Aff3">3</xref><xref ref-type="aff" rid="Aff5">5</xref></contrib><contrib contrib-type="author" id="Au4"><contrib-id contrib-id-type="orcid">http://orcid.org/0000-0003-2152-2257</contrib-id><name><surname>Clark</surname><given-names>Matthew J.</given-names></name><xref ref-type="aff" rid="Aff6">6</xref></contrib><contrib contrib-type="author" id="Au5"><name><surname>Jayasuriya</surname><given-names>Raveen L.</given-names></name><xref ref-type="aff" rid="Aff6">6</xref></contrib><contrib contrib-type="author" id="Au6"><name><surname>Swift</surname><given-names>Diane</given-names></name><xref ref-type="aff" rid="Aff6">6</xref></contrib><contrib contrib-type="author" id="Au7"><contrib-id contrib-id-type="orcid">http://orcid.org/0000-0001-9909-6409</contrib-id><name><surname>Shah</surname><given-names>Karan M.</given-names></name><xref ref-type="aff" rid="Aff6">6</xref></contrib><contrib contrib-type="author" id="Au8"><contrib-id contrib-id-type="orcid">http://orcid.org/0000-0002-5209-7508</contrib-id><name><surname>Butterfield</surname><given-names>Natalie C.</given-names></name><xref ref-type="aff" rid="Aff7">7</xref></contrib><contrib contrib-type="author" id="Au9"><name><surname>Brooks</surname><given-names>Roger A.</given-names></name><xref ref-type="aff" rid="Aff8">8</xref></contrib><contrib contrib-type="author" id="Au10"><name><surname>McCaskie</surname><given-names>Andrew W.</given-names></name><xref ref-type="aff" rid="Aff8">8</xref></contrib><contrib contrib-type="author" id="Au11"><contrib-id contrib-id-type="orcid">http://orcid.org/0000-0003-0817-0082</contrib-id><name><surname>Bassett</surname><given-names>J. H. Duncan</given-names></name><xref ref-type="aff" rid="Aff7">7</xref></contrib><contrib contrib-type="author" id="Au12"><contrib-id contrib-id-type="orcid">http://orcid.org/0000-0002-8555-8219</contrib-id><name><surname>Williams</surname><given-names>Graham R.</given-names></name><xref ref-type="aff" rid="Aff7">7</xref></contrib><contrib contrib-type="author" id="Au13"><contrib-id contrib-id-type="orcid">http://orcid.org/0000-0003-0881-5477</contrib-id><name><surname>Choudhary</surname><given-names>Jyoti S.</given-names></name><xref ref-type="aff" rid="Aff3">3</xref><xref ref-type="aff" rid="Aff5">5</xref></contrib><contrib contrib-type="author" corresp="yes" id="Au14"><contrib-id contrib-id-type="orcid">http://orcid.org/0000-0001-5577-3674</contrib-id><name><surname>Wilkinson</surname><given-names>J. Mark</given-names></name><xref ref-type="aff" rid="Aff6">6</xref><xref ref-type="aff" rid="Aff9">9</xref><xref ref-type="corresp" rid="IDs41467021215937_cor14">q</xref></contrib><contrib contrib-type="author" corresp="yes" id="Au15"><contrib-id contrib-id-type="orcid">http://orcid.org/0000-0003-4238-659X</contrib-id><name><surname>Zeggini</surname><given-names>Eleftheria</given-names></name><xref ref-type="aff" rid="Aff1">1</xref><xref ref-type="aff" rid="Aff3">3</xref><xref ref-type="aff" rid="Aff10">10</xref><xref ref-type="corresp" rid="IDs41467021215937_cor15">r</xref></contrib><aff id="Aff1"><label>1</label><institution-wrap><institution-id institution-id-type="GRID">grid.4567.0</institution-id><institution-id institution-id-type="ISNI">0000 0004 0483 2525</institution-id><institution content-type="org-division">Institute of Translational Genomics</institution><institution content-type="org-name">Helmholtz Zentrum München – German Research Center for Environmental Health</institution></institution-wrap><addr-line content-type="city">Neuherberg</addr-line><country country="DE">Germany</country></aff><aff id="Aff2"><label>2</label><institution-wrap><institution-id institution-id-type="GRID">grid.420082.c</institution-id><institution-id institution-id-type="ISNI">0000 0001 2166 6280</institution-id><institution content-type="org-division">Cancer Research Division</institution><institution content-type="org-name">Cancer Council NSW</institution></institution-wrap><addr-line content-type="city">Sydney</addr-line><addr-line content-type="state">NSW</addr-line><country country="AU">Australia</country></aff><aff id="Aff3"><label>3</label><institution-wrap><institution-id institution-id-type="GRID">grid.10306.34</institution-id><institution-id institution-id-type="ISNI">0000 0004 0606 5382</institution-id><institution content-type="org-name">Wellcome Sanger Institute</institution></institution-wrap><addr-line content-type="city">Hinxton</addr-line><country country="GB">United Kingdom</country></aff><aff id="Aff4"><label>4</label><institution-wrap><institution-id institution-id-type="GRID">grid.1013.3</institution-id><institution-id institution-id-type="ISNI">0000 0004 1936 834X</institution-id><institution content-type="org-division">School of Public Health</institution><institution content-type="org-name">The University of Sydney</institution></institution-wrap><addr-line content-type="city">Sydney</addr-line><addr-line content-type="state">NSW</addr-line><country country="AU">Australia</country></aff><aff id="Aff5"><label>5</label><institution-wrap><institution-id institution-id-type="GRID">grid.18886.3f</institution-id><institution-id institution-id-type="ISNI">0000 0001 1271 4623</institution-id><institution content-type="org-name">The Institute of Cancer Research</institution></institution-wrap><addr-line content-type="city">London</addr-line><country country="GB">United Kingdom</country></aff><aff id="Aff6"><label>6</label><institution-wrap><institution-id institution-id-type="GRID">grid.11835.3e</institution-id><institution-id institution-id-type="ISNI">0000 0004 1936 9262</institution-id><institution content-type="org-division">Department of Oncology and Metabolism</institution><institution content-type="org-name">University of Sheffield</institution></institution-wrap><addr-line content-type="city">Sheffield</addr-line><country country="GB">United Kingdom</country></aff><aff id="Aff7"><label>7</label><institution-wrap><institution-id institution-id-type="GRID">grid.7445.2</institution-id><institution-id institution-id-type="ISNI">0000 0001 2113 8111</institution-id><institution content-type="org-division">Molecular Endocrinology Laboratory, Department of Metabolism, Digestion and Reproduction</institution><institution content-type="org-name">Imperial College London</institution></institution-wrap><addr-line content-type="city">London</addr-line><country country="GB">United Kingdom</country></aff><aff id="Aff8"><label>8</label><institution-wrap><institution-id institution-id-type="GRID">grid.5335.0</institution-id><institution-id institution-id-type="ISNI">0000000121885934</institution-id><institution content-type="org-division">Division of Trauma &amp; Orthopaedic Surgery, Department of Surgery</institution><institution content-type="org-name">University of Cambridge</institution></institution-wrap><addr-line content-type="city">Cambridge</addr-line><country country="GB">United Kingdom</country></aff><aff id="Aff9"><label>9</label><institution-wrap><institution-id institution-id-type="GRID">grid.11835.3e</institution-id><institution-id institution-id-type="ISNI">0000 0004 1936 9262</institution-id><institution content-type="org-division">Centre for Integrated Research into Musculoskeletal Ageing and Sheffield Healthy Lifespan Institute</institution><institution content-type="org-name">University of Sheffield</institution></institution-wrap><addr-line content-type="city">Sheffield</addr-line><country country="GB">United Kingdom</country></aff><aff id="Aff10"><label>10</label><institution-wrap><institution-id institution-id-type="GRID">grid.15474.33</institution-id><institution-id institution-id-type="ISNI">0000 0004 0477 2438</institution-id><institution content-type="org-name">TUM School of Medicine, Technical University of Munich and Klinikum Rechts der Isar</institution></institution-wrap><addr-line content-type="city">Munich</addr-line><country country="DE">Germany</country></aff></contrib-group><author-notes><corresp id="IDs41467021215937_cor14"><label>q</label><email>j.m.wilkinson@sheffield.ac.uk</email></corresp><corresp id="IDs41467021215937_cor15"><label>r</label><email>eleftheria.zeggini@helmholtz-muenchen.de</email></corresp></author-notes><pub-date date-type="pub" publication-format="electronic"><day>26</day><month>2</month><year>2021</year></pub-date><pub-date date-type="collection" publication-format="electronic"><month>12</month><year>2021</year></pub-date><volume>12</volume><issue seq="21216">1</issue><elocation-id>1309</elocation-id><history><date date-type="registration"><day>4</day><month>2</month><year>2021</year></date><date date-type="received"><day>9</day><month>7</month><year>2020</year></date><date date-type="accepted"><day>3</day><month>2</month><year>2021</year></date><date date-type="online"><day>26</day><month>2</month><year>2021</year></date></history><pub-history><event event-type="Update"><event-desc>Open Access funding information has been added to this article.</event-desc><date><day>12</day><month>4</month><year>2021</year></date></event></pub-history><permissions><copyright-statement>© The Author(s) 2021. corrected publication 2021</copyright-statement><copyright-year>2021</copyright-year><license license-type="open-access" xlink:href="http://creativecommons.org/licenses/by/4.0/"><license-p><bold>Open Access</bold> This article is licensed under a Creative Commons Attribution 4.0 International License, which permits use, sharing, adaptation, distribution and reproduction in any medium or format, as long as you give appropriate credit to the original author(s) and the source, provide a link to the Creative Commons license, and indicate if changes were made. The images or other third party material in this article are included in the article’s Creative Commons license, unless indicated otherwise in a credit line to the material. If material is not included in the article’s Creative Commons license and your intended use is not permitted by statutory regulation or exceeds the permitted use, you will need to obtain permission directly from the copyright holder. To view a copy of this license, visit <ext-link xlink:href="http://creativecommons.org/licenses/by/4.0/" ext-link-type="url">http://creativecommons.org/licenses/by/4.0/</ext-link>.</license-p></license></permissions><abstract id="Abs1" xml:lang="en"><title>Abstract</title><p id="Par2">Osteoarthritis causes pain and functional disability for over 500 million people worldwide. To develop disease-stratifying tools and modifying therapies, we need a better understanding of the molecular basis of the disease in relevant tissue and cell types. Here, we study primary cartilage and synovium from 115 patients with osteoarthritis to construct a deep molecular signature map of the disease. By integrating genetics with transcriptomics and proteomics, we discover molecular trait loci in each tissue type and omics level, identify likely effector genes for osteoarthritis-associated genetic signals and highlight high-value targets for drug development and repurposing. These findings provide insights into disease aetiopathology, and offer translational opportunities in response to the global clinical challenge of osteoarthritis.</p></abstract><abstract id="Abs2" xml:lang="en"><p id="Par3">Understanding the molecular effects of disease variants in relevant tissues is essential to understanding and treating disease. Here, the authors discover expression and protein quantitative trait loci in cartilage and synovium from 115 osteoarthritis patients to pinpoint genes of action and potential drug treatments.</p></abstract><funding-group><award-group><funding-source><institution-wrap><institution>Medical Research Council Centre for Integrated Research into Musculoskeletal Ageing grant (148985)</institution></institution-wrap></funding-source></award-group><award-group><funding-source><institution-wrap><institution>Versus Arthritis; Tissue Engineering and Regenerative Therapies Centre (21156)</institution></institution-wrap></funding-source></award-group><award-group><funding-source><institution-wrap><institution>Wellcome Trust (Wellcome)</institution><institution-id institution-id-type="doi" vocab="open-funder-registry">https://doi.org/10.13039/100004440</institution-id></institution-wrap></funding-source><award-id award-type="FundRef grant">101123</award-id><award-id award-type="FundRef grant">110140 and 110141</award-id><principal-award-recipient><name><surname>Williams</surname><given-names>Graham R.</given-names></name></principal-award-recipient><principal-award-recipient><name><surname>Williams</surname><given-names>Graham R.</given-names></name></principal-award-recipient></award-group><award-group><funding-source><institution-wrap><institution>EC | Horizon 2020 Framework Programme (EU Framework Programme for Research and Innovation H2020)</institution><institution-id institution-id-type="doi" vocab="open-funder-registry">https://doi.org/10.13039/100010661</institution-id></institution-wrap></funding-source><award-id award-type="FundRef grant">666869, THYRAGE</award-id><principal-award-recipient><name><surname>Williams</surname><given-names>Graham R.</given-names></name></principal-award-recipient></award-group></funding-group><custom-meta-group><custom-meta><meta-name>publisher-imprint-name</meta-name><meta-value>Nature Portfolio</meta-value></custom-meta><custom-meta><meta-name>volume-issue-count</meta-name><meta-value>1</meta-value></custom-meta><custom-meta><meta-name>issue-article-count</meta-name><meta-value>2307</meta-value></custom-meta><custom-meta><meta-name>issue-toc-levels</meta-name><meta-value>0</meta-value></custom-meta><custom-meta><meta-name>issue-pricelist-year</meta-name><meta-value>2021</meta-value></custom-meta><custom-meta><meta-name>issue-copyright-holder</meta-name><meta-value>The Author(s)</meta-value></custom-meta><custom-meta><meta-name>issue-copyright-year</meta-name><meta-value>2021</meta-value></custom-meta><custom-meta><meta-name>article-contains-esm</meta-name><meta-value>Yes</meta-value></custom-meta><custom-meta><meta-name>article-numbering-style</meta-name><meta-value>Unnumbered</meta-value></custom-meta><custom-meta><meta-name>article-registration-date-year</meta-name><meta-value>2021</meta-value></custom-meta><custom-meta><meta-name>article-registration-date-month</meta-name><meta-value>2</meta-value></custom-meta><custom-meta><meta-name>article-registration-date-day</meta-name><meta-value>4</meta-value></custom-meta><custom-meta><meta-name>article-toc-levels</meta-name><meta-value>0</meta-value></custom-meta><custom-meta><meta-name>toc-levels</meta-name><meta-value>0</meta-value></custom-meta><custom-meta><meta-name>volume-type</meta-name><meta-value>Regular</meta-value></custom-meta><custom-meta><meta-name>journal-product</meta-name><meta-value>NonStandardArchiveJournal</meta-value></custom-meta><custom-meta><meta-name>numbering-style</meta-name><meta-value>Unnumbered</meta-value></custom-meta><custom-meta><meta-name>article-grants-type</meta-name><meta-value>OpenChoice</meta-value></custom-meta><custom-meta><meta-name>metadata-grant</meta-name><meta-value>OpenAccess</meta-value></custom-meta><custom-meta><meta-name>abstract-grant</meta-name><meta-value>OpenAccess</meta-value></custom-meta><custom-meta><meta-name>bodypdf-grant</meta-name><meta-value>OpenAccess</meta-value></custom-meta><custom-meta><meta-name>bodyhtml-grant</meta-name><meta-value>OpenAccess</meta-value></custom-meta><custom-meta><meta-name>bibliography-grant</meta-name><meta-value>OpenAccess</meta-value></custom-meta><custom-meta><meta-name>esm-grant</meta-name><meta-value>OpenAccess</meta-value></custom-meta><custom-meta><meta-name>online-first</meta-name><meta-value>false</meta-value></custom-meta><custom-meta><meta-name>pdf-file-reference</meta-name><meta-value>BodyRef/PDF/41467_2021_Article_21593.pdf</meta-value></custom-meta><custom-meta><meta-name>pdf-type</meta-name><meta-value>Typeset</meta-value></custom-meta><custom-meta><meta-name>target-type</meta-name><meta-value>OnlinePDF</meta-value></custom-meta><custom-meta><meta-name>issue-type</meta-name><meta-value>Regular</meta-value></custom-meta><custom-meta><meta-name>article-type</meta-name><meta-value>OriginalPaper</meta-value></custom-meta><custom-meta><meta-name>journal-subject-primary</meta-name><meta-value>Science, Humanities and Social Sciences, multidisciplinary</meta-value></custom-meta><custom-meta><meta-name>journal-subject-secondary</meta-name><meta-value>Science, Humanities and Social Sciences, multidisciplinary</meta-value></custom-meta><custom-meta><meta-name>journal-subject-secondary</meta-name><meta-value>Science, multidisciplinary</meta-value></custom-meta><custom-meta><meta-name>journal-subject-collection</meta-name><meta-value>Science (multidisciplinary)</meta-value></custom-meta><custom-meta><meta-name>open-access</meta-name><meta-value>true</meta-value></custom-meta></custom-meta-group></article-meta><notes notes-type="AuthorContribution"><p>These authors contributed equally: J. Mark Wilkinson, Eleftheria Zeggini.</p></notes><notes notes-type="CopyrightComment"><title>Copyright comment</title><p>corrected publication 2021</p></notes></front><body><sec id="Sec1" sec-type="introduction"><title>Introduction</title><p id="Par4">Osteoarthritis is a severe, debilitating disease that affects the whole joint organ and is hallmarked by cartilage degeneration and synovial hypertrophy. As of 2019, osteoarthritis is estimated to affect 528 Million people worldwide and be the 15th leading cause of years lived with disability<sup><xref ref-type="bibr" rid="CR1">1</xref></sup>. The lifetime risk of developing symptomatic knee and hip osteoarthritis is estimated to be 45 and 25%, respectively<sup><xref ref-type="bibr" rid="CR2">2</xref>,<xref ref-type="bibr" rid="CR3">3</xref></sup>, and is on an upward trajectory commensurate with rises in obesity and the ageing population. Between 2010 and 2019, the global prevalence of osteoarthritis and the resulting years lived with disability have both risen by 27.5%<sup><xref ref-type="bibr" rid="CR1">1</xref></sup>. In 2013, osteoarthritis cost the United States $304 Billion<sup><xref ref-type="bibr" rid="CR4">4</xref></sup> and was the second most costly health condition treated at US hospitals, accounting for 4.3% of the combined costs for all hospitalisations<sup><xref ref-type="bibr" rid="CR5">5</xref></sup>. Osteoarthritis-associated reduced physical activity results in a standardised all-cause mortality ratio of 1.55 (95% confidence interval 1.41 to 1.70) for its sufferers versus the general population<sup><xref ref-type="bibr" rid="CR6">6</xref></sup>. Disease management focusses on alleviating pain, and in end-stage disease the only treatment is joint replacement surgery, emphasising the clear and urgent need to develop new therapies. To achieve this, we need to improve our understanding of the underlying molecular pathophysiology.</p><p id="Par5">Epidemiological risk factors for the disease have been well-established and include older age, female sex, obesity, joint morphology and injury, and family history. The heritability of osteoarthritis has been estimated to range between 40 (for knee osteoarthritis) and 60% (for hip osteoarthritis)<sup><xref ref-type="bibr" rid="CR7">7</xref></sup>. Genome-wide association studies (GWAS) have identified ~90 robustly-replicating risk loci<sup><xref ref-type="bibr" rid="CR8">8</xref></sup>. However, the molecular landscape of osteoarthritis-relevant tissue has not been similarly characterised by large-scale efforts such as the GTEx<sup><xref ref-type="bibr" rid="CR9">9</xref></sup>, ENCODE<sup><xref ref-type="bibr" rid="CR10">10</xref></sup> and RoadMap Epigenomics<sup><xref ref-type="bibr" rid="CR11">11</xref></sup> projects.</p><p id="Par6">In this work, we perform a deep characterisation of the transcriptional and proteomic landscape of disease in chondrocytes and synoviocytes extracted from primary joint tissue of osteoarthritis patients. We define molecular quantitative trait loci, identify likely effector genes for GWAS signals, characterise molecular features of cartilage degradation and highlight drug development and repurposing opportunities through analysis of transcriptional signature changes.</p></sec><sec id="Sec2" sec-type="results"><title>Results</title><sec id="Sec3"><title>Generation of human joint tissue molecular profiles</title><p id="Par7">We collected and characterised low-grade (intact; preserved) and high-grade (highly degraded; lesioned) cartilage, and synovial tissue from patients undergoing joint replacement for osteoarthritis (see Methods). The availability of paired low-grade and high-grade cartilage samples enables the comparison between these two disease states in affected primary tissue within the same individual, with high-grade cartilage showing more advanced cartilage degradation. All three tissues were profiled by RNA sequencing, and cartilage samples were additionally profiled by quantitative proteomics (Fig. <xref rid="Fig1" ref-type="fig">1</xref> and Supplementary Fig. <xref ref-type="supplementary-material" rid="MOESM1">1</xref>). We generated genome-wide genotype data from peripheral blood to define molecular quantitative trait loci (molQTLs) in each tissue type and omics level.<fig id="Fig1"><label>Fig. 1</label><caption xml:lang="en"><title>Study design leveraging multi-omics profiling from osteoarthritis patient tissues.</title><p>We examined the molecular characteristics of osteoarthritis by profiling mRNA and proteins from low-grade cartilage, high-grade cartilage and synovium tissue of over 100 patients undergoing total joint replacement for osteoarthritis, and combining these data with patient genotypes. We identified genetic variants influencing mRNA or protein levels, several of which co-localise with genetic risk variants for osteoarthritis. We also identified molecular markers of cartilage degeneration, creating a gene expression profile of degeneration, and shortlisting existing drugs or compounds that reverse this profile in cell experiments.</p></caption><p><graphic specific-use="HTML" mime-subtype="PNG" xlink:href="MediaObjects/41467_2021_21593_Fig1_HTML.png"/></p></fig></p></sec><sec id="Sec4"><title>Molecular QTLs in osteoarthritis tissues</title><p id="Par8">Identification of molecular QTLs can provide a better understanding of the transcriptional regulation of key cell types across disease stages. We identified <italic>cis</italic> expression QTLs (<italic>cis</italic>-eQTLs) for 1891 genes in at least one tissue (Fig. <xref rid="Fig2" ref-type="fig">2</xref>), with high correlation of effects across the tissues studied (Supplementary Fig. <xref ref-type="supplementary-material" rid="MOESM1">2</xref>). The direction of effect was concordant across all <italic>cis-</italic>eQTLs detected in both low-grade and high-grade cartilage. We identified <italic>cis</italic> protein QTLs (<italic>cis</italic>-pQTLs) for 38 genes in at least one tissue, with a similarly strong correlation across low-grade and high-grade cartilage (Supplementary Figs. <xref ref-type="supplementary-material" rid="MOESM1">3</xref>, <xref ref-type="supplementary-material" rid="MOESM1">4</xref> and Supplementary Note <xref ref-type="supplementary-material" rid="MOESM1">1</xref>).<fig id="Fig2"><label>Fig. 2</label><caption xml:lang="en"><title>Molecular QTLs in osteoarthritis disease tissue.</title><p><bold>a</bold> eQTL overlap between tissues, for a total of 1891 genes with a least one eQTL (left) and 219,709 eQTL gene-variant pairs (right). 49% of detected eQTLs are not tissue-specific. <bold>b</bold> An example of differential QTL effect: a molecular QTL present in high-grade, but not low-grade cartilage (posterior probability <italic>m</italic> &gt; 0.9 and <italic>m</italic> &lt; 0.1, respectively), or vice versa. The boxplots show expression (residuals after regressing the normalised expression data on the 15 PEER factors, sex and array) at 25th to 75th percentiles, with centre at the median and whiskers extend to 1.5 times the interquartile range. NES FastQTL normalised effect size, <italic>P</italic> FastQTL association <italic>P</italic> value, n number of individuals included in the analysis for each genotype. <bold>c</bold> Genes with ≥5 differential eQTL variants (all genes see Supplementary Data <xref ref-type="supplementary-material" rid="MOESM5">1</xref>).</p></caption><p><graphic specific-use="HTML" mime-subtype="PNG" xlink:href="MediaObjects/41467_2021_21593_Fig2_HTML.png"/></p></fig></p></sec><sec id="Sec5"><title>Differential genetic regulation of gene expression</title><p id="Par9">We identified differential regulation of gene expression between high-grade and low-grade cartilage, with 172 variants showing strong evidence for an eQTL effect in one tissue grade (posterior probability &gt;0.9), but not in the other (posterior probability &lt;0.1), termed ‘differential eQTLs' (Fig. <xref rid="Fig2" ref-type="fig">2</xref>, Supplementary Figs. <xref ref-type="supplementary-material" rid="MOESM1">5</xref>, <xref ref-type="supplementary-material" rid="MOESM1">6</xref>, and Supplementary Data <xref ref-type="supplementary-material" rid="MOESM5">1</xref>). We detected 32 genes with differential eQTLs (Supplementary Data <xref ref-type="supplementary-material" rid="MOESM5">1</xref>). These genes function in the regulation of gene expression (high-grade eQTL: <italic>HOXB2, IFITM3, EIF2B3TRAF2, HLCS, APBA1, HLX</italic>; low-grade eQTL: <italic>EARS2, TCEB1, USP16</italic>), nervous system development (high-grade eQTL: <italic>HOXB2, CRLF1, EIF2B3, APBA1, HLX, NEGR1, ARGHAP11B</italic>; low-grade eQTL: <italic>SZT2, NRN1</italic>), response to stress (high-grade eQTL: <italic>IFITM3, EIF2B3, TRAF2, ICAM3</italic>; low-grade eQTL: <italic>PNKP, SZT2, REV1, USP16</italic>), immune response (high-grade eQTL: <italic>IFITM3, IL4I1, CRLF1, TRAF2, ICAM3, HLX</italic>), cell adhesion (high-grade eQTL: <italic>TRAF2, ICAM3, APBA1, HLX, NEGR1</italic>) and catabolic processes (high-grade eQTL: <italic>IL4I1, TRAF2, WDR91, HAAO</italic>; low-grade eQTL: <italic>USP16</italic>). For most of these processes, some genes with differential eQTLs show gain of genetic regulatory associations in high-grade cartilage, while others show loss of such associations compared to low-grade cartilage, suggesting a broader rewiring of regulatory processes (see Supplementary Note <xref ref-type="supplementary-material" rid="MOESM1">2</xref>). Sixteen genes have differential eQTLs located in a regulatory region.</p></sec><sec id="Sec6"><title>Co-localisation of GWAS signals and molecular QTLs in disease tissue</title><p id="Par10">Having established these cell- and disease stage-specific maps of molecular QTLs, we used them to identify effector genes driving GWAS signals, the majority of which reside in non-coding sequence. Co-localisation analysis can indicate whether the same variant underpins association with both disease and gene expression levels. We found strong evidence for co-localisation of five osteoarthritis signals with cartilage molQTLs for <italic>ALDH1A2</italic>, <italic>NPC1</italic>, <italic>SMAD3</italic>, <italic>FAM53A</italic> and <italic>SLC44A2</italic> (Table <xref rid="Tab1" ref-type="table">1</xref>, Fig. <xref rid="Fig3" ref-type="fig">3</xref>, and Supplementary Fig. <xref ref-type="supplementary-material" rid="MOESM1">7</xref>). In all five instances, the GWAS index variant is non-coding. In three cases (<italic>ALDH1A2, SMAD3</italic> and <italic>SLC44A2</italic>), the likely effector gene is the one closest to the lead variant. For the <italic>NCP1</italic> and <italic>FAM53A</italic> loci, the lead variants reside in introns of the <italic>TMEM241</italic> and <italic>SLBP</italic> genes, 141 and 18 kb away from the likely effector gene, respectively.<table-wrap id="Tab1"><label>Table 1</label><caption xml:lang="en"><p>Osteoarthritis GWAS signals with high posterior probability for co-localisation with molecular QTLs.</p></caption><table frame="hsides" rules="groups"><thead><tr><th><p>GWAS variant<sup>a</sup></p></th><th><p>rs10502437</p></th><th><p>rs11732213</p></th><th><p>rs12901372</p></th><th><p>rs1560707</p></th><th><p>rs4775006</p></th></tr></thead><tbody><tr><td><p>Osteoarthritis phenotype</p></td><td><p>All</p></td><td><p>Hip/Knee</p></td><td><p>Hip</p></td><td><p>All</p></td><td><p>Knee</p></td></tr><tr><td><p>Risk allele frequency</p></td><td><p>0.6</p></td><td><p>0.81</p></td><td><p>0.53</p></td><td><p>0.37</p></td><td><p>0.41</p></td></tr><tr><td><p>Odds ratio</p></td><td><p>1.03</p></td><td><p>1.06</p></td><td><p>1.08</p></td><td><p>1.04</p></td><td><p>1.06</p></td></tr><tr><td><p><italic>P</italic> value</p></td><td><p>2.50 × 10<sup>–8</sup></p></td><td><p>8.81 × 10<sup>–10</sup></p></td><td><p>3.46 × 10<sup>–11</sup></p></td><td><p>1.35 × 10<sup>–13</sup></p></td><td><p>8.40 × 10<sup>–10</sup></p></td></tr><tr><td><p>Risk allele</p></td><td><p>A</p></td><td><p>T</p></td><td><p>C</p></td><td><p>T</p></td><td><p>A</p></td></tr><tr><td><p>Gene</p></td><td><p><italic>NPC1</italic></p></td><td><p><italic>FAM53A</italic></p></td><td><p><italic>SMAD3</italic></p></td><td><p><italic>SLC44A2</italic></p></td><td><p><italic>ALDH1A2</italic></p></td></tr><tr><td><p>Risk allele effect on gene expression</p></td><td><p>decrease</p></td><td><p>decrease</p></td><td><p>decrease</p></td><td><p>increase</p></td><td><p>increase</p></td></tr></tbody></table><table-wrap-foot><p><sup>a</sup>Signals are denoted by their lead variants.</p></table-wrap-foot></table-wrap><fig id="Fig3"><label>Fig. 3</label><caption xml:lang="en"><title>GWAS and molecular QTL <italic>P</italic> values in regions with co-localisation of the associations.</title><p>GWAS and molecular QTL <italic>P</italic> values in regions with co-localisation of the associations. PP4 posterior probability for co-localisation. <bold>a</bold><italic>NCP1</italic> eQTLs in high-grade cartilage (<bold>b</bold>) <italic>SLC44A2</italic> eQTLs in high-grade cartilage (<bold>c</bold>) <italic>FAM53A</italic> eQTLs in low-grade cartilage (<bold>d</bold>) <italic>ALDH1A2</italic> pQTLs in low-grade cartilage (<bold>e</bold>) <italic>SMAD3</italic> eQTLs in high-grade cartilage. For <italic>NPC1</italic> and <italic>SMAD3</italic>, Supplementary Fig. <xref ref-type="supplementary-material" rid="MOESM1">7</xref> shows co-localisation with low-grade cartilage molecular QTLs.</p></caption><p><graphic specific-use="HTML" mime-subtype="PNG" xlink:href="MediaObjects/41467_2021_21593_Fig3_HTML.png"/></p></fig></p></sec><sec id="Sec7"><title>Molecular hallmarks of cartilage degradation</title><p id="Par11">To identify molecular signatures associated with disease severity, we characterised gene expression and protein abundance differences between high-grade and low-grade cartilage in the largest sample set to date. We detected significant expression differences for 2557 genes and abundance differences for 2233 proteins at 5% false discovery rate (FDR) (Supplementary Figs. <xref ref-type="supplementary-material" rid="MOESM1">8</xref>, <xref ref-type="supplementary-material" rid="MOESM1">9</xref>, and <xref ref-type="supplementary-material" rid="MOESM1">Supplementary Notes</xref> <xref ref-type="supplementary-material" rid="MOESM1">3</xref>, <xref ref-type="supplementary-material" rid="MOESM1">4</xref>). A total of 409 genes (Supplementary Data <xref ref-type="supplementary-material" rid="MOESM6">2</xref>) demonstrated significant differential expression at both the RNA and protein levels, lending robust cross-omics evidence for their involvement in disease progression. In keeping with previous smaller-scale reports<sup><xref ref-type="bibr" rid="CR12">12</xref>–<xref ref-type="bibr" rid="CR14">14</xref></sup>, extracellular matrix (ECM)-receptor interaction was the primarily activated pathway in high-grade compared to low-grade cartilage (Supplementary Note <xref ref-type="supplementary-material" rid="MOESM1">5</xref>, Supplementary Fig. <xref ref-type="supplementary-material" rid="MOESM1">10</xref>, and Supplementary Data <xref ref-type="supplementary-material" rid="MOESM7">3</xref>).</p><p id="Par12">We used an independent, previously published analysis of RNA sequencing data from 35 osteoarthritis patients<sup><xref ref-type="bibr" rid="CR15">15</xref></sup> to replicate the observed molecular differences between high-grade and low-grade cartilage (within the power constraints of the smaller replication set). Of the differentially expressed genes and proteins, 65.9 and 68.3% showed a concordant direction of effect in the replication data, respectively (both Fisher’s <italic>p</italic> &lt; 10<sup>−10</sup>, Supplementary Data <xref ref-type="supplementary-material" rid="MOESM8">4</xref>). This concordance increased to 77.9% for genes with cross-omics concordant differential expression in this study, indicating additional robustness afforded by cross-omics data integration. We found significantly higher concordance where replication power was highest (88.5% for genes with cross-omics higher expression in high-grade cartilage, compared to 66.7% for genes with cross-omics lower expression in high-grade cartilage, Fisher’s <italic>p</italic> = 8.6 × 10<sup>−6</sup>, Supplementary Data <xref ref-type="supplementary-material" rid="MOESM8">4</xref>).</p></sec><sec id="Sec8"><title>Differentially expressed genes with genetic associations</title><p id="Par13">Ninety-one of the genes with significantly different expression profiles between high- and low-grade cartilage were also associated with genetic risk of osteoarthritis (i.e. among the 238 genes with gene-level significant association in a recent meta-analysis<sup><xref ref-type="bibr" rid="CR8">8</xref></sup>; Supplementary Data <xref ref-type="supplementary-material" rid="MOESM9">5</xref>). A total of 54 genes also showed concordant molecular differences in the independent replication dataset, providing further evidence for the potential involvement of these genes in osteoarthritis disease processes. For example, variants in <italic>ALDH1A2</italic> are associated with knee osteoarthritis, and we found significantly higher <italic>ALDH1A2</italic> gene expression and lower protein abundance in high-grade cartilage (suggesting a potential role for post-transcriptional regulation for this gene). For <italic>SLC39A8</italic>, the GWAS signal was fine-mapped with posterior probability of 0.999 to a single missense variant predicted to be possibly deleterious by PolyPhen-2<sup><xref ref-type="bibr" rid="CR16">16</xref></sup>, and the gene demonstrated higher expression levels in high-grade cartilage.</p></sec><sec id="Sec9"><title>Candidate therapeutic compounds and drug targets</title><p id="Par14">Having characterised the molecular differences between low-grade and high-grade cartilage in osteoarthritis, we sought to identify compounds capable of reversing these changes.</p><p id="Par15">Using in vitro drug screen data from ConnectivityMap<sup><xref ref-type="bibr" rid="CR17">17</xref></sup>, we identified 19 compounds that induced strong opposing gene expression signatures, reducing the expression of genes with cross-omics higher expression in high-grade cartilage (Table <xref rid="Tab2" ref-type="table">2</xref>, Supplementary Data <xref ref-type="supplementary-material" rid="MOESM10">6</xref>, and Supplementary Note <xref ref-type="supplementary-material" rid="MOESM1">6</xref>). Several of these compounds foster known biological relevance to osteoarthritis, including the oestrogen receptor agonists diethylstilbestrol and alpha-estradiol, consistent with established epidemiological data showing an association between osteoarthritis and oestrogen deficiency<sup><xref ref-type="bibr" rid="CR18">18</xref></sup>. Although studies of oestrogen therapy for osteoarthritis have been inconclusive<sup><xref ref-type="bibr" rid="CR19">19</xref>,<xref ref-type="bibr" rid="CR20">20</xref></sup>, this screening approach could facilitate further refinement of existing drug groups and allow more focussed investigational molecule development.<table-wrap id="Tab2"><label>Table 2</label><caption xml:lang="en"><p>Compounds with strongest evidence for inducing gene expression signatures that counter differences between high-grade and low-grade cartilage, based on data from ConnectivityMap<sup><xref ref-type="bibr" rid="CR17">17</xref></sup> and the clue.io platform.</p></caption><table frame="hsides" rules="groups"><thead><tr><th><p>Name<sup>a</sup></p></th><th><p>Description</p></th><th><p>DE targets<sup>b</sup></p></th></tr></thead><tbody><tr><td><p>Emetine</p></td><td><p>protein synthesis inhibitor</p></td><td><p>RPS2 (P+)</p></td></tr><tr><td><p>Rucaparib</p></td><td><p>PARP inhibitor</p></td><td><p>PARP2 (R−)</p></td></tr><tr><td><p>Alpha-estradiol</p></td><td><p>oestrogen receptor agonist</p></td><td><p>KCNMA1 (R−, P−)</p></td></tr><tr><td><p>VEGF-receptor-2-kinase-inhibitor-IV</p></td><td><p>VEGFR inhibitor</p></td><td/></tr><tr><td><p>IB-MECA</p></td><td><p>adenosine receptor agonist, granulocyte colony stimulating factor agonist</p></td><td/></tr><tr><td><p>Diethylstilbestrol</p></td><td><p>oestrogen receptor agonist, chloride channel blocker</p></td><td/></tr><tr><td><p>KIN001-220</p></td><td><p>Aurora kinase inhibitor</p></td><td/></tr><tr><td><p>SB-216763</p></td><td><p>glycogen synthase kinase inhibitor</p></td><td><p>GSK3B (R+, P+), CDK2 (P−)</p></td></tr><tr><td><p>RHO-kinase-inhibitor-III[rockout]</p></td><td><p>ROCK inhibitor</p></td><td><p>IMPDH2 (P−)</p></td></tr><tr><td><p>Nornicotine</p></td><td><p>acetylcholine receptor agonist</p></td><td/></tr></tbody></table><table-wrap-foot><p><sup>a</sup>The 10 compounds with strongest evidence based on median of cell lines are shown; full results see Supplementary Data <xref ref-type="supplementary-material" rid="MOESM10">6</xref>.</p><p><sup>b</sup>Drug targets as listed in ConnectivityMap with differential expression between high-grade and low-grade cartilage on RNA (R) or protein (P) level, ‘+' and ‘−' indicate higher or lower expression high-grade cartilage, respectively.</p></table-wrap-foot></table-wrap></p><p id="Par16">Further notable examples of compounds with the potential to reverse molecular changes include IB-MECA, VEGF-receptor-2-kinase-inhibitor-IV and nornicotine. IB-MECA is used as an anti-inflammatory drug in rheumatoid arthritis, and has been shown to prevent cartilage damage, osteoclast/osteophyte formation and bone destruction in a rat model of chemically-induced osteoarthritis<sup><xref ref-type="bibr" rid="CR21">21</xref></sup>. VEGF-receptor-2-kinase-inhibitor-IV is a <italic>VEGF</italic> receptor inhibitor. <italic>VEGF</italic> modulates chondrocyte survival during development, is essential for bone formation and skeletal growth, and dysregulated in osteoarthritis<sup><xref ref-type="bibr" rid="CR22">22</xref></sup>. A VEGFR2 kinase inhibitor has demonstrated therapeutic potential in mice<sup><xref ref-type="bibr" rid="CR23">23</xref></sup>. Nornicotine is a demethylated derivative of nicotine found in plants, and there is a well-established epidemiological inverse relationship between smoking and osteoarthritis<sup><xref ref-type="bibr" rid="CR8">8</xref>,<xref ref-type="bibr" rid="CR24">24</xref></sup>. Thus, synthetic nicotine derivatives may have a candidate role in osteoarthritis prevention.</p><p id="Par17">We further identified 36 genes for which knock-down or overexpression has the potential to reduce the molecular differences between high-grade and low-grade cartilage (Supplementary Data <xref ref-type="supplementary-material" rid="MOESM10">6</xref>), for example knock-down of <italic>IL11</italic>. IL11 is a cytokine with a key role in inflammation, and several therapeutics that inhibit IL11 signalling are in development against a range of inflammatory and fibrotic diseases<sup><xref ref-type="bibr" rid="CR25">25</xref>,<xref ref-type="bibr" rid="CR26">26</xref></sup>. Variation in <italic>IL11</italic> is associated with increased risk of hip osteoarthritis<sup><xref ref-type="bibr" rid="CR8">8</xref></sup>, and the gene is upregulated in osteoarthritis knee tissue<sup><xref ref-type="bibr" rid="CR27">27</xref></sup>, showing the most significant upregulation in high-grade cartilage in the independent replication dataset<sup><xref ref-type="bibr" rid="CR15">15</xref></sup> (22.8-fold higher expression, FDR = 1.5 × 10<sup>−20</sup>). These findings provide strong supportive evidence for downregulation of <italic>IL11</italic> as a potential therapeutic intervention for osteoarthritis.</p></sec></sec><sec id="Sec10" sec-type="discussion"><title>Discussion</title><p id="Par18">To date, deep molecular profiling of cells of relevance to osteoarthritis has been absent from large-scale efforts such as GTEx, ENCODE and RoadMap, in part because of the challenges associated with low cellularity and high ECM content of cartilage<sup><xref ref-type="bibr" rid="CR28">28</xref>,<xref ref-type="bibr" rid="CR29">29</xref></sup>. Here, we address this gap by the systematic study of paired low-grade and high-grade cartilage and synovium tissues from 115 patients with osteoarthritis to provide the first deep molecular QTL map of cell types directly involved in disease. These data are made publically available (see Data availability). We identify genotype-dependent, divergent patterns of gene regulation between diseased and healthy cartilage, and between cartilage and synovium, that underline the biological specificity of disease stage and cell type when investigating regulatory variant function.</p><p id="Par19">The co-localisation of GWAS signals and molecular QTLs in disease tissue helps pinpoint the identity of causal genes for hithert unsolved association signals. A previous study<sup><xref ref-type="bibr" rid="CR8">8</xref></sup> examined the co-localisation of osteoarthritis GWAS signals with eQTLs from the GTEx resource<sup><xref ref-type="bibr" rid="CR9">9</xref></sup>, which does not include cartilage. Several osteoarthritis GWAS signals were found to co-localise with eQTLs in only 1 or 2 of 53 tissues (e.g. <italic>ALDH1A2</italic>: ovary and tibial artery, <italic>SMAD3</italic>: skeletal muscle; <italic>SLC44A2</italic>: adrenal gland), without clear transferability of results to disease-relevant tissue. Here, we provide robust evidence that these three GWAS signals co-localise with molecular QTLs in primary cartilage. This availability of molQTL data from chondrocytes and synoviocytes offers a resource that will enable the resolution of further genetic association signals emerging from ongoing large-scale efforts in osteoarthritis (e.g. <ext-link xlink:href="https://www.genetics-osteoarthritis.com" ext-link-type="url">https://www.genetics-osteoarthritis.com</ext-link>) and further traits of musculoskeletal relevance, in which these cell types are of importance.</p><p id="Par20">The strong correlation between RNA and protein level signatures within cell type and disease state and the contrast between disease states provides both strong internal consistency and with previous observations of ECM remodelling by chondrocytes in an inflammatory environment as a key modifiable molecular signature in the degeneration process<sup><xref ref-type="bibr" rid="CR30">30</xref>,<xref ref-type="bibr" rid="CR31">31</xref></sup>. We further demonstrate the external validity of these hallmark signatures by showing good concordance with publically-available RNA-seq data from a smaller, independent, human osteoarthritis cohort<sup><xref ref-type="bibr" rid="CR15">15</xref></sup>, particularly for genes whose expression is increased in high-grade cartilage.</p><p id="Par21">On a related theme, our identification of an association between 91 established osteoarthritis variants and disease state, cell-specific gene expression and protein profiles highlights the value of integrating multi-omics data with genetic association summary statistics to identify likely effector genes for GWAS signals, and hence targets for development of new therapies. The subsequent ConnectivityMap analyses built upon this theme to identify molecules with potential clinical application to reverse the molecular signatures characteristic of diseased chondrocytes, an in-silico complementary approach to conventional compound library screening. Here, we identify 19 molecules and 36 genes with biological supportive evidence for validation as candidate drugs in experimental models, such as that recently applied by Shi et al. to demonstrate that the small molecule BNTA stimulates expression of ECM components while supressing inflammatory mediators in human osteoarthritic cartilage, an effect mediated through superoxide dismutase 3<sup><xref ref-type="bibr" rid="CR30">30</xref></sup>.</p><p id="Par22">Although these data demonstrate the context-specificity of molecular indicators of disease, this work also has limitations. Cartilage remains a difficult tissue to study at the molecular level, given the low cell accessibility. The analyses were performed on extracted tissues characterised by macroscopic grading using the International Cartilage Repair Society (ICRS) scoring system<sup><xref ref-type="bibr" rid="CR32">32</xref></sup> and although all tissues were collected from weight-bearing areas (i.e. mechanically-loaded), a degree of cell disease state heterogeneity is inevitable in the pooled sample collection. However, the collection of both disease states from the same individual eliminates inter-individual variation as a source of confounding. Finally, although these data represent the largest cohort of its depth and kind in osteoarthritis, they are observational in nature. Future biological experiments (e.g. by gene knock-out or over-expression) will be key to prove which gene expression changes are causal in disease development and progression. The genes with convergent evidence from this study as well as osteoarthritis GWAS are presented to help such research and accelerate the pathway to translation.</p><p id="Par23">In summary, by integrating multiple layers of omics data, we have generated a first molecular QTL map for osteoarthritis-relevant tissues and have helped resolve genetic association signals by identifying likely effector genes. We demonstrate how integrating multi-omics data in primary human complex disease tissue can serve as a valuable approach that moves from basic discovery to accelerated translational opportunities. Our findings identify drug repurposing opportunities and allow novel investigational avenues for therapy development, responding to the urgent clinical need of patients suffering from osteoarthritis.</p></sec><sec id="Sec11" sec-type="methods"><title>Methods</title><sec id="Sec12"><title>Study participants</title><p id="Par24">We collected tissue samples from 115 patients undergoing total joint replacement surgery in 4 cohorts: 12 knee osteoarthritis patients (cohort 1; 2 women, 10 men, age 50–88 years, mean 68 years); 20 knee osteoarthritis patients (cohort 2; 14 women, 6 men, age 54–82 years, mean 70 years); 13 hip osteoarthritis patients (cohort 3; 8 women, 5 men, age 44–84 years, mean 62 years); 70 knee osteoarthritis patients (cohort 4; 42 women, 28 men, age 38–84 years, mean 70 years).</p><p id="Par25">All patients provided written, informed consent prior to participation in the study. Matched low-grade and high-grade cartilage samples were collected from each patient, while synovial lining samples were collected from patients in cohorts 2 and 4. All cartilage samples were collected from weight-bearing areas of the joint to ensure that any differences observed between low- and high-grade cartilage reflect disease progression stage rather than differential biomechanical stress.</p><sec id="Sec13"><title>Cohorts 1, 2, 4 (knee osteoarthritis)</title><p id="Par26">This work was approved by Oxford NHS REC C (10/H0606/20 and 15/SC/0132), and samples were collected under Human Tissue Authority license 12182, Sheffield Musculoskeletal Biobank, University of Sheffield, UK.</p><p id="Par27">We confirmed a joint replacement for osteoarthritis, with no history of significant knee surgery (apart from meniscectomy), knee infection, or fracture, and no malignancy within the previous 5 years. We further confirmed that no patient used glucocorticoid use (systemic or intra-articular) within the previous 6 months, or any other drug associated with immune modulation. For cohort 1, cartilage samples were scored using the OARSI cartilage classification system<sup><xref ref-type="bibr" rid="CR33">33</xref>,<xref ref-type="bibr" rid="CR34">34</xref></sup>. From each patient, we obtained one sample with high OARSI grade signifying high-grade degeneration ('high-grade sample'), and one cartilage sample with low OARSI grade signifying healthy tissue or low-grade degeneration ('low-grade sample').</p><p id="Par28">For cohorts 2 and 4, cartilage samples were scored macroscopically using the ICRS scoring system<sup><xref ref-type="bibr" rid="CR32">32</xref></sup>. From each patient, we obtained one sample of ICRS grade 3 or 4 signifying high-grade degeneration ('high-grade sample'), and one cartilage sample of ICRS grade 0 or 1 signifying healthy tissue or low-grade degeneration ('low-grade sample'). For cohorts 2 and 4, we also collected synovial membrane from the suprapatellar region of the knee joint.</p><p id="Par29">Finally, from all patients in cohorts 1, 2 and 4, we also obtained a blood sample to extract DNA for genotyping.</p></sec><sec id="Sec14"><title>Cohort 3 (hip osteoarthritis)</title><p id="Par30">Samples were collected under National Research Ethics approval reference 11/EE/0011, Cambridge Biomedical Research Centre Human Research Tissue Bank, Cambridge University Hospitals, UK.</p><p id="Par31">We confirmed osteoarthritis disease status by examination of the excised femoral head. From each patient, we obtained a cartilage sample showing a fibrillated or fissured surface signifying high-grade degeneration ('high-grade sample'), one cartilage sample showing a smooth shiny appearance signifying healthy tissue or low-grade degeneration ('low-grade sample').</p></sec></sec><sec id="Sec15"><title>Isolation of chondrocytes</title><p id="Par32">For cohorts 1, 2 and 4, we followed a previously established protocol to isolate chondrocytes<sup><xref ref-type="bibr" rid="CR12">12</xref></sup> with the details as follows. Osteochondral samples were transported in Dulbecco’s modified Eagle’s medium (DMEM)/F-12 (1:1) (Life Technologies) supplemented with 2 mM glutamine (Life Technologies), 100 U/ml penicillin, 100 μg/ml streptomycin (Life Technologies), 2.5 μg/ml amphotericin B (Sigma-Aldrich) and 50 μg/ml ascorbic acid (Sigma-Aldrich) (serum free media). Half of each sample was then taken forward for chondrocyte extraction. Cartilage was removed from the bone, dissected and washed twice in 1xPBS. Tissue was digested in 3 mg/ml collagenase type I (Sigma-Aldrich) in serum free media overnight at 37 °C on a flatbed shaker. The resulting cell suspension was passed through a 70 μm cell strainer (Fisher Scientific) and centrifuged at 400×<italic>g</italic> for 10 min. Subsequently, the cell pellet was washed twice in serum free media and centrifuged at 400×<italic>g</italic> for 10 min. The resulting cell pellet was resuspended in serum free media. Cells were counted using a haemocytometer and the viability checked using trypan blue exclusion (Invitrogen). The optimal cell number for spin column extraction from cells was between 4 × 10<sup>6</sup> and 1 × 10<sup>7</sup>. Cells were then pelleted and homogenised.</p><p id="Par33">For cohort 3, the extraction of chondrocytes has previously been described<sup><xref ref-type="bibr" rid="CR35">35</xref></sup> in the majority of these samples, with the remaining samples following the same protocol. The protocol was based on that for cohorts 1, 2, 4 and highly similar as described in the following. Each cartilage portion was minced with a scalpel and placed in 20 ml of Dulbecco’s modified Eagle medium (Invitrogen) containing 10% foetal bovine serum (Invitrogen) and 6 mgml<sup>−1</sup> collagenase A (Sigma). The tissue culture flasks were incubated overnight to digest the cartilage pieces. The resulting cell suspension was passed through a 30 μm filter (Miltenyi) and centrifuged at 400×<italic>g</italic> for 10 min. The cell pellet was then re-suspended in 1 ml of PBS and counted on a haemocytometer following 1:1 mixing with trypan blue to determine cell viability.</p></sec><sec id="Sec16"><title>Isolation of synoviocytes</title><p id="Par34">We followed a previously established protocol to process synovial samples<sup><xref ref-type="bibr" rid="CR36">36</xref></sup>, with details as follows. Synovial samples were transported in serum free media, as described above. The synovial membrane was dissected from underlying tissue then trypsinised for 1 h. Tissue was then digested in 1 mg/ml Collagenase Blend H (Sigma Aldrich) in serum free media overnight at 37 °C on a flatbed shaker. The resulting cell suspension was passed through a 100 μm cell strainer (Fisher Scientific) and centrifuged at 400×<italic>g</italic> for 10 min. Subsequently, the cell pellet was washed twice in serum free media and centrifuged at 400×<italic>g</italic> for 10 min. The resulting cell pellet was resuspended in serum free media. Cells were counted using a haemocytometer and the viability checked using trypan blue exclusion (Invitrogen). The optimal cell number for spin column extraction from cells was between 4 × 10<sup>6</sup> and 1 × 10<sup>7</sup>. Cells were then pelleted and homogenised.</p></sec><sec id="Sec17"><title>DNA, RNA and protein extraction</title><p id="Par35">DNA, RNA, and protein extraction was carried out using Qiagen AllPrep DNA/RNA/Protein Mini Kit following the manufacturer’s instructions for cohorts 1, 2 and 4, with small variations for cohort 3 as previously described<sup><xref ref-type="bibr" rid="CR35">35</xref></sup> and recapitulated in the Supplementary Methods. Samples were frozen at −80 °C (cohorts 1, 2, 4) or −70 °C (cohort 3) prior to assays.</p></sec><sec id="Sec18"><title>RNA sequencing</title><p id="Par36">We performed a gene expression analysis on samples from 113 patients (Supplementary Data <xref ref-type="supplementary-material" rid="MOESM11">7</xref>). We purified poly-A tailed RNA (mRNA) from total RNA using Illumina’s TruSeq RNA Sample Prep v2 kits. We then fragmented the mRNA using metal ion-catalysed hydrolysis and synthesised a random-primed cDNA library. The resulting double-strand cDNA was used as the input to a standard Illumina library prep, whereby ends were repaired to produce blunt ends by a combination of fill-in reactions and exonuclease activity. We performed A-tailing to allow samples to be pooled, by adding an ‘A' base to the blunt ends and ligation to Illumina Paired-end Sequencing adaptors containing unique index sequences. Due to better performance, the 10-cycle PCR amplification of libraries was carried out using KAPA Hifi Polymerase. A post-PCR Agilent Bioanalyzer was used to quantify samples, followed by sample pooling and size-selection of pools using the LabChip XT Caliper. The multiplexed libraries were sequenced on the Illumina HiSeq 2000 for cohort 1 and HiSeq 4000 for cohorts 2–4 (75 bp paired-ends). Sequenced data underwent initial analysis and quality control (QC) on reads as standard. The sequencing depth was similar across samples, with 90% of samples passing final QC (see below) having 87.2–129.2 million reads.</p></sec><sec id="Sec19"><title>Proteomics</title><p id="Par37">Proteomics analysis was performed on cartilage samples from 103 patients (Supplementary Data <xref ref-type="supplementary-material" rid="MOESM11">7</xref>).</p><p id="Par38">For cohort 1, all steps of protein digestion, 6-plex TMT labelling, peptide fractionation and LC-MS analysis on the Dionex Ultimate 3000 UHPLC system coupled with the high-resolution LTQ Orbitrap Velos mass spectrometer (Thermo Scientific), were previously described<sup><xref ref-type="bibr" rid="CR12">12</xref></sup> and are recapitulated in the Supplementary Methods. The sample preparation protocol formed the basis of processing for cohorts 2–4 using 10-plex TMT labelling and an Orbitrap Fusion Tribrid Mass Spectrometer (Thermo Scientific) with otherwise only minor alterations as described in the Supplementary Methods.</p></sec><sec id="Sec20"><title>Genotyping</title><p id="Par39">We used Illumina HumanCoreExome-12v1-1 for genoting cohort 1 and Illumina InfiniumCoreExome-24v1-1 for genotyping cohort 2–4 patients.</p></sec><sec id="Sec21"><title>Quantification of RNA levels</title><p id="Par40">We used samtools v1.3.1<sup><xref ref-type="bibr" rid="CR37">37</xref></sup> and biobambam v0.0.191<sup><xref ref-type="bibr" rid="CR38">38</xref></sup> to convert cram to fasq files after exclusion of reads that failed QC. We applied FastQC v0.11.5 to check sample quality<sup><xref ref-type="bibr" rid="CR39">39</xref></sup> and excluded nine samples (Supplementary Data <xref ref-type="supplementary-material" rid="MOESM11">7</xref>).</p><p id="Par41">We obtained transcript-level quantification using salmon 0.8.2<sup><xref ref-type="bibr" rid="CR40">40</xref></sup> (with–gcBias and –seqBias flags to account for potential biases) and the GRCh38 cDNA assembly release 87 downloaded from Ensembl [http://ftp.ensembl.org/pub/release-87/fasta/homo_sapiens/cdna/]. We used tximport<sup><xref ref-type="bibr" rid="CR41">41</xref></sup> to convert transcript-level to gene-level scaled transcripts per million (TPM) estimates, with estimates for 39,037 genes based on Ensembl gene IDs.</p><p id="Par42">We excluded four samples due to low mapping rate (&lt;80%), three samples due to non-European ancestry recorded in the clinic, 18 samples due to low RIN (&lt;5), two samples as duplicates, eight samples due to abnormal gene read density plots (detected separately in cartilage and synovium for three cartilage and five synovium samples; all exclusions are listed in Supplementary Data <xref ref-type="supplementary-material" rid="MOESM11">7</xref>).</p><p id="Par43">The final gene expression dataset included 259 samples (Supplementary Fig. <xref ref-type="supplementary-material" rid="MOESM1">1</xref>; 87 patients’ low-grade and 95 high-grade cartilage samples with 15,249 genes that showed counts per million (CPM) of ≥1 in ≥40 samples, and 77 patients’ synovium samples with 16,004 genes that showed CPM ≥1 in ≥20 samples).</p></sec><sec id="Sec22"><title>Quantification of protein levels</title><p id="Par44">To carry out protein identification and quantification, we submitted the mass spectra to SequestHT search in Proteome Discoverer 2.1. The precursor mass tolerance was set at 30 ppm (Orbitrap Velos data, cohort 1) or 20 ppm (Fusion data, cohorts 2–4). For the CID spectra, we set the fragment ion mass tolerance to 0.5 Da; for the HCD spectra, to 0.02 Da. Spectra were searched for fully tryptic peptides with maximum two miss-cleavages and minimum length of six amino acids. We specified static modifications as TMT6plex at N-termimus, K and Carbamidomethyl at C; dynamic modifications included deamidation of N,Q and oxidation of M. For each peptide, we allowed for a maximum two different dynamic modifications with a maximum of two repetitions. We used the Percolator node to estimate peptide confidence. We set the peptide FDR at 1% and based validation on the <italic>q</italic> value and decoy database search. We searched all spectra against a UniProt fasta file that contained 20,165 reviewed human entries. The Reporter Ion Quantifier node included a TMT-6plex (Velos data, cohort 1) or TMT-10plex (Fusion data, cohorts 2–4) custom Quantification Method with integration window tolerance at 20 or 15 ppm, respectively. As integration methods, we used the Most Confident Centroid at the MS2 or MS3 level. We only used peptides uniquely belonging to protein groups for quantification.</p><p id="Par45">We excluded samples from four patients due to non-European ancestry (Supplementary Data <xref ref-type="supplementary-material" rid="MOESM11">7</xref>). The final dataset included low-grade and high-grade cartilage samples each from 99 patients, with 4801 proteins was observed in ≥30% of samples, and 1677 proteins in all samples, in line with the resolution depth of the isobaric labelling method employed. To account for protein loading, abundance values were normalised by the sum of all protein abundances in a given sample, then log2-transformed and quantile normalised.</p></sec><sec id="Sec23"><title>Genotype analysis and QC</title><p id="Par46">Genotypes were called using GenCall (Illumina) and mapped to GRC37/hg19 using online tools (<ext-link xlink:href="http://www.well.ox.ac.uk/~wrayner/strand/index.html" ext-link-type="url">http://www.well.ox.ac.uk/~wrayner/strand/index.html</ext-link>). QC was carried out using the same method for both arrays. Briefly, we performed a pre-filtering step to exclude samples and variants with a call rate &lt;90%. Sample QC included identity checks correlating the array genotypes to Fluidigm genotypes obtained at sample reception (no samples had a concordance &lt;0.95). We excluded samples based on call rate &lt;98%, heterozygosity distribution outliers performed using two different minor allele frequency (MAF) bins (≥1% MAF and &lt;1% MAF) and sex discrepancies. We performed pairwise identity by descent (IBD) in PLINK<sup><xref ref-type="bibr" rid="CR42">42</xref>,<xref ref-type="bibr" rid="CR43">43</xref></sup> after filtering out variants with MAF &lt; 1% and carrying out linkage disequilibrium based pruning using R<sup>2</sup> 0.2. We retained only patients with pairwise PI_HAT ≤ 0.2. To look at ethnicity we combined all patients from both arrays with data from the 1000 Genomes Project individuals (<ext-link xlink:href="https://www.internationalgenome.org" ext-link-type="url">https://www.internationalgenome.org</ext-link>)<sup><xref ref-type="bibr" rid="CR44">44</xref></sup>. We included overlapping variants only and conducting IBD, as described above, followed by multidimensional scaling using PLINK. Visual ethnic outliers were excluded following examination of the first two components. Variants were excluded if call rate &lt;98% and/or Hardy–Weinberg <italic>p</italic> value (pHWE) &lt;1 × 10<sup>−4</sup>. The final datasets contained 12 patients and 534,694 variants and 99 patients and 527,717 variants for cohorts 1 and 2–4, respectively.</p><p id="Par47">Prior to imputation, all genotypes were combined into a single dataset containing 111 patients and 504,235 overlapping variants. Further QC was performed to exclude any variants with strand, position and allele frequency differences compared to the HRC panel<sup><xref ref-type="bibr" rid="CR45">45</xref></sup> using a HRC preparation checking tool (<ext-link xlink:href="http://www.well.ox.ac.uk/~wrayner/tools/" ext-link-type="url">http://www.well.ox.ac.uk/~wrayner/tools/</ext-link>; v4.2.7). The resulting dataset contained 111 patients and 389,511 variants. We imputed up to HRC panel (v1.1 2016) using the Michigan imputation server (<ext-link xlink:href="https://imputationserver.sph.umich.edu/index.html" ext-link-type="url">https://imputationserver.sph.umich.edu/index.html</ext-link>)<sup><xref ref-type="bibr" rid="CR46">46</xref></sup> with Eagle2 (v2.3) phasing. Post-HRC imputation we used a post-imputation data checking programme (<ext-link xlink:href="http://www.well.ox.ac.uk/~wrayner/tools/Post-Imputation.html" ext-link-type="url">http://www.well.ox.ac.uk/~wrayner/tools/Post-Imputation.html</ext-link>; v1.0.2) to visualise the results and we excluded variants with poor imputation quality (<italic>R</italic><sup>2</sup> &lt; 0.3) and pHWE &lt;1 × 10<sup>−4</sup>. We excluded two patients due to absence of RNA and protein data. The resulting final dataset contained 10,249,108 autosomal variants and 109 patients.</p></sec><sec id="Sec24"><title>Identification of <italic>cis</italic>-eQTLs</title><p id="Par48">For each gene, we considered genetic variants within 1 Mb of the transcription start site (TSS) (definition see below), and followed a similar method to GTEx<sup><xref ref-type="bibr" rid="CR9">9</xref>,<xref ref-type="bibr" rid="CR47">47</xref></sup> as follows.</p><p id="Par49">For each tissue, we included only genes with ≥1 count per million in at least 20% samples and we normalised between samples using TMM (weighted trimmed mean of <italic>M</italic> values)<sup><xref ref-type="bibr" rid="CR48">48</xref></sup> implemented in edgeR<sup><xref ref-type="bibr" rid="CR49">49</xref></sup>. To facilitate cartilage comparisons post-analysis, the previous two steps (exclusions of low expressed genes and the between sample normalisation) were performed with high-grade and low-grade cartilage samples combined. For each tissue separately, we then normalised across samples using an inverse normalisation transformation for each gene. To infer hidden factors associated with cohort, sequencing batch, or other technical differences, we applied probabilistic estimation of expression residuals (PEER)<sup><xref ref-type="bibr" rid="CR50">50</xref></sup> separately to each tissue (PEER C + + version with standard parameters from the R version, i.e. iteration = 1000, bound = 0.001, variance = 0.00001, Alpha a = 0.001, Alpha b = 0.1, Eps a = 0.1, Eps b = 10). We used the GTEx modified version of FastQTL<sup><xref ref-type="bibr" rid="CR51">51</xref></sup> (<ext-link xlink:href="https://github.com/francois-a/fastqtl" ext-link-type="url">https://github.com/francois-a/fastqtl</ext-link>; v6p), which allows for minor allele count filtering, reporting of MAF and calculation of FDR. We determined the TSS for each gene using empirical transcript level expression information from synovium, high-grade and low-grade cartilage samples (see below) and defined the <italic>cis</italic>-mapping region to be 1 Mb in either direction from the TSS. We restricted the analysis to variants with minor allele count of at least 10 in a given tissue. Nominal <italic>p</italic> values for each gene-variant pair were based on linear regression, including 15 PEER factors for the given tissue, sex and genotype array as covariates. We then employed the adaptive permutation scheme with the—permute 1000 10000 option to generate empirical <italic>p</italic> values. Genes with significant eQTLs ('eGenes') were defined at the 5% Storey–Tibshirani FDR using the <italic>q</italic> values generated from the empirical <italic>p</italic> values<sup><xref ref-type="bibr" rid="CR52">52</xref></sup>. For each eGene, significant eQTLs were defined as variants with nominal <italic>p</italic> value below the nominal <italic>p</italic> value threshold for that gene generated in FastQTL.</p><p id="Par50">The normalised effect size (NES) of the eQTL is reported for the alternate allele according to GRC37/hg19.</p></sec><sec id="Sec25"><title>Identification of <italic>cis</italic>-pQTLs</title><p id="Par51">We followed a similar protocol as for <italic>cis</italic>-eQTL analysis, also considering genetic variants within 1 Mb of the TSS (definition see below). For low-grade and high-grade cartilage, we included 1677 proteins that were measured across all samples. We normalised across samples using an inverse normalisation transformation for each gene separately in each tissue. To account for possible technical variation, we used PEER<sup><xref ref-type="bibr" rid="CR50">50</xref></sup> (with parameters as for the eQTL analysis above) and included 26 PEER factors, sex and genotype array as covariates using the GTEx modified version of FastQTL (<ext-link xlink:href="https://github.com/francois-a/fastqtl" ext-link-type="url">https://github.com/francois-a/fastqtl</ext-link>; v6p). We used the TSS established for the eQTL analysis, yielding a unique mapping for 1461 proteins, which were then taken forward. For each protein, we considered variants within a 1 Mb region in either direction from the TSS, restricting further to minor allele count of 10 or higher. We then followed the same procedure as for <italic>cis</italic>-eQTLs to identify variant-protein pairs with significant <italic>cis</italic>-pQTL effects.</p><p id="Par52">We carried out a protein–protein network analysis for the proteins with significant pQTL effects using the STRING v11.0 database<sup><xref ref-type="bibr" rid="CR53">53</xref></sup> (<ext-link xlink:href="https://string-db.org/" ext-link-type="url">https://string-db.org/</ext-link>; see Supplementary Note <xref ref-type="supplementary-material" rid="MOESM1">1</xref>).</p></sec><sec id="Sec26"><title>Sensitivity analysis for <italic>cis</italic>-eQTLs and <italic>cis</italic>-pQTLs</title><p id="Par53">For both eQTLs and pQTLs, we verified that the results were robust by carrying out a sensitivity analysis including patient age and osteoarthritis joint (knee or hip) as covariates in addition to the PEER factors, sex and array (see Supplementary Note <xref ref-type="supplementary-material" rid="MOESM1">1</xref>).</p></sec><sec id="Sec27"><title>Transcription start site (TSS) definition</title><p id="Par54">To determine which transcript to use to define the TSS for each gene, we established the most abundant transcript for each gene. To have the same definition across all tissues, we analysed cartilage and synovium tissues jointly, considering 16,886 genes that passed quantification QC in at least one tissue (based on 15,249 genes in cartilage and 16,004 genes in synovium). For each transcript, we calculated the expression in each sample as scaled TPM using tximport<sup><xref ref-type="bibr" rid="CR41">41</xref></sup>. For each gene, we then obtained the most abundant transcript in each tissue in each patient, and calculated the proportion of samples in which each transcript was the most abundant. The transcript that was the most abundant in the largest proportion of samples was used to define the TSS for the gene. For genes in which more than one transcript was the most abundant, we chose one of the most abundant transcripts at random to define the TSS.</p><p id="Par55">Across all genes, 47.9% had the same most abundant transcript in at least 90% of samples in both cartilage and synovium; 71% of genes had the same most abundant transcript in 60% of samples in cartilage and synovium. We mapped the most abundant transcript (using the ENST identifier) to GRCh37 using ftp://ftp.ensembl.org/pub/grch37/release-87/gtf/homo_sapiens/Homo_sapiens.GRCh37.87.chr.gtf.gz. For 93 transcripts, ENSG identifiers differed between the builds, therefore genes were identified by a composite ID in the format GeneName(b37)_ENSG. We excluded 1940 transcripts with missing start or end positions (largely in patched genome build regions), and ~500 transcripts mapped to chromosomes X, Y or mitochondrial DNA, we established the TSS for 13,180 autosomal genes included in the cartilage and 13,708 genes included in the synovium molQTL analysis.</p></sec><sec id="Sec28"><title>Differential gene regulation in low-grade and high-grade cartilage</title><p id="Par56">To identify differential regulation of gene expression between high- and low-grade cartilage, we used Meta-Tissue v0.5<sup><xref ref-type="bibr" rid="CR54">54</xref></sup> (downloaded from <ext-link xlink:href="http://genetics.cs.ucla.edu/metatissue/index.html" ext-link-type="url">http://genetics.cs.ucla.edu/metatissue/index.html</ext-link>), which implements METASOFT<sup><xref ref-type="bibr" rid="CR55">55</xref></sup>. The <italic>m</italic> value calculated by METASOFT for gene-variant pair in each tissue provides a posterior probability (<italic>m</italic> value) of an effect in that tissue. Consequently, we aimed to identify eQTLs present in one tissue (defined as <italic>m</italic> &gt; 0.9), and absent in the other (defined as <italic>m</italic> &lt; 0.1). We note that there were no <italic>cis</italic>-eQTLs present in both tissues (<italic>m</italic> &gt; 0.9) with opposing direction of effect.</p><p id="Par57">Meta-Tissue restricts covariates input to the same values for each patient across tissues, while different PEER covariates were provided for each tissue in the FastQTL analysis. Hence, for each tissue, we obtained residuals from regressing the normalised expression data on the 15 PEER factors, sex and array, then used the residuals as input for Meta-Tissue. We included genotype dosages based on both the low- and high-grade results for each analysis. We ran METASOFT using the default settings provided in the output script from Meta-Tissue. We only considered eQTLs that were identified in the FastQTL analysis in the appropriate tissue. To identify variants located in regulatory regions, we used Ensembl Variant Effect Predictor (<ext-link xlink:href="http://grch37.ensembl.org/Homo_sapiens/Tools/VEP/" ext-link-type="url">http://grch37.ensembl.org/Homo_sapiens/Tools/VEP/</ext-link>). For 32 genes with differential QTLs, we considered gene annotations from several sources: using Gene Ontology biological process terms, through gene annotations in GOSeq<sup><xref ref-type="bibr" rid="CR56">56</xref></sup> (accessed on 31 May 2020), and from the summaries and functions sections on GeneCards gene pages (<ext-link xlink:href="https://www.genecards.org" ext-link-type="url">https://www.genecards.org</ext-link>, accessed 3 June 2020).</p><p id="Par58">We also carried out a formal enrichment analysis of the 32 genes using GoSeq (see Supplementary Note <xref ref-type="supplementary-material" rid="MOESM1">2</xref>).</p></sec><sec id="Sec29"><title>Co-localisation between molecular QTLs (molQTLs) and osteoarthritis GWAS associations</title><p id="Par59">To examine co-localisation between molQTLs and GWAS associations, we used genome-wide summary statistics from the largest osteoarthritis meta-analysis to date, based on UK Biobank and arcOGEN data<sup><xref ref-type="bibr" rid="CR8">8</xref></sup>. We analysed all 64 genome-wide significant signals using coloc<sup><xref ref-type="bibr" rid="CR57">57</xref></sup>, separately for each tissue and omics level.</p><p id="Par60">In the co-localisation analysis for each signal, we considered the region spanning 100 kb either side of the index variants. If that region overlapped any genes with significant <italic>cis</italic>-eQTLs or <italic>cis</italic>-pQTLs, we extended the region to encompass all variants included in the molQTL analysis for these genes. To formally obtain a posterior probability for co-localisation, we used coloc.fast (<ext-link xlink:href="https://github.com/tobyjohnson/gtx/blob/526120435bb3e29c39fc71604eee03a371ec3753/R/coloc.R" ext-link-type="url">https://github.com/tobyjohnson/gtx/blob/526120435bb3e29c39fc71604eee03a371ec3753/R/coloc.R</ext-link>), a Bayesian statistical test which implements the coloc<sup><xref ref-type="bibr" rid="CR57">57</xref></sup> method. We used the default settings for coloc.fast. We considered a 80% posterior probability of GWAS and molQTL shared association at a single variant ('PP4 ≥ 0.8') to indicate evidence of co-localisation.</p></sec><sec id="Sec30"><title>Differential RNA expression between high-grade and low-grade cartilage</title><p id="Par61">We tested differential expression of 15,249 genes between high-grade and low-grade cartilage using paired samples from 83 patients. To detect robust gene expression differences, we carried out analyses using different software packages as recommended in a landmark survey of best practices<sup><xref ref-type="bibr" rid="CR58">58</xref></sup>, applying limma<sup><xref ref-type="bibr" rid="CR59">59</xref></sup>, edgeR<sup><xref ref-type="bibr" rid="CR60">60</xref></sup> and DESeq2<sup><xref ref-type="bibr" rid="CR61">61</xref></sup>. We also tested five analysis designs with different options to account for technical variation, including SVAseq<sup><xref ref-type="bibr" rid="CR62">62</xref></sup>. In particular, we tested for differential expression using</p><p id="Par62">(1) a paired analysis of intact and degraded samples (i.e. specifying patient ID as covariate);</p><p id="Par63">(2) a paired analysis of intact and degraded samples, with ten additional covariates accounting for technical variation identified by SVAseq<sup><xref ref-type="bibr" rid="CR62">62</xref></sup>;</p><p id="Par64">(3) a paired analysis of intact and degraded samples, with ten RNA sequencing batches as covariates;</p><p id="Par65">(4) an unpaired analysis of intact and degraded samples;</p><p id="Par66">(5) an unpaired analysis of intact and degraded samples, with 19 additional covariates accounting for technical variation identified by SVAseq.</p><p id="Par67">We tested for differential expression using the following R packages:</p><p id="Par68">(i) limma<sup><xref ref-type="bibr" rid="CR59">59</xref></sup> (with lmFit and eBayes), after applying limma-voom<sup><xref ref-type="bibr" rid="CR63">63</xref></sup> to remove heteroscedasticity;</p><p id="Par69">(ii) DESeq2<sup><xref ref-type="bibr" rid="CR61">61</xref></sup>, separately with and without outlier filtering/replacement (minReplicatesForReplace = Inf, cooksCutoff = FALSE options);</p><p id="Par70">(iii) edgeR<sup><xref ref-type="bibr" rid="CR60">60</xref></sup>, using the likelihood ratio test (glmFit and glmLRT functions), and separately, using the F test (glmQLFit and glmQLFTest functions).</p><p id="Par71">Here and elsewhere, we used Ensembl38p10 to identify genes with uniquely corresponding Ensembl gene ID and gene name (13,737 of 15,249 genes in the RNA data).</p><p id="Par72">In each analysis design and method, we used a 5% FDR threshold to correct for multiple testing. As the final step, we applied a conservative approach and considered a gene ‘significantly differentially expressed' between low-grade and high-grade cartilage if it showed significant differential expression across all analysis designs and testing methods (2557 genes, including 2418 with uniquely corresponding Ensembl gene ID and gene name). As discussed in the Supplementary Note <xref ref-type="supplementary-material" rid="MOESM1">3</xref>, for each analysis design, the number of significant genes at 5% FDR was similar across all tests.</p><p id="Par73">Furthermore, we have performed a sensitivity analysis testing differential expression based on the patients in cohort 4 only and verified that results were highly consistent with the analysis of all patients (see Supplementary Note <xref ref-type="supplementary-material" rid="MOESM1">3</xref>).</p><p id="Par74">We used the pheatmap v1.0.12 function in R to plot the gene expression for the 20 genes with the highest absolute log-fold differences between high-grade and low-grade cartilage (selected from 290 genes with significant and concordant cross-omics differences, see Supplementary Methods).</p></sec><sec id="Sec31"><title>Differential protein abundance between high-grade and low-grade cartilage</title><p id="Par75">We performed differential analysis for 4801 proteins that were measured in ≥30% of patients, applying limma<sup><xref ref-type="bibr" rid="CR59">59</xref></sup> to paired samples from 99 patients. Significance was defined at 5% FDR to correct for multiple testing, yielding 2233 proteins with significant differential abundance (2019 proteins with uniquely corresponding Ensembl gene ID and gene name).</p><p id="Par76">As batch effects in proteomics data can be pervasive<sup><xref ref-type="bibr" rid="CR64">64</xref>,<xref ref-type="bibr" rid="CR65">65</xref></sup>, paired samples from any patient were always assayed in the same 6-plex (cohort 1) or 10-plex (cohorts 2–4). As a sensitivity analysis, we carried out a differential abundance analysis for the proteomics data with explicit adjustment for the plexes, to confirm that adjustment for patient effects was captured between-plex batch effects. For each protein, we calculated the log2 of normalised abundance values plus 1. We then obtained residuals from linear regression of these values on the 13 batches used in the proteomics data. These residuals were quantile normalised. We then used limma to test for differential abundance between low-grade and high-grade cartilage using the quantile normalised residuals, applying a paired design and the same method as in the main analysis. The results were highly similar to the main analysis (see Supplementary Note <xref ref-type="supplementary-material" rid="MOESM1">4</xref>).</p><p id="Par77">We used the pheatmap v1.0.12 function in R to plot the abundance of the 20 proteins with the highest absolute log-fold differences between high-grade and low-grade cartilage (selected from 290 genes with significant and concordant cross-omics differences, see Supplementary Methods).</p></sec><sec id="Sec32"><title>Replication of differences between high-grade and low-grade cartilage</title><p id="Par78">We used an independent dataset from the RAAK study (35 osteoarthritis patients, of which 28 knee, 7 hip osteoarthritis) to replicate the molecular differences between high-grade and low-grade cartilage. The RAAK dataset and analysis were taken from a previous publication<sup><xref ref-type="bibr" rid="CR15">15</xref></sup>. Briefly, mRNA samples were sequenced on Illumina HiSeq 2000/4000 (paired-end 2 × 100 bp RNA-sequencing), aligned using GSNAP<sup><xref ref-type="bibr" rid="CR66">66</xref></sup>, and quantified using HTSeq<sup><xref ref-type="bibr" rid="CR67">67</xref></sup>. We used the differential expression analysis as originally published<sup><xref ref-type="bibr" rid="CR15">15</xref></sup>, based on DESeq2, with removal of batch effects using the function removeBatchEffect from the limma R package, then application of a general linear model assuming a negative binomial distribution and a paired Wald-test between preserved and lesioned OA cartilage samples. The data were downloaded from the Github repository <ext-link xlink:href="https://git.lumc.nl/rcoutinhodealmeida/miRNAmRNA" ext-link-type="url">https://git.lumc.nl/rcoutinhodealmeida/miRNAmRNA</ext-link> on 20 June 2020. The analysis as published included 20,165 genes, of which 2387 were differentially expressed at 5% FDR. Of these genes, 12,663 had uniquely mapping Ensembl gene ID and gene name and were also assayed in our data (including 1830 genes with significant differential expression at 5% FDR in the RAAK study). For the genes with RNA-level, protein-level or cross-omics level differential expression in the discovery analysis above, we calculated the proportion with concordant direction of effect in the RAAK data, and applied Fisher’s test to determine whether this proportion was higher than for other genes (RNA-level: compared to all genes not differentially expressed; protein- and cross-omics level: compared to all genes not differentially expressed, but assayed in proteomics). We also calculated the proportion of genes with <italic>p</italic> &lt; 0.05 and FDR &lt; 5% in the RAAK analysis, and the proportion of these with concordant direction of effect.</p></sec><sec id="Sec33"><title>Pathway associations for differences between high-grade and low-grade cartilage</title><p id="Par79">To identify the biological processes with significant molecular differences between high-grade and low-grade cartilage, we carried out gene set enrichment analyses based on the differential expression on RNA, protein and cross-omics levels. We tested for association of the differentially expressed (DE) genes on RNA and/or protein levels at 5% FDR, with robustness checks using more stringent FDR thresholds (1, 0.5 and 0.1%). We restricted this analysis to genes with unique mapping between Ensembl gene ID and Gene Name in Ensembl38p10, and Ensembl gene ID and Entrez ID in HUGO (<ext-link xlink:href="https://www.genenames.org/" ext-link-type="url">https://www.genenames.org/</ext-link> accessed 05/03/2018; 13,094 genes measured on RNA level, 4390 genes measured on protein level, 4387 genes measured on both RNA and protein levels).</p><p id="Par80">We applied Signalling Pathway Impact Analysis (SPIA)<sup><xref ref-type="bibr" rid="CR68">68</xref></sup> to test for association with KEGG signalling pathways. SPIA combines enrichment <italic>p</italic> values with perturbation impact on the pathway based on log-fold differences of the DE genes; perturbation <italic>p</italic> values are obtained by bootstrapping. Enrichment and perturbation <italic>p</italic> values were combined using a normal inversion method which only gives low <italic>p</italic> values when both over-representation and pathway impact <italic>p</italic> values are low (function option combine = 'norminv'). Significance of pathway association was defined as a threshold of 5% FDR applied to the combined <italic>p</italic> values in each analysis. For the analysis of genes DE on both RNA and protein levels, we carried out tests using the log-fold differences from the RNA data (based on the limma analysis with paired samples and SVAseq covariates), and separately, from the protein data.</p><p id="Par81">We also tested enrichment in Gene Ontology terms using GOseq<sup><xref ref-type="bibr" rid="CR56">56</xref></sup>, separately for genes with higher or lower expression in high-grade compared to low-grade cartilage. We accounted for gene length (pwf function options ‘hg19' and ‘geneSymbol'). Significance was defined as a threshold of 5% FDR in each analysis. The results showed broad agreement with the results of the SPIA analysis (see Supplementary Note <xref ref-type="supplementary-material" rid="MOESM1">5</xref> and Supplementary Data <xref ref-type="supplementary-material" rid="MOESM7">3</xref>).</p></sec><sec id="Sec34"><title>Identification of genes with osteoarthritis GWAS gene-level association</title><p id="Par82">From the recent UK Biobank and arcOGEN GWAS meta-analysis<sup><xref ref-type="bibr" rid="CR8">8</xref></sup>, we obtained the results of a gene-level analysis for each of the four osteoarthritis phenotypes (self-reported plus hospital diagnosed, hospital diagnosed knee or hip, hospital diagnosed knee, hospital diagnosed hip), as described in the GWAS paper. Briefly, this analysis used MAGMA v1.06<sup><xref ref-type="bibr" rid="CR69">69</xref></sup> and was based on the mean SNP log-<italic>p</italic>-value in the gene, accounting for LD.</p><p id="Par83">To calculate the effective number of tests across phenotypes, we calculated the correlation matrix between the gene <italic>p</italic> values for the four osteoarthritis phenotypes, and obtained the eigenvalues of this matrix. The effective number of tests <italic>N</italic><sub>eff</sub> for phenotypes was then calculated as <inline-formula id="IEq1"><alternatives><mml:math><mml:msub><mml:mrow><mml:mi>N</mml:mi></mml:mrow><mml:mrow><mml:mi>e</mml:mi><mml:mi>f</mml:mi><mml:mi>f</mml:mi></mml:mrow></mml:msub><mml:mo>=</mml:mo><mml:mi>N</mml:mi><mml:mo>−</mml:mo><mml:msub><mml:mrow><mml:mo>∑</mml:mo></mml:mrow><mml:mrow><mml:mi>λ</mml:mi></mml:mrow></mml:msub><mml:mi>I</mml:mi><mml:mrow><mml:mo>(</mml:mo><mml:mrow><mml:mi>λ</mml:mi><mml:mo>&gt;</mml:mo><mml:mn>1</mml:mn></mml:mrow><mml:mo>)</mml:mo></mml:mrow><mml:mo>*</mml:mo><mml:mrow><mml:mo>(</mml:mo><mml:mrow><mml:mi>λ</mml:mi><mml:mo>−</mml:mo><mml:mn>1</mml:mn></mml:mrow><mml:mo>)</mml:mo></mml:mrow></mml:math><tex-math id="IEq1_TeX">\documentclass[12pt]{minimal}
				\usepackage{amsmath}
				\usepackage{wasysym}
				\usepackage{amsfonts}
				\usepackage{amssymb}
				\usepackage{amsbsy}
				\usepackage{mathrsfs}
				\usepackage{upgreek}
				\setlength{\oddsidemargin}{-69pt}
				\begin{document}$$N_{eff} = N - \mathop {\sum}\nolimits_\lambda {I(\lambda &gt; 1) \ast (\lambda - 1)}$$\end{document}</tex-math><inline-graphic xlink:href="41467_2021_21593_Article_IEq1.gif"/></alternatives></inline-formula>, where <italic>N</italic> = 4 is the number of phenotypes, and <italic>λ</italic> denotes the eigenvalues. Across the Pearson and Spearman correlation matrices, we obtained <italic>N</italic><sub>eff</sub> &lt; 2.65. With 18,449 genes per phenotype, the significance threshold for gene-level <italic>p</italic> values across genes and phenotypes was thus set as 0.05/(18,449 × 2.65) = 1.02 × 10<sup>−6</sup>.</p><p id="Par84">After accounting for the effective number of tests across phenotypes and genes using a Bonferroni correction, 320 of 18,449 genes showed significant association with at least one phenotype. Of these genes, 238 genes were compared between low-grade and high-grade cartilage on at least one omics level and had uniquely corresponding Ensembl gene ID and gene name.</p></sec><sec id="Sec35"><title>ConnectivityMap analysis</title><p id="Par85">To identify opportunities for drug repurposing, we used ConnectivityMap<sup><xref ref-type="bibr" rid="CR17">17</xref></sup> to identify compounds and perturbagen classes (PCLs) that could possibly reverse the differences identified between high-grade and low-grade cartilage. Using the online interface clue.io (accessed 3 March 2019), we submitted the 148 genes with significantly higher expression on both RNA and protein level to calculate a ‘tau’ connectivity score to gene expression signatures experimentally induced by various perturbations in nine cell lines. A positive tau score indicates similarity between the gene expression signature of a perturbation and the submitted query (i.e. upregulation of the genes with higher expression in high-grade compared to low-grade cartilage). A negative tau score indicates that gene expression signature of a perturbation opposes the submitted query (i.e. downregulation of the genes with higher expression in high-grade compared to low-grade cartilage). Recommended thresholds for further consideration of results are tau of at least 90, or below −90, respectively (<ext-link xlink:href="https://clue.io/connectopedia/connectivity_scores" ext-link-type="url">https://clue.io/connectopedia/connectivity_scores</ext-link>, accessed 3 March 2019). A total of 2837 compound and 171 PCL perturbations were evaluated in clue.io. We shortlisted perturbations where both the summary tau and the median tau across cell lines were higher than 90 or lower than −90 for PCLs, with more conservative thresholds of higher than 95 or lower than −95 for compounds. The clue.io platform also contained perturbation data from 3799 gene knock-down and 2160 over-expression experiments (with 2111 genes in both, i.e. 3848 genes total). These data were used to shortlist genes where both the summary and median tau were higher than 95 or lower than −95.</p></sec><sec id="Sec36"><title>Reporting Summary</title><p id="Par86">Further information on research design is available in the <xref ref-type="supplementary-material" rid="MOESM3">Nature Research Reporting Summary</xref> linked to this article.</p></sec></sec></body><back><ack><title>Acknowledgements</title><p>We thank the study participants who made this work possible by their generous donation of samples. This research was conducted using the UK Biobank Resource under application number 9979. The authors are grateful to Dr. Iris Fischer for helpful edits. We are grateful to the RAAK study<sup><xref ref-type="bibr" rid="CR15">15</xref></sup> for making their data available. This work was funded by the Wellcome Trust (206194). M.J.C. was funded through a Medical Research Council Centre for Integrated Research into Musculoskeletal Ageing grant (148985). R.A.B. and the Human Research Tissue Bank are supported by the NIHR Cambridge Biomedical Research Centre. J.H.D.B. and G.R.W. are funded by a Wellcome Trust Strategic Award (101123), a Wellcome Trust Joint Investigator Award (110140 and 110141) and a European Commission Horizon 2020 Grant (666869, THYRAGE). A.W.M. receives funding from Versus Arthritis; Tissue Engineering and Regenerative Therapies Centre (21156).</p></ack><sec sec-type="author-contribution"><title>Author contributions</title><p>Study design: E.Z., J.M.W. and J.S. Collection of knee samples: M.J.C., R.L.J., D.S., K.S. and J.M.W. Collection of hip samples: R.A.B. and A.W.M. Proteomics assays: T.I.R. and J.S.C. Molecular QTL and co-localisation analyses: L.S. Differential expression analyses: J.S. and L.S. Pathway association and drug repurposing analyses: J.S. Writing—original draft: J.S., L.S., N.B., J.H.D.B., G.R.W., J.M.W. and E.Z. Writing—comments and review: all authors.</p></sec><sec><title>Funding</title><p>Open Access funding enabled and organized by Projekt DEAL.</p></sec><sec sec-type="data-availability"><title>Data availability</title><p>The RNA sequencing data reported in this paper have been deposited to the EGA (accession numbers <ext-link xlink:href="https://ega-archive.org/datasets/EGAD00001005215" ext-link-type="url">EGAD00001005215</ext-link>, <ext-link xlink:href="https://ega-archive.org/datasets/EGAD00001003355" ext-link-type="url">EGAD00001003355</ext-link>, <ext-link xlink:href="https://ega-archive.org/datasets/EGAD00001003354" ext-link-type="url">EGAD00001003354</ext-link>, <ext-link xlink:href="https://ega-archive.org/datasets/EGAD00001001331" ext-link-type="url">EGAD00001001331</ext-link>. The proteomics data reported in this paper have been deposited to PRIDE (accession numbers <ext-link xlink:href="https://www.ebi.ac.uk/pride/archive?keyword=PXD014666" ext-link-type="url">PXD014666</ext-link>, <ext-link xlink:href="https://www.ebi.ac.uk/pride/archive?keyword=PXD006673" ext-link-type="url">PXD006673</ext-link>, <ext-link xlink:href="https://www.ebi.ac.uk/pride/archive?keyword=PXD002014" ext-link-type="url">PXD002014</ext-link>. The genotype data reported in this paper have been deposited to the EGA (accession numbers <ext-link xlink:href="https://ega-archive.org/datasets/EGAD00010001746" ext-link-type="url">EGAD00010001746</ext-link>, <ext-link xlink:href="https://ega-archive.org/datasets/EGAD00010001285" ext-link-type="url">EGAD00010001285</ext-link>, <ext-link xlink:href="https://ega-archive.org/datasets/EGAD00010001292" ext-link-type="url">EGAD00010001292</ext-link>, <ext-link xlink:href="https://ega-archive.org/datasets/EGAD00010000722" ext-link-type="url">EGAD00010000722</ext-link>. Data from the 1000 Genomes Project are publicly available (<ext-link xlink:href="https://www.internationalgenome.org" ext-link-type="url">https://www.internationalgenome.org</ext-link>). We also used publicly available data on osteoarthritis differential gene expression from the RAAK Study (<ext-link xlink:href="https://git.lumc.nl/rcoutinhodealmeida/miRNAmRNA" ext-link-type="url">https://git.lumc.nl/rcoutinhodealmeida/miRNAmRNA</ext-link>, accessed 20 June 2020).</p><p>Further data including the TSS information, all significant molQTLs and full co-localisation results, can be obtained online from <ext-link xlink:href="https://hmgubox.helmholtz-muenchen.de/d/fc1fcf65a6724152b7f9/" ext-link-type="url">https://hmgubox.helmholtz-muenchen.de/d/fc1fcf65a6724152b7f9/</ext-link>. The full molecular QTL data and molecular differences between high-grade and low-grade cartilage are available through the Downloads page of the Musculoskelatal Knowledge Portal (mskkp.org).</p></sec><sec sec-type="data-availability"><title>Code availability</title><p>All software used in this study is available from free repositories or from manufacturers as referenced in the Methods section.</p></sec><sec sec-type="ethics-statement"><sec id="FPar1" sec-type="COI-statement"><title>Competing interests</title><p id="Par87">The authors declare no competing interests.</p></sec></sec><ref-list id="Bib1"><title>References</title><ref-list><ref id="CR1"><label>1.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><collab>GBD 2019 Diseases and Injuries Collaborators.</collab></person-group><article-title xml:lang="en">Global burden of 369 diseases and injuries in 204 countries and territories, 1990–2019: a systematic analysis for the Global Burden of Disease Study 2019</article-title><source>Lancet</source><year>2020</year><volume>396</volume><fpage>1204</fpage><lpage>1222</lpage><pub-id pub-id-type="doi">10.1016/S0140-6736(20)30925-9</pub-id></mixed-citation></ref><ref id="CR2"><label>2.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Murphy</surname><given-names>L</given-names></name><etal/></person-group><article-title xml:lang="en">Lifetime risk of symptomatic knee osteoarthritis</article-title><source>Arthritis Rheum.</source><year>2008</year><volume>59</volume><fpage>1207</fpage><lpage>1213</lpage><pub-id pub-id-type="pmid">18759314</pub-id><pub-id pub-id-type="pmcid">4516049</pub-id><pub-id pub-id-type="doi">10.1002/art.24021</pub-id></mixed-citation></ref><ref id="CR3"><label>3.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Murphy</surname><given-names>LB</given-names></name><etal/></person-group><article-title xml:lang="en">One in four people may develop symptomatic hip osteoarthritis in his or her lifetime</article-title><source>Osteoarthr. Cartil.</source><year>2010</year><volume>18</volume><fpage>1372</fpage><lpage>1379</lpage><pub-id pub-id-type="coi">1:STN:280:DC%2BC3cbltFyjtQ%3D%3D</pub-id><pub-id pub-id-type="pmcid">2998063</pub-id><pub-id pub-id-type="doi">10.1016/j.joca.2010.08.005</pub-id></mixed-citation></ref><ref id="CR4"><label>4.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Murphy</surname><given-names>LB</given-names></name><name><surname>Cisternas</surname><given-names>MG</given-names></name><name><surname>Pasta</surname><given-names>DJ</given-names></name><name><surname>Helmick</surname><given-names>CG</given-names></name><name><surname>Yelin</surname><given-names>EH</given-names></name></person-group><article-title xml:lang="en">Medical expenditures and earnings losses among US dults with arthritis in 2013</article-title><source>Arthritis Care Res.</source><year>2018</year><volume>70</volume><fpage>869</fpage><lpage>876</lpage><pub-id pub-id-type="doi">10.1002/acr.23425</pub-id></mixed-citation></ref><ref id="CR5"><label>5.</label><mixed-citation publication-type="other">Torio, C. M. &amp; Moore, B. J. <italic>Statistical Brief #204. National Inpatient Hospital Costs: The Most Expensive Conditions by Payer, 2013</italic>. (Agency for Healthcare Research and Quality, 2016).</mixed-citation></ref><ref id="CR6"><label>6.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Nüesch</surname><given-names>E</given-names></name><etal/></person-group><article-title xml:lang="en">All cause and disease specific mortality in patients with knee or hip osteoarthritis: population based cohort study</article-title><source>BMJ</source><year>2011</year><volume>342</volume><fpage>d1165</fpage><pub-id pub-id-type="pmid">21385807</pub-id><pub-id pub-id-type="pmcid">3050438</pub-id><pub-id pub-id-type="doi">10.1136/bmj.d1165</pub-id></mixed-citation></ref><ref id="CR7"><label>7.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Spector</surname><given-names>TD</given-names></name><name><surname>MacGregor</surname><given-names>AJ</given-names></name></person-group><article-title xml:lang="en">Risk factors for osteoarthritis: genetics</article-title><source>Osteoarthr. Cartil.</source><year>2004</year><volume>12</volume><fpage>39</fpage><lpage>44</lpage><pub-id pub-id-type="doi">10.1016/j.joca.2003.09.005</pub-id></mixed-citation></ref><ref id="CR8"><label>8.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Tachmazidou</surname><given-names>I</given-names></name><etal/></person-group><article-title xml:lang="en">Identification of new therapeutic targets for osteoarthritis through genome-wide analyses of UK Biobank data</article-title><source>Nat. Genet.</source><year>2019</year><volume>51</volume><fpage>230</fpage><lpage>236</lpage><pub-id pub-id-type="coi">1:CAS:528:DC%2BC1MXmtFGruro%3D</pub-id><pub-id pub-id-type="pmid">30664745</pub-id><pub-id pub-id-type="pmcid">6400267</pub-id><pub-id pub-id-type="doi">10.1038/s41588-018-0327-1</pub-id></mixed-citation></ref><ref id="CR9"><label>9.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>GTEx</surname><given-names>Consortium</given-names></name><etal/></person-group><article-title xml:lang="en">Genetic effects on gene expression across human tissues</article-title><source>Nature</source><year>2017</year><volume>550</volume><fpage>204</fpage><lpage>213</lpage><pub-id pub-id-type="doi">10.1038/nature24277</pub-id></mixed-citation></ref><ref id="CR10"><label>10.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Dunham</surname><given-names>I</given-names></name><etal/></person-group><article-title xml:lang="en">An integrated encyclopedia of DNA elements in the human genome</article-title><source>Nature</source><year>2012</year><volume>489</volume><fpage>57</fpage><lpage>74</lpage><pub-id pub-id-type="bibcode">2012Natur.489...57T</pub-id><pub-id pub-id-type="coi">1:CAS:528:DC%2BC38XhtlGnsbzN</pub-id><pub-id pub-id-type="doi">10.1038/nature11247</pub-id></mixed-citation></ref><ref id="CR11"><label>11.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Kundaje</surname><given-names>A</given-names></name><etal/></person-group><article-title xml:lang="en">Integrative analysis of 111 reference human epigenomes</article-title><source>Nature</source><year>2015</year><volume>518</volume><fpage>317</fpage><lpage>330</lpage><pub-id pub-id-type="coi">1:CAS:528:DC%2BC2MXjtVSktbc%3D</pub-id><pub-id pub-id-type="pmid">25693563</pub-id><pub-id pub-id-type="pmcid">4530010</pub-id><pub-id pub-id-type="doi">10.1038/nature14248</pub-id></mixed-citation></ref><ref id="CR12"><label>12.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Steinberg</surname><given-names>J</given-names></name><etal/></person-group><article-title xml:lang="en">Integrative epigenomics, transcriptomics and proteomics of patient chondrocytes reveal genes and pathways involved in osteoarthritis</article-title><source>Sci. Rep.</source><year>2017</year><volume>7</volume><pub-id pub-id-type="bibcode">2017NatSR...7.8935S</pub-id><pub-id pub-id-type="pmid">28827734</pub-id><pub-id pub-id-type="pmcid">5566454</pub-id><pub-id pub-id-type="doi">10.1038/s41598-017-09335-6</pub-id><pub-id pub-id-type="coi">1:CAS:528:DC%2BC1cXhtlOiur3O</pub-id></mixed-citation></ref><ref id="CR13"><label>13.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Karlsson</surname><given-names>C</given-names></name><etal/></person-group><article-title xml:lang="en">Genome-wide expression profiling reveals new candidate genes associated with osteoarthritis</article-title><source>Osteoarthr. Cartil.</source><year>2010</year><volume>18</volume><fpage>581</fpage><lpage>592</lpage><pub-id pub-id-type="coi">1:STN:280:DC%2BC3c3itF2mtg%3D%3D</pub-id><pub-id pub-id-type="doi">10.1016/j.joca.2009.12.002</pub-id></mixed-citation></ref><ref id="CR14"><label>14.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Ramos</surname><given-names>YFM</given-names></name><etal/></person-group><article-title xml:lang="en">Genes involved in the osteoarthritis process identified through genome wide expression analysis in articular cartilage; the RAAK Study</article-title><source>PLoS ONE</source><year>2014</year><volume>9</volume><fpage>e103056</fpage><pub-id pub-id-type="bibcode">2014PLoSO...9j3056R</pub-id><pub-id pub-id-type="pmid">25054223</pub-id><pub-id pub-id-type="pmcid">4108379</pub-id><pub-id pub-id-type="doi">10.1371/journal.pone.0103056</pub-id><pub-id pub-id-type="coi">1:CAS:528:DC%2BC2cXhs1alt7fI</pub-id></mixed-citation></ref><ref id="CR15"><label>15.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Coutinho de Almeida</surname><given-names>R</given-names></name><etal/></person-group><article-title xml:lang="en">RNA sequencing data integration reveals an miRNA interactome of osteoarthritis cartilage</article-title><source>Ann. Rheum. Dis.</source><year>2019</year><volume>78</volume><fpage>270</fpage><lpage>277</lpage><pub-id pub-id-type="pmid">30504444</pub-id><pub-id pub-id-type="doi">10.1136/annrheumdis-2018-213882</pub-id><pub-id pub-id-type="coi">1:CAS:528:DC%2BC1MXht1Ghu7bP</pub-id></mixed-citation></ref><ref id="CR16"><label>16.</label><mixed-citation publication-type="other">Adzhubei, I., Jordan, D. M. &amp; Sunyaev, S. R. Predicting functional effect of human missense mutations using PolyPhen-2. <italic>Curr. Protoc. Hum. Genet</italic>. <ext-link xlink:href="https://doi.org/10.1002/0471142905.hg0720s76" ext-link-type="doi">https://doi.org/10.1002/0471142905.hg0720s76</ext-link> (2013).</mixed-citation></ref><ref id="CR17"><label>17.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Subramanian</surname><given-names>A</given-names></name><etal/></person-group><article-title xml:lang="en">A next generation connectivity map: L1000 platform and the first 1,000,000 profiles</article-title><source>Cell</source><year>2017</year><volume>171</volume><fpage>1437</fpage><lpage>1452.e1417</lpage><pub-id pub-id-type="coi">1:CAS:528:DC%2BC2sXhvFWmsbrE</pub-id><pub-id pub-id-type="pmid">29195078</pub-id><pub-id pub-id-type="pmcid">5990023</pub-id><pub-id pub-id-type="doi">10.1016/j.cell.2017.10.049</pub-id></mixed-citation></ref><ref id="CR18"><label>18.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Roman-Blas</surname><given-names>JA</given-names></name><name><surname>Castañeda</surname><given-names>S</given-names></name><name><surname>Largo</surname><given-names>R</given-names></name><name><surname>Herrero-Beaumont</surname><given-names>G</given-names></name></person-group><article-title xml:lang="en">Osteoarthritis associated with estrogen deficiency</article-title><source>Arthritis Res. Ther.</source><year>2009</year><volume>11</volume><fpage>241</fpage><pub-id pub-id-type="pmid">19804619</pub-id><pub-id pub-id-type="pmcid">2787275</pub-id><pub-id pub-id-type="doi">10.1186/ar2791</pub-id><pub-id pub-id-type="coi">1:CAS:528:DC%2BD1MXht1egtrnI</pub-id></mixed-citation></ref><ref id="CR19"><label>19.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>de Klerk</surname><given-names>BM</given-names></name><etal/></person-group><article-title xml:lang="en">Limited evidence for a protective effect of unopposed oestrogen therapy for osteoarthritis of the hip: a systematic review</article-title><source>Rheumatol.</source><year>2009</year><volume>48</volume><fpage>104</fpage><lpage>112</lpage><pub-id pub-id-type="doi">10.1093/rheumatology/ken390</pub-id><pub-id pub-id-type="coi">1:CAS:528:DC%2BD1MXosVSmtQ%3D%3D</pub-id></mixed-citation></ref><ref id="CR20"><label>20.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Watt</surname><given-names>FE</given-names></name></person-group><article-title xml:lang="en">Hand osteoarthritis, menopause and menopausal hormone therapy</article-title><source>Maturitas</source><year>2016</year><volume>83</volume><fpage>13</fpage><lpage>18</lpage><pub-id pub-id-type="pmid">26471929</pub-id><pub-id pub-id-type="doi">10.1016/j.maturitas.2015.09.007</pub-id></mixed-citation></ref><ref id="CR21"><label>21.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Bar-Yehuda</surname><given-names>S</given-names></name><etal/></person-group><article-title xml:lang="en">Induction of an antiinflammatory effect and prevention of cartilage damage in rat knee osteoarthritis by CF101 treatment</article-title><source>Arthritis Rheumatol.</source><year>2009</year><volume>60</volume><fpage>3061</fpage><lpage>3071</lpage><pub-id pub-id-type="coi">1:CAS:528:DC%2BD1MXhsFOmtbfI</pub-id><pub-id pub-id-type="doi">10.1002/art.24817</pub-id></mixed-citation></ref><ref id="CR22"><label>22.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Yuan</surname><given-names>Q</given-names></name><name><surname>Sun</surname><given-names>L</given-names></name><name><surname>Li</surname><given-names>J-J</given-names></name><name><surname>An</surname><given-names>C-H</given-names></name></person-group><article-title xml:lang="en">Elevated VEGF levels contribute to the pathogenesis of osteoarthritis</article-title><source>BMC Musculoskelet. Disord.</source><year>2014</year><volume>15</volume><pub-id pub-id-type="pmid">25515407</pub-id><pub-id pub-id-type="pmcid">4391471</pub-id><pub-id pub-id-type="doi">10.1186/1471-2474-15-437</pub-id><pub-id pub-id-type="coi">1:CAS:528:DC%2BC2MXitlCntbw%3D</pub-id></mixed-citation></ref><ref id="CR23"><label>23.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Nagao</surname><given-names>M</given-names></name><etal/></person-group><article-title xml:lang="en">Vascular endothelial growth factor in cartilage development and osteoarthritis</article-title><source>Sci. Rep.</source><year>2017</year><volume>7</volume><pub-id pub-id-type="bibcode">2017NatSR...713027N</pub-id><pub-id pub-id-type="pmid">29026147</pub-id><pub-id pub-id-type="pmcid">5638804</pub-id><pub-id pub-id-type="doi">10.1038/s41598-017-13417-w</pub-id><pub-id pub-id-type="coi">1:CAS:528:DC%2BC1cXhsVOksbfJ</pub-id></mixed-citation></ref><ref id="CR24"><label>24.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Kong</surname><given-names>L</given-names></name><name><surname>Wang</surname><given-names>L</given-names></name><name><surname>Meng</surname><given-names>F</given-names></name><name><surname>Cao</surname><given-names>J</given-names></name><name><surname>Shen</surname><given-names>Y</given-names></name></person-group><article-title xml:lang="en">Association between smoking and risk of knee osteoarthritis: a systematic review and meta-analysis</article-title><source>Osteoarthr. Cartil.</source><year>2017</year><volume>25</volume><fpage>809</fpage><lpage>816</lpage><pub-id pub-id-type="coi">1:STN:280:DC%2BC1c%2FotFWntA%3D%3D</pub-id><pub-id pub-id-type="doi">10.1016/j.joca.2016.12.020</pub-id></mixed-citation></ref><ref id="CR25"><label>25.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Nguyen</surname><given-names>PM</given-names></name><name><surname>Abdirahman</surname><given-names>SM</given-names></name><name><surname>Putoczki</surname><given-names>TL</given-names></name></person-group><article-title xml:lang="en">Emerging roles for Interleukin-11 in disease</article-title><source>Growth Factors</source><year>2019</year><volume>37</volume><fpage>1</fpage><lpage>11</lpage><pub-id pub-id-type="coi">1:CAS:528:DC%2BC1MXhtFSmtbfN</pub-id><pub-id pub-id-type="pmid">31161823</pub-id><pub-id pub-id-type="doi">10.1080/08977194.2019.1620227</pub-id></mixed-citation></ref><ref id="CR26"><label>26.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Corden</surname><given-names>B</given-names></name><name><surname>Adami</surname><given-names>E</given-names></name><name><surname>Sweeney</surname><given-names>M</given-names></name><name><surname>Schafer</surname><given-names>S</given-names></name><name><surname>Cook</surname><given-names>SA</given-names></name></person-group><article-title xml:lang="en">IL-11 in cardiac and renal fibrosis: late to the party but a central player</article-title><source>Br. J. Pharmacol.</source><year>2020</year><volume>177</volume><fpage>1695</fpage><lpage>1708</lpage><pub-id pub-id-type="coi">1:CAS:528:DC%2BB3cXjsFSrtbw%3D</pub-id><pub-id pub-id-type="pmid">32022251</pub-id><pub-id pub-id-type="pmcid">7070163</pub-id><pub-id pub-id-type="doi">10.1111/bph.15013</pub-id></mixed-citation></ref><ref id="CR27"><label>27.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Chou</surname><given-names>CH</given-names></name><etal/></person-group><article-title xml:lang="en">Insights into osteoarthritis progression revealed by analyses of both knee tibiofemoral compartments</article-title><source>Osteoarthr. Cartil.</source><year>2015</year><volume>23</volume><fpage>571</fpage><lpage>580</lpage><pub-id pub-id-type="pmcid">4814163</pub-id><pub-id pub-id-type="doi">10.1016/j.joca.2014.12.020</pub-id><pub-id pub-id-type="pmid">4814163</pub-id></mixed-citation></ref><ref id="CR28"><label>28.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Ruettger</surname><given-names>A</given-names></name><name><surname>Neumann</surname><given-names>S</given-names></name><name><surname>Wiederanders</surname><given-names>B</given-names></name><name><surname>Huber</surname><given-names>R</given-names></name></person-group><article-title xml:lang="en">Comparison of different methods for preparation and characterization of total RNA from cartilage samples to uncover osteoarthritis in vivo</article-title><source>BMC Res. Notes</source><year>2010</year><volume>3</volume><pub-id pub-id-type="pmid">20180968</pub-id><pub-id pub-id-type="pmcid">2841606</pub-id><pub-id pub-id-type="doi">10.1186/1756-0500-3-7</pub-id><pub-id pub-id-type="coi">1:CAS:528:DC%2BC3cXht1egt70%3D</pub-id></mixed-citation></ref><ref id="CR29"><label>29.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Le Bleu</surname><given-names>HK</given-names></name><etal/></person-group><article-title xml:lang="en">Extraction of high-quality RNA from human articular cartilage</article-title><source>Anal. Biochem</source><year>2017</year><volume>518</volume><fpage>134</fpage><lpage>138</lpage><pub-id pub-id-type="pmid">27913164</pub-id><pub-id pub-id-type="doi">10.1016/j.ab.2016.11.018</pub-id><pub-id pub-id-type="coi">1:CAS:528:DC%2BC28XitVKqu7vO</pub-id></mixed-citation></ref><ref id="CR30"><label>30.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Shi</surname><given-names>Y</given-names></name><etal/></person-group><article-title xml:lang="en">A small molecule promotes cartilage extracellular matrix generation and inhibits osteoarthritis development</article-title><source>Nat. Commun.</source><year>2019</year><volume>10</volume><pub-id pub-id-type="bibcode">2019NatCo..10.1914S</pub-id><pub-id pub-id-type="pmid">31015473</pub-id><pub-id pub-id-type="pmcid">6478911</pub-id><pub-id pub-id-type="doi">10.1038/s41467-019-09839-x</pub-id><pub-id pub-id-type="coi">1:CAS:528:DC%2BC1MXovVKlsLw%3D</pub-id></mixed-citation></ref><ref id="CR31"><label>31.</label><mixed-citation publication-type="other">Maldonado, M. &amp; Nam, J. The role of changes in extracellular matrix of cartilage in the presence of inflammation on the pathology of osteoarthritis. <italic>Biomed. Res. Int</italic>. 2013, 284873 (2013).</mixed-citation></ref><ref id="CR32"><label>32.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Mainil-Varlet</surname><given-names>P</given-names></name><etal/></person-group><article-title xml:lang="en">Histological assessment of cartilage repair: a report by the Histology Endpoint Committee of the International Cartilage Repair Society (ICRS)</article-title><source>J. Bone Jt. Surg. Am. 85-A Suppl.</source><year>2003</year><volume>2</volume><fpage>45</fpage><lpage>57</lpage><pub-id pub-id-type="doi">10.2106/00004623-200300002-00007</pub-id></mixed-citation></ref><ref id="CR33"><label>33.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Mankin</surname><given-names>HJ</given-names></name><name><surname>Dorfman</surname><given-names>H</given-names></name><name><surname>Lippiello</surname><given-names>L</given-names></name><name><surname>Zarins</surname><given-names>A</given-names></name></person-group><article-title xml:lang="en">Biochemical and metabolic abnormalities in articular cartilage from osteo-arthritic human hips. II. Correlation of morphology with biochemical and metabolic data</article-title><source>J. Bone Joint Surg. Am.</source><year>1971</year><volume>53</volume><fpage>523</fpage><lpage>537</lpage><pub-id pub-id-type="coi">1:CAS:528:DyaE3MXktFOmtrg%3D</pub-id><pub-id pub-id-type="pmid">5580011</pub-id><pub-id pub-id-type="doi">10.2106/00004623-197153030-00009</pub-id><pub-id pub-id-type="pmcid">5580011</pub-id></mixed-citation></ref><ref id="CR34"><label>34.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Pearson</surname><given-names>RG</given-names></name><name><surname>Kurien</surname><given-names>T</given-names></name><name><surname>Shu</surname><given-names>KS</given-names></name><name><surname>Scammell</surname><given-names>BE</given-names></name></person-group><article-title xml:lang="en">Histopathology grading systems for characterisation of human knee osteoarthritis–reproducibility, variability, reliability, correlation, and validity</article-title><source>Osteoarthr. Cartil.</source><year>2011</year><volume>19</volume><fpage>324</fpage><lpage>331</lpage><pub-id pub-id-type="coi">1:STN:280:DC%2BC3M3hvFymsQ%3D%3D</pub-id><pub-id pub-id-type="doi">10.1016/j.joca.2010.12.005</pub-id></mixed-citation></ref><ref id="CR35"><label>35.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Steinberg</surname><given-names>J</given-names></name><etal/></person-group><article-title xml:lang="en">Widespread epigenomic, transcriptomic and proteomic differences between hip osteophytic and articular chondrocytes in osteoarthritis</article-title><source>Rheumatol.</source><year>2018</year><volume>57</volume><fpage>1481</fpage><lpage>1489</lpage><pub-id pub-id-type="coi">1:CAS:528:DC%2BC1MXhtlCnu7jF</pub-id><pub-id pub-id-type="doi">10.1093/rheumatology/key101</pub-id></mixed-citation></ref><ref id="CR36"><label>36.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Hawtree</surname><given-names>S</given-names></name><name><surname>Muthana</surname><given-names>M</given-names></name><name><surname>Wilkinson</surname><given-names>JM</given-names></name><name><surname>Akil</surname><given-names>M</given-names></name><name><surname>Wilson</surname><given-names>AG</given-names></name></person-group><article-title xml:lang="en">Histone deacetylase 1 regulates tissue destruction in rheumatoid arthritis</article-title><source>Hum. Mol. Genet.</source><year>2015</year><volume>24</volume><fpage>5367</fpage><lpage>5377</lpage><pub-id pub-id-type="coi">1:CAS:528:DC%2BC2MXitVeltbzO</pub-id><pub-id pub-id-type="pmid">26152200</pub-id><pub-id pub-id-type="doi">10.1093/hmg/ddv258</pub-id><pub-id pub-id-type="pmcid">26152200</pub-id></mixed-citation></ref><ref id="CR37"><label>37.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Li</surname><given-names>H</given-names></name><etal/></person-group><article-title xml:lang="en">The Sequence Alignment/Map format and SAMtools</article-title><source>Bioinformatics</source><year>2009</year><volume>25</volume><fpage>2078</fpage><lpage>2079</lpage><pub-id pub-id-type="pmid">2723002</pub-id><pub-id pub-id-type="pmcid">2723002</pub-id><pub-id pub-id-type="doi">10.1093/bioinformatics/btp352</pub-id><pub-id pub-id-type="coi">1:CAS:528:DC%2BD1MXpslertr8%3D</pub-id></mixed-citation></ref><ref id="CR38"><label>38.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Tischler</surname><given-names>G</given-names></name><name><surname>Leonard</surname><given-names>S</given-names></name></person-group><article-title xml:lang="en">biobambam: tools for read pair collation based algorithms on BAM files</article-title><source>Source Code Biol. Med.</source><year>2014</year><volume>9</volume><fpage>13</fpage><lpage>13</lpage><pub-id pub-id-type="pmcid">4075596</pub-id><pub-id pub-id-type="doi">10.1186/1751-0473-9-13</pub-id><pub-id pub-id-type="pmid">4075596</pub-id></mixed-citation></ref><ref id="CR39"><label>39.</label><mixed-citation publication-type="other">Andrews, S. FastQC: a quality control tool for high throughput sequence data <ext-link xlink:href="http://www.bioinformatics.babraham.ac.uk/projects/fastqc" ext-link-type="url">http://www.bioinformatics.babraham.ac.uk/projects/fastqc</ext-link> (2010).</mixed-citation></ref><ref id="CR40"><label>40.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Patro</surname><given-names>R</given-names></name><name><surname>Duggal</surname><given-names>G</given-names></name><name><surname>Love</surname><given-names>MI</given-names></name><name><surname>Irizarry</surname><given-names>RA</given-names></name><name><surname>Kingsford</surname><given-names>C</given-names></name></person-group><article-title xml:lang="en">almon provides fast and bias-aware quantification of transcript expression</article-title><source>Nat. Meth.</source><year>2017</year><volume>14</volume><fpage>417</fpage><lpage>419</lpage><pub-id pub-id-type="coi">1:CAS:528:DC%2BC2sXltVWgtL8%3D</pub-id><pub-id pub-id-type="doi">10.1038/nmeth.4197</pub-id></mixed-citation></ref><ref id="CR41"><label>41.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Soneson</surname><given-names>C</given-names></name><name><surname>Love</surname><given-names>M</given-names></name><name><surname>Robinson</surname><given-names>M</given-names></name></person-group><article-title xml:lang="en">Differential analyses for RNA-seq: transcript-level estimates improve gene-level inferences [version 1; referees: 2 approved]</article-title><source>F1000Res.</source><year>2015</year><volume>4</volume><fpage>1521</fpage><pub-id pub-id-type="pmid">26925227</pub-id><pub-id pub-id-type="doi">10.12688/f1000research.7563.1</pub-id><pub-id pub-id-type="coi">1:CAS:528:DC%2BC1cXmtF2ksr0%3D</pub-id><pub-id pub-id-type="pmcid">26925227</pub-id></mixed-citation></ref><ref id="CR42"><label>42.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Purcell</surname><given-names>S</given-names></name><etal/></person-group><article-title xml:lang="en">PLINK: a tool set for whole-genome association and population-based linkage analyses</article-title><source>Am. J. Hum. Genet.</source><year>2007</year><volume>81</volume><fpage>559</fpage><lpage>575</lpage><pub-id pub-id-type="coi">1:CAS:528:DC%2BD2sXhtVSqurrL</pub-id><pub-id pub-id-type="pmid">1950838</pub-id><pub-id pub-id-type="pmcid">1950838</pub-id><pub-id pub-id-type="doi">10.1086/519795</pub-id></mixed-citation></ref><ref id="CR43"><label>43.</label><mixed-citation publication-type="other">Chang, C. C. et al. Second-generation PLINK: rising to the challenge of larger and richer datasets. GigaScience 4, 7 (2015).</mixed-citation></ref><ref id="CR44"><label>44.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><collab>1000 Genomes Project Consortium.</collab><etal/></person-group><article-title xml:lang="en">A global reference for human genetic variation</article-title><source>Nature</source><year>2015</year><volume>526</volume><fpage>68</fpage><lpage>74</lpage><pub-id pub-id-type="doi">10.1038/nature15393</pub-id><pub-id pub-id-type="coi">1:CAS:528:DC%2BC2MXhs1SitLjO</pub-id></mixed-citation></ref><ref id="CR45"><label>45.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>McCarthy</surname><given-names>S</given-names></name><etal/></person-group><article-title xml:lang="en">A reference panel of 64,976 haplotypes for genotype imputation</article-title><source>Nat. Genet.</source><year>2016</year><volume>48</volume><fpage>1279</fpage><lpage>1283</lpage><pub-id pub-id-type="coi">1:CAS:528:DC%2BC28Xhtlykt7zO</pub-id><pub-id pub-id-type="pmid">27548312</pub-id><pub-id pub-id-type="pmcid">5388176</pub-id><pub-id pub-id-type="doi">10.1038/ng.3643</pub-id></mixed-citation></ref><ref id="CR46"><label>46.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Das</surname><given-names>S</given-names></name><etal/></person-group><article-title xml:lang="en">Next-generation genotype imputation service and methods</article-title><source>Nat. Genet.</source><year>2016</year><volume>48</volume><fpage>1284</fpage><lpage>1287</lpage><pub-id pub-id-type="coi">1:CAS:528:DC%2BC28XhsVWksL%2FK</pub-id><pub-id pub-id-type="pmid">27571263</pub-id><pub-id pub-id-type="pmcid">5157836</pub-id><pub-id pub-id-type="doi">10.1038/ng.3656</pub-id></mixed-citation></ref><ref id="CR47"><label>47.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Brown</surname><given-names>AA</given-names></name><etal/></person-group><article-title xml:lang="en">Predicting causal variants affecting expression by using whole-genome sequencing and RNA-seq from multiple human tissues</article-title><source>Nat. Genet.</source><year>2017</year><volume>49</volume><fpage>1747</fpage><lpage>1751</lpage><pub-id pub-id-type="coi">1:CAS:528:DC%2BC2sXhslehtrjM</pub-id><pub-id pub-id-type="pmid">29058714</pub-id><pub-id pub-id-type="doi">10.1038/ng.3979</pub-id><pub-id pub-id-type="pmcid">29058714</pub-id></mixed-citation></ref><ref id="CR48"><label>48.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Robinson</surname><given-names>MD</given-names></name><name><surname>Oshlack</surname><given-names>A</given-names></name></person-group><article-title xml:lang="en">A scaling normalization method for differential expression analysis of RNA-seq data</article-title><source>Genome Biol.</source><year>2010</year><volume>11</volume><pub-id pub-id-type="pmid">20196867</pub-id><pub-id pub-id-type="pmcid">2864565</pub-id><pub-id pub-id-type="doi">10.1186/gb-2010-11-3-r25</pub-id><pub-id pub-id-type="coi">1:CAS:528:DC%2BC3cXkvVKjtro%3D</pub-id></mixed-citation></ref><ref id="CR49"><label>49.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Robinson</surname><given-names>MD</given-names></name><name><surname>McCarthy</surname><given-names>DJ</given-names></name><name><surname>Smyth</surname><given-names>GK</given-names></name></person-group><article-title xml:lang="en">edgeR: a Bioconductor package for differential expression analysis of digital gene expression data</article-title><source>Bioinformatics</source><year>2010</year><volume>26</volume><fpage>139</fpage><lpage>140</lpage><pub-id pub-id-type="coi">1:CAS:528:DC%2BD1MXhs1WlurvO</pub-id><pub-id pub-id-type="doi">10.1093/bioinformatics/btp616</pub-id></mixed-citation></ref><ref id="CR50"><label>50.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Stegle</surname><given-names>O</given-names></name><name><surname>Parts</surname><given-names>L</given-names></name><name><surname>Piipari</surname><given-names>M</given-names></name><name><surname>Winn</surname><given-names>J</given-names></name><name><surname>Durbin</surname><given-names>R</given-names></name></person-group><article-title xml:lang="en">Using probabilistic estimation of expression residuals (PEER) to obtain increased power and interpretability of gene expression analyses</article-title><source>Nat. Protoc.</source><year>2012</year><volume>7</volume><fpage>500</fpage><lpage>507</lpage><pub-id pub-id-type="coi">1:CAS:528:DC%2BC38XivVChu7s%3D</pub-id><pub-id pub-id-type="pmid">22343431</pub-id><pub-id pub-id-type="pmcid">3398141</pub-id><pub-id pub-id-type="doi">10.1038/nprot.2011.457</pub-id></mixed-citation></ref><ref id="CR51"><label>51.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Ongen</surname><given-names>H</given-names></name><name><surname>Buil</surname><given-names>A</given-names></name><name><surname>Brown</surname><given-names>AA</given-names></name><name><surname>Dermitzakis</surname><given-names>ET</given-names></name><name><surname>Delaneau</surname><given-names>O</given-names></name></person-group><article-title xml:lang="en">Fast and efficient QTL mapper for thousands of molecular phenotypes</article-title><source>Bioinformatics</source><year>2016</year><volume>32</volume><fpage>1479</fpage><lpage>1485</lpage><pub-id pub-id-type="coi">1:CAS:528:DC%2BC28XhsVGnt73E</pub-id><pub-id pub-id-type="pmid">26708335</pub-id><pub-id pub-id-type="doi">10.1093/bioinformatics/btv722</pub-id></mixed-citation></ref><ref id="CR52"><label>52.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Storey</surname><given-names>JD</given-names></name><name><surname>Tibshirani</surname><given-names>R</given-names></name></person-group><article-title xml:lang="en">Statistical significance for genomewide studies</article-title><source>Proc. Natl Acad. Sci. USA</source><year>2003</year><volume>100</volume><fpage>9440</fpage><lpage>9445</lpage><pub-id pub-id-type="bibcode">2003PNAS..100.9440S</pub-id><pub-id pub-id-type="amsid">1994856</pub-id><pub-id pub-id-type="coi">1:CAS:528:DC%2BD3sXmtlyktbY%3D</pub-id><pub-id pub-id-type="pmid">12883005</pub-id><pub-id pub-id-type="pmcid">12883005</pub-id><pub-id pub-id-type="zbl">1130.62385</pub-id><pub-id pub-id-type="doi">10.1073/pnas.1530509100</pub-id></mixed-citation></ref><ref id="CR53"><label>53.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Szklarczyk</surname><given-names>D</given-names></name><etal/></person-group><article-title xml:lang="en">STRING v11: protein-protein association networks with increased coverage, supporting functional discovery in genome-wide experimental datasets</article-title><source>Nucleic Acids Res.</source><year>2019</year><volume>47</volume><fpage>D607</fpage><lpage>D613</lpage><pub-id pub-id-type="coi">1:CAS:528:DC%2BC1MXhs1Gqtr%2FM</pub-id><pub-id pub-id-type="pmid">30476243</pub-id><pub-id pub-id-type="doi">10.1093/nar/gky1131</pub-id></mixed-citation></ref><ref id="CR54"><label>54.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Sul</surname><given-names>JH</given-names></name><name><surname>Han</surname><given-names>B</given-names></name><name><surname>Ye</surname><given-names>C</given-names></name><name><surname>Choi</surname><given-names>T</given-names></name><name><surname>Eskin</surname><given-names>E</given-names></name></person-group><article-title xml:lang="en">Effectively identifying eQTLs from multiple tissues by combining mixed model and meta-analytic approaches</article-title><source>PLoS Genet.</source><year>2013</year><volume>9</volume><fpage>e1003491</fpage><pub-id pub-id-type="coi">1:CAS:528:DC%2BC3sXhtFersLjE</pub-id><pub-id pub-id-type="pmid">23785294</pub-id><pub-id pub-id-type="pmcid">3681686</pub-id><pub-id pub-id-type="doi">10.1371/journal.pgen.1003491</pub-id></mixed-citation></ref><ref id="CR55"><label>55.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Han</surname><given-names>B</given-names></name><name><surname>Eskin</surname><given-names>E</given-names></name></person-group><article-title xml:lang="en">Interpreting meta-analyses of genome-wide association studies</article-title><source>PLoS Genet.</source><year>2012</year><volume>8</volume><fpage>e1002555</fpage><pub-id pub-id-type="coi">1:CAS:528:DC%2BC38XjslOltrs%3D</pub-id><pub-id pub-id-type="pmid">22396665</pub-id><pub-id pub-id-type="pmcid">3291559</pub-id><pub-id pub-id-type="doi">10.1371/journal.pgen.1002555</pub-id></mixed-citation></ref><ref id="CR56"><label>56.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Young</surname><given-names>MD</given-names></name><name><surname>Wakefield</surname><given-names>MJ</given-names></name><name><surname>Smyth</surname><given-names>GK</given-names></name><name><surname>Oshlack</surname><given-names>A</given-names></name></person-group><article-title xml:lang="en">Gene ontology analysis for RNA-seq: accounting for selection bias</article-title><source>Genome Biol.</source><year>2010</year><volume>11</volume><pub-id pub-id-type="pmid">20132535</pub-id><pub-id pub-id-type="pmcid">2872874</pub-id><pub-id pub-id-type="doi">10.1186/gb-2010-11-2-r14</pub-id><pub-id pub-id-type="coi">1:CAS:528:DC%2BC3cXjtVKhtr8%3D</pub-id></mixed-citation></ref><ref id="CR57"><label>57.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Giambartolomei</surname><given-names>C</given-names></name><etal/></person-group><article-title xml:lang="en">Bayesian test for colocalisation between pairs of genetic association studies using summary statistics</article-title><source>PLoS Genet.</source><year>2014</year><volume>10</volume><fpage>e1004383</fpage><pub-id pub-id-type="pmid">24830394</pub-id><pub-id pub-id-type="pmcid">24830394</pub-id><pub-id pub-id-type="doi">10.1371/journal.pgen.1004383</pub-id><pub-id pub-id-type="coi">1:CAS:528:DC%2BC2cXhsVGku7rK</pub-id></mixed-citation></ref><ref id="CR58"><label>58.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Conesa</surname><given-names>A</given-names></name><etal/></person-group><article-title xml:lang="en">A survey of best practices for RNA-seq data analysis</article-title><source>Genome Biol.</source><year>2016</year><volume>17</volume><pub-id pub-id-type="pmid">26813401</pub-id><pub-id pub-id-type="pmcid">4728800</pub-id><pub-id pub-id-type="doi">10.1186/s13059-016-0881-8</pub-id><pub-id pub-id-type="coi">1:CAS:528:DC%2BC28XpvVCltrk%3D</pub-id></mixed-citation></ref><ref id="CR59"><label>59.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Ritchie</surname><given-names>ME</given-names></name><etal/></person-group><article-title xml:lang="en">limma powers differential expression analyses for RNA-sequencing and microarray studies</article-title><source>Nucleic Acids Res.</source><year>2015</year><volume>43</volume><fpage>e47</fpage><pub-id pub-id-type="pmid">25605792</pub-id><pub-id pub-id-type="pmcid">25605792</pub-id><pub-id pub-id-type="doi">10.1093/nar/gkv007</pub-id><pub-id pub-id-type="coi">1:CAS:528:DC%2BC2MXhsFaiu7%2FN</pub-id></mixed-citation></ref><ref id="CR60"><label>60.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>McCarthy</surname><given-names>DJ</given-names></name><name><surname>Chen</surname><given-names>Y</given-names></name><name><surname>Smyth</surname><given-names>GK</given-names></name></person-group><article-title xml:lang="en">Differential expression analysis of multifactor RNA-Seq experiments with respect to biological variation</article-title><source>Nucleic Acids Res.</source><year>2012</year><volume>40</volume><fpage>4288</fpage><lpage>4297</lpage><pub-id pub-id-type="coi">1:CAS:528:DC%2BC38XnsF2ks74%3D</pub-id><pub-id pub-id-type="pmid">22287627</pub-id><pub-id pub-id-type="pmcid">3378882</pub-id><pub-id pub-id-type="doi">10.1093/nar/gks042</pub-id></mixed-citation></ref><ref id="CR61"><label>61.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Love</surname><given-names>MI</given-names></name><name><surname>Huber</surname><given-names>W</given-names></name><name><surname>Anders</surname><given-names>S</given-names></name></person-group><article-title xml:lang="en">Moderated estimation of fold change and dispersion for RNA-seq data with DESeq2</article-title><source>Genome Biol.</source><year>2014</year><volume>15</volume><pub-id pub-id-type="pmid">25516281</pub-id><pub-id pub-id-type="pmcid">25516281</pub-id><pub-id pub-id-type="doi">10.1186/s13059-014-0550-8</pub-id><pub-id pub-id-type="coi">1:CAS:528:DC%2BC2MXjtVCrsL8%3D</pub-id></mixed-citation></ref><ref id="CR62"><label>62.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Leek</surname><given-names>JT</given-names></name></person-group><article-title xml:lang="en">svaseq: removing batch effects and other unwanted noise from sequencing data</article-title><source>Nucleic Acids Res.</source><year>2014</year><volume>42</volume><fpage>e161</fpage><pub-id pub-id-type="pmcid">4245966</pub-id><pub-id pub-id-type="doi">10.1093/nar/gku864</pub-id><pub-id pub-id-type="coi">1:CAS:528:DC%2BC28XhtVKkurbP</pub-id><pub-id pub-id-type="pmid">4245966</pub-id></mixed-citation></ref><ref id="CR63"><label>63.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Law</surname><given-names>CW</given-names></name><name><surname>Chen</surname><given-names>Y</given-names></name><name><surname>Shi</surname><given-names>W</given-names></name><name><surname>Smyth</surname><given-names>GK</given-names></name></person-group><article-title xml:lang="en">voom: precision weights unlock linear model analysis tools for RNA-seq read counts</article-title><source>Genome Biol.</source><year>2014</year><volume>15</volume><pub-id pub-id-type="pmid">24485249</pub-id><pub-id pub-id-type="pmcid">4053721</pub-id><pub-id pub-id-type="doi">10.1186/gb-2014-15-2-r29</pub-id><pub-id pub-id-type="coi">1:CAS:528:DC%2BC2MXjsVeltbw%3D</pub-id></mixed-citation></ref><ref id="CR64"><label>64.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Gregori</surname><given-names>J</given-names></name><etal/></person-group><article-title xml:lang="en">Batch effects correction improves the sensitivity of significance tests in spectral counting-based comparative discovery proteomics</article-title><source>J. Proteom.</source><year>2012</year><volume>75</volume><fpage>3938</fpage><lpage>3951</lpage><pub-id pub-id-type="coi">1:CAS:528:DC%2BC38XotVehtLo%3D</pub-id><pub-id pub-id-type="doi">10.1016/j.jprot.2012.05.005</pub-id></mixed-citation></ref><ref id="CR65"><label>65.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Kuligowski</surname><given-names>J</given-names></name><etal/></person-group><article-title xml:lang="en">Detection of batch effects in liquid chromatography-mass spectrometry metabolomic data using guided principal component analysis</article-title><source>Talanta</source><year>2014</year><volume>130</volume><fpage>442</fpage><lpage>448</lpage><pub-id pub-id-type="coi">1:CAS:528:DC%2BC2cXhtlWisrvF</pub-id><pub-id pub-id-type="pmid">25159433</pub-id><pub-id pub-id-type="doi">10.1016/j.talanta.2014.07.031</pub-id></mixed-citation></ref><ref id="CR66"><label>66.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Wu</surname><given-names>TD</given-names></name><name><surname>Watanabe</surname><given-names>CK</given-names></name></person-group><article-title xml:lang="en">GMAP: a genomic mapping and alignment program for mRNA and EST sequences</article-title><source>Bioinformatics</source><year>2005</year><volume>21</volume><fpage>1859</fpage><lpage>1875</lpage><pub-id pub-id-type="coi">1:CAS:528:DC%2BD2MXjsl2ntLw%3D</pub-id><pub-id pub-id-type="pmid">15728110</pub-id><pub-id pub-id-type="doi">10.1093/bioinformatics/bti310</pub-id></mixed-citation></ref><ref id="CR67"><label>67.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Anders</surname><given-names>S</given-names></name><name><surname>Pyl</surname><given-names>PT</given-names></name><name><surname>Huber</surname><given-names>W</given-names></name></person-group><article-title xml:lang="en">HTSeq–a Python framework to work with high-throughput sequencing data</article-title><source>Bioinformatics</source><year>2015</year><volume>31</volume><fpage>166</fpage><lpage>169</lpage><pub-id pub-id-type="coi">1:CAS:528:DC%2BC28Xht1Sjt7vL</pub-id><pub-id pub-id-type="pmid">25260700</pub-id><pub-id pub-id-type="pmcid">25260700</pub-id><pub-id pub-id-type="doi">10.1093/bioinformatics/btu638</pub-id></mixed-citation></ref><ref id="CR68"><label>68.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Tarca</surname><given-names>AL</given-names></name><etal/></person-group><article-title xml:lang="en">A novel signaling pathway impact analysis</article-title><source>Bioinformatics</source><year>2008</year><volume>25</volume><fpage>75</fpage><lpage>82</lpage><pub-id pub-id-type="pmid">18990722</pub-id><pub-id pub-id-type="pmcid">2732297</pub-id><pub-id pub-id-type="doi">10.1093/bioinformatics/btn577</pub-id><pub-id pub-id-type="coi">1:CAS:528:DC%2BD1MXmvFWr</pub-id></mixed-citation></ref><ref id="CR69"><label>69.</label><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>de Leeuw</surname><given-names>CA</given-names></name><name><surname>Mooij</surname><given-names>JM</given-names></name><name><surname>Heskes</surname><given-names>T</given-names></name><name><surname>Posthuma</surname><given-names>D</given-names></name></person-group><article-title xml:lang="en">MAGMA: generalized gene-set analysis of GWAS data</article-title><source>PLoS Comput. Biol.</source><year>2015</year><volume>11</volume><fpage>e1004219</fpage><pub-id pub-id-type="pmid">25885710</pub-id><pub-id pub-id-type="pmcid">4401657</pub-id><pub-id pub-id-type="doi">10.1371/journal.pcbi.1004219</pub-id><pub-id pub-id-type="coi">1:CAS:528:DC%2BC28XktlOguro%3D</pub-id></mixed-citation></ref></ref-list></ref-list><app-group><app id="App1"><sec id="Sec37"><title>Supplementary information</title><p id="Par88"><supplementary-material content-type="local-data" id="MOESM1" xlink:title="Supplementary information"><media xlink:href="MediaObjects/41467_2021_21593_MOESM1_ESM.pdf" mimetype="application" mime-subtype="pdf"><caption xml:lang="en"><p>Supplementary Information</p></caption></media></supplementary-material><supplementary-material content-type="local-data" id="MOESM2" xlink:title="Supplementary information"><media xlink:href="MediaObjects/41467_2021_21593_MOESM2_ESM.pdf" mimetype="application" mime-subtype="pdf"><caption xml:lang="en"><p>Peer Review File</p></caption></media></supplementary-material><supplementary-material content-type="local-data" id="MOESM3" xlink:title="Supplementary information"><media xlink:href="MediaObjects/41467_2021_21593_MOESM3_ESM.pdf" mimetype="application" mime-subtype="pdf"><caption xml:lang="en"><p>Reporting Summary</p></caption></media></supplementary-material><supplementary-material content-type="local-data" id="MOESM4" xlink:title="Supplementary information"><media xlink:href="MediaObjects/41467_2021_21593_MOESM4_ESM.pdf" mimetype="application" mime-subtype="pdf"><caption xml:lang="en"><p>Description of Additional Supplementary Files</p></caption></media></supplementary-material><supplementary-material content-type="local-data" id="MOESM5" xlink:title="Supplementary information"><media xlink:href="MediaObjects/41467_2021_21593_MOESM5_ESM.xlsx" mimetype="application" mime-subtype="vnd.ms-excel"><caption xml:lang="en"><p>Supplementary Data 1</p></caption></media></supplementary-material><supplementary-material content-type="local-data" id="MOESM6" xlink:title="Supplementary information"><media xlink:href="MediaObjects/41467_2021_21593_MOESM6_ESM.xlsx" mimetype="application" mime-subtype="vnd.ms-excel"><caption xml:lang="en"><p>Supplementary Data 2</p></caption></media></supplementary-material><supplementary-material content-type="local-data" id="MOESM7" xlink:title="Supplementary information"><media xlink:href="MediaObjects/41467_2021_21593_MOESM7_ESM.xlsx" mimetype="application" mime-subtype="vnd.ms-excel"><caption xml:lang="en"><p>Supplementary Data 3</p></caption></media></supplementary-material><supplementary-material content-type="local-data" id="MOESM8" xlink:title="Supplementary information"><media xlink:href="MediaObjects/41467_2021_21593_MOESM8_ESM.xlsx" mimetype="application" mime-subtype="vnd.ms-excel"><caption xml:lang="en"><p>Supplementary Data 4</p></caption></media></supplementary-material><supplementary-material content-type="local-data" id="MOESM9" xlink:title="Supplementary information"><media xlink:href="MediaObjects/41467_2021_21593_MOESM9_ESM.xlsx" mimetype="application" mime-subtype="vnd.ms-excel"><caption xml:lang="en"><p>Supplementary Data 5</p></caption></media></supplementary-material><supplementary-material content-type="local-data" id="MOESM10" xlink:title="Supplementary information"><media xlink:href="MediaObjects/41467_2021_21593_MOESM10_ESM.xlsx" mimetype="application" mime-subtype="vnd.ms-excel"><caption xml:lang="en"><p>Supplementary Data 6</p></caption></media></supplementary-material><supplementary-material content-type="local-data" id="MOESM11" xlink:title="Supplementary information"><media xlink:href="MediaObjects/41467_2021_21593_MOESM11_ESM.xlsx" mimetype="application" mime-subtype="vnd.ms-excel"><caption xml:lang="en"><p>Supplementary Data 7</p></caption></media></supplementary-material></p></sec></app></app-group><notes notes-type="ESMHint"><title>Supplementary information</title><p>The online version contains supplementary material available at <ext-link xlink:href="https://doi.org/10.1038/s41467-021-21593-7" ext-link-type="doi">https://doi.org/10.1038/s41467-021-21593-7</ext-link>.</p></notes><notes notes-type="Misc"><p><bold>Peer review information</bold><italic>Nature Communications</italic> thanks Caroline Ospelt and the other anonymous reviewers for their contribution to the peer review of this work. Peer reviewer reports are available.</p></notes><notes notes-type="Misc"><p><bold>Publisher’s note</bold> Springer Nature remains neutral with regard to jurisdictional claims in published maps and institutional affiliations.</p></notes></back></article>