Full text
51,104 characters
· extracted from
preprint-html
· click to expand
Development and validation of a risk prediction algorithm for high-risk populations combining genetic and conventional risk factors of cardiovascular disease | medRxiv /* */ /* */ <!-- <!-- /*! * yepnope1.5.4 * (c) WTFPL, GPLv2 */ (function(a,b,c){function d(a){return"[object Function]"==o.call(a)}function e(a){return"string"==typeof a}function f(){}function g(a){return!a||"loaded"==a||"complete"==a||"uninitialized"==a}function h(){var a=p.shift();q=1,a?a.t?m(function(){("c"==a.t?B.injectCss:B.injectJs)(a.s,0,a.a,a.x,a.e,1)},0):(a(),h()):q=0}function i(a,c,d,e,f,i,j){function k(b){if(!o&&g(l.readyState)&&(u.r=o=1,!q&&h(),l.onload=l.onreadystatechange=null,b)){"img"!=a&&m(function(){t.removeChild(l)},50);for(var d in y[c])y[c].hasOwnProperty(d)&&y[c][d].onload()}}var j=j||B.errorTimeout,l=b.createElement(a),o=0,r=0,u={t:d,s:c,e:f,a:i,x:j};1===y[c]&&(r=1,y[c]=[]),"object"==a?l.data=c:(l.src=c,l.type=a),l.width=l.height="0",l.onerror=l.onload=l.onreadystatechange=function(){k.call(this,r)},p.splice(e,0,u),"img"!=a&&(r||2===y[c]?(t.insertBefore(l,s?null:n),m(k,j)):y[c].push(l))}function j(a,b,c,d,f){return q=0,b=b||"j",e(a)?i("c"==b?v:u,a,b,this.i++,c,d,f):(p.splice(this.i++,0,a),1==p.length&&h()),this}function k(){var a=B;return a.loader={load:j,i:0},a}var l=b.documentElement,m=a.setTimeout,n=b.getElementsByTagName("script")[0],o={}.toString,p=[],q=0,r="MozAppearance"in l.style,s=r&&!!b.createRange().compareNode,t=s?l:n.parentNode,l=a.opera&&"[object Opera]"==o.call(a.opera),l=!!b.attachEvent&&!l,u=r?"object":l?"script":"img",v=l?"script":u,w=Array.isArray||function(a){return"[object Array]"==o.call(a)},x=[],y={},z={timeout:function(a,b){return b.length&&(a.timeout=b[0]),a}},A,B;B=function(a){function b(a){var a=a.split("!"),b=x.length,c=a.pop(),d=a.length,c={url:c,origUrl:c,prefixes:a},e,f,g;for(f=0;f<d;f++)g=a[f].split("="),(e=z[g.shift()])&&(c=e(c,g));for(f=0;f<b;f++)c=x[f](c);return c}function g(a,e,f,g,h){var i=b(a),j=i.autoCallback;i.url.split(".").pop().split("?").shift(),i.bypass||(e&&(e=d(e)?e:e[a]||e[g]||e[a.split("/").pop().split("?")[0]]),i.instead?i.instead(a,e,f,g,h):(y[i.url]?i.noexec=!0:y[i.url]=1,f.load(i.url,i.forceCSS||!i.forceJS&&"css"==i.url.split(".").pop().split("?").shift()?"c":c,i.noexec,i.attrs,i.timeout),(d(e)||d(j))&&f.load(function(){k(),e&&e(i.origUrl,h,g),j&&j(i.origUrl,h,g),y[i.url]=2})))}function h(a,b){function c(a,c){if(a){if(e(a))c||(j=function(){var a=[].slice.call(arguments);k.apply(this,a),l()}),g(a,j,b,0,h);else if(Object(a)===a)for(n in m=function(){var b=0,c;for(c in a)a.hasOwnProperty(c)&&b++;return b}(),a)a.hasOwnProperty(n)&&(!c&&!--m&&(d(j)?j=function(){var a=[].slice.call(arguments);k.apply(this,a),l()}:j[n]=function(a){return function(){var b=[].slice.call(arguments);a&&a.apply(this,b),l()}}(k[n])),g(a[n],j,b,n,h))}else!c&&l()}var h=!!a.test,i=a.load||a.both,j=a.callback||f,k=j,l=a.complete||f,m,n;c(h?a.yep:a.nope,!!i),i&&c(i)}var i,j,l=this.yepnope.loader;if(e(a))g(a,0,l,0);else if(w(a))for(i=0;i (function(w,d,s,l,i){w[l]=w[l]||[];w[l].push({'gtm.start':new Date().getTime(),event:'gtm.js'});var f=d.getElementsByTagName(s)[0];var j=d.createElement(s);var dl=l!='dataLayer'?'&l='+l:'';j.src='//www.googletagmanager.com/gtm.js?id='+i+dl;j.type='text/javascript';j.async=true;f.parentNode.insertBefore(j,f);})(window,document,'script','dataLayer','GTM-P4HH5NV'); Skip to main content Home About Submit ALERTS / RSS Search for this keyword Advanced Search Development and validation of a risk prediction algorithm for high-risk populations combining genetic and conventional risk factors of cardiovascular disease View ORCID Profile Tuuli Puusepp , View ORCID Profile Ave Põld , View ORCID Profile Lili Milani , View ORCID Profile Aet Elken , Estonian Biobank Research Team , View ORCID Profile Mikk Jürisson , View ORCID Profile Krista Fischer doi: https://doi.org/10.1101/2025.04.02.25324383 Tuuli Puusepp 1 Institute of Mathematics and Statistics, University of Tartu 2 Institute of Genomics, University of Tartu Find this author on Google Scholar Find this author on PubMed Search for this author on this site ORCID record for Tuuli Puusepp For correspondence: tuuli.puusepp{at}ut.ee Ave Põld 3 Institute of Family Medicine and Public Health, University of Tartu Find this author on Google Scholar Find this author on PubMed Search for this author on this site ORCID record for Ave Põld Lili Milani 2 Institute of Genomics, University of Tartu Find this author on Google Scholar Find this author on PubMed Search for this author on this site ORCID record for Lili Milani Aet Elken 4 Heart Clinic, University of Tartu 5 Cardiology Centre, North Estonia Medical Centre Find this author on Google Scholar Find this author on PubMed Search for this author on this site ORCID record for Aet Elken 2 Institute of Genomics, University of Tartu Mikk Jürisson 3 Institute of Family Medicine and Public Health, University of Tartu Find this author on Google Scholar Find this author on PubMed Search for this author on this site ORCID record for Mikk Jürisson Krista Fischer 1 Institute of Mathematics and Statistics, University of Tartu 2 Institute of Genomics, University of Tartu Find this author on Google Scholar Find this author on PubMed Search for this author on this site ORCID record for Krista Fischer Abstract Full Text Info/History Metrics Preview PDF Abstract Aim To develop a model for cardiovascular disease (CVD) risk, combining polygenic risk score (PRS) with traditional risk factors while assessing the added value of PRS in two cohorts of biobank participants. Methods Data of 128 209 participants from the Estonian Biobank recruited between 2003– 2011 and 2018–2019 without prevalent cardiovascular disease, was included. Hazard ratios (HR) for polygenic risk versus conventional risk factors were estimated with Cox proportional hazards models, cumulative incidence was assessed with Aalen-Johansen curves. Predictive performance was tested using a split-sample approach and competing risk modelling. Age at CVD event served as the outcome, and the impact of the PRS was evaluated by age group (25–59 vs. 60+), sex, and recruitment period, using HRs, Harrell’s C-index, and net reclassification indices (NRI). Results The estimated HR per one standard deviation (SD) of PRS ranged from 1.1, 95% CI 1.06–1.15 (age 60+, earlier cohort) to 1.36, 95% CI 1.24–1.49 (men 25–59, later cohort). Adding PRS to the conventional risk factors in the age group 25–59 increased the C-statistic by 0.028 (p<0.0001) for men. In the age group 60+, the increase was 0.016 (p=0.0002) across all. In the independent validation set, the continuous NRI was 19.1% (95% CI 13.3%–24.9%) in the 25–59 group and 13.9% (95% CI 8.1%–19.6%) in the 60+ group. Conclusions In a high-risk population, PRS is a strong independent risk factor for CVD and should be considered in routine risk assessment, starting at a relatively young age. Introduction Atherosclerotic cardiovascular diseases (CVD), including coronary artery disease and cerebrovascular disease, are the leading cause of death in numerous European countries [ 1 ]. Polygenic risk scores (PRSs) have demonstrated their effectiveness as a valuable and innovative method for assessing genetic risk related to CVD and improving the accuracy of disease risk prediction[ 2 ]. Several studies indicate that incorporating genetic risk assessment into existing risk stratification algorithms could significantly enhance their efficiency [ 3 ], [ 4 ], [ 5 ]. Polygenic risk scores have been validated in several studies and have been found to enhance CVD risk prediction independently of many traditional factors such as smoking, hypercholesterolemia, hypertension, obesity, and family history of CVD [ 2 ], [ 6 ], [ 7 ] [ 8 ]. Typically, a PRS combines the effect of a large number (from hundreds to millions) of single nucleotide polymorphisms (SNPs) as a weighted sum of allele counts [ 9 ]. It has been shown that elevated polygenic scores contribute to a significantly higher percentage of early-onset myocardial infarction cases than monogenic variants for familial hypercholesterolemia [ 10 ], [ 11 ]. This implies that integrating genetic predisposition complements CVD risk prediction and, when combined with traditional factors, can significantly improve disease risk prediction and facilitate decision making in primary prevention of CVD [ 12 ]. Studies assessing CVD risk combining PRS with clinical and lifestyle data show promising results, yet more rigorous validation and comparisons between existing models are necessary to justify their clinical utility [ 12 ], [ 13 ], [ 14 ], [ 15 ], [ 16 ]. Several large cohort and country-specific prognostic CVD risk models have been developed based on traditional risk factors such as the European SCORE2, the American pooled cohort equations (PCE), and the UK-specific QRISK3 algorithm, yet their generalisability across different populations remains limited [ 17 ], [ 18 ], [ 19 ], [ 20 ], [ 21 ]. A comparable challenge arises when considering the questionable effectiveness of PRS-based risk assessment algorithms that utilize UK Biobank data, given that the United Kingdom has a low prevalence of cardiovascular disease compared to high-risk populations in Middle and Eastern Europe [ 22 ], [ 23 ], [ 24 ]. This study describes the development and validation of the novel risk assessment model combining traditional cardiovascular risk factors with PRSs. METHODS Sources of data Data from the Estonian Biobank (EstBB) was used to compose the study cohort [ 25 ]. EstBB is a volunteer-based biobank that includes genotype and clinical events data of more than 210 000 participants. The health records are regularly updated using national registries, hospital databases and the database of the national health insurance fund which covers data from both primary and secondary care. Additionally, the cholesterol data was quantified using nuclear magnetic resonance (NMR) spectroscopy. Participants have been recruited during two distinct phases: 1) an earlier cohort (2002–2017) including 52 266 participants mainly recruited by general practitioners and 2) a later cohort (2018–2022) including 159 102 participants that joined the biobank during a national campaign [ 25 ]. The PRS used in this study is the multi-ancestry PRS for coronary artery disease developed by Patel, et al. and obtained from the PGS catalog [ 26 ], [ 27 ]. This PRS was selected from a pool of 151 candidate CAD PRSs for its highest z-score in predicting prevalent CVD via logistic regression adjusted for age at recruitment and sex, utilising EstBB data comprising 15 095 CVD cases and 119 694 controls at baseline. Participants Sample size All 185 760 EstBB participants aged at least 25 years at recruitment with genotyping data available were considered for the analysis. After applying the inclusion and exclusion criteria (see below), the total number of individuals included in the study was n=128 209. Exclusion criteria – Prevalent CVD cases i.e. individuals diagnosed with non-fatal CVD (ICD-10 codes I20, I21–I25, I60–I69 excluding I60, I62, I67.1, I67.5, I68.2) before recruitment (n=25 894) – Individuals with diabetes mellitus (E10–E14) at baseline (n=11 555) – Individuals with familial hypercholesterolemia (FH) (n=76) – Individuals with missing lipid values (Total cholesterol (Total-C), HDL cholesterol (HDL-C)) (n=3 569) – Individuals with missing systolic blood pressure (SBP) or with SBP300 mmHg (n=27 212) – Individuals with missing body mass index (BMI) or with BMI less than 15 kg/m2 or more than 50 kg/m2 (n=752) – Individuals with missing smoking data (n=1 591) Outcome The outcome of interest was incident non-fatal or fatal CVD event. Incident CVD events were identified using the atherosclerotic CVD definition provided by the SCORE2 working group [ 17 ]. A detailed list of the ICD-codes is in Table S1. As the outcome data is based on Electronic Health Records (EHR) linkage, we assumed there is no missing outcome data. The data from EHR was available up to December 31 st , 2023. The outcome event was observed in 6 893 individuals and deaths from non-CVD causes (n=2 124) were treated as competing events. Predictors The model combined conventional predictors of CVD and pre-calculated PRS for CAD. Predictor management has been described in the Supplementary file. List of the included predictors: – Age (years) – Sex (M/F) – Current smoking (y/n) – SBP (mmHg) – Total cholesterol (mmol/L) – HDL cholesterol (mmol/L) – BMI (kg/m2) – PRS for CAD Statistical methods a) Impact of PRS on the outcome Aalen-Johansen curves, accounting for competing causes of death and using age as time scale, were estimated to assess how PRS differences impact cumulative CVD event incidence in men and women aged 25–70 at recruitment across three PRS groups (bottom 10%, 10%– 90%, top 10%). Crude hazard ratios with 95% CIs were calculated using the Cox proportional hazards models, separately for the earlier and later cohort. b) PRS effect and discrimination compared with that of conventional risk factors Separate models were fitted for earlier and later cohort and two age groups (25–59 and 60+ years) to estimate the effect of PRS and compare its discrimination with that of conventional risk factors (current smoking, SBP, total cholesterol, HDL cholesterol, BMI). In the younger group, models were sex-specific due to significant interactions, while in the older group, sex was a stratification variable. Cox models used age at CVD event as the time scale to avoid bias due to left truncation. Harrel’s C-index was used to assess discrimination for single-risk-factor models and models with all covariates, with and without PRS. c) Model’s predictive ability in an independent sample An independent dataset was used to assess the model’s predictive ability using a split-sample approach, dividing the cohort into a training set (one-third) and a validation set (two-thirds of the cohort). In each group (defined by age, sex, and recruitment period), two Cox proportional hazard models were fitted in the training datasets with time from recruitment to the incident CVD event as the main outcome of interest. The models accounted for competing risks using the multi-state modelling principle, where other causes of death were considered as competing events [ 28 ]. The event times for participants with no events or with an event occurring later than 5 or 10 years after recruitment were censored at 5 or 10 years for the later and the earlier cohort, respectively. Models were first fitted using traditional CVD risk factors, then with these factors plus PRS. Model calibration was assessed by comparing the number of observed CVD events within quintiles of predicted risk with those predicted from the models. Continuous net reclassification index (NRI) and categorical NRI were computed to compare the 5-year predicted risk with observed event rates. To convert the linear predictors from the survival models adjusted for traditional risk factors and the PRS into absolute 5– or 10-year CVD risk predictions, the absolute risk formula of Benichou and Gail was used (implemented in the riskRegression package version 2023.9.20 of the R software) [ 29 ], [ 30 ]. All analyses were done using R version 4.2.2 [ 31 ]. RESULTS Cohort description and outcomes A total of 128 209 Estonian Biobank participants were included in the analysis. Baseline characteristics of the earlier cohort (n=32 554, recruited in 2002–2017) and later cohort (n=95 655, recruited in 2018–2022) are presented in Table 1 . The mean age at recruitment in both cohorts is similar, approximately 44 years (44.4 in the earlier cohort and 43.7 in the later) and the sex distributions are comparable (about two thirds of the cohorts being women). The median follow-up period is 14.9 years for the earlier cohort and 5.1 years for the later cohort. Regarding risk factors, there is a large difference in smoking prevalence: 30% of the earlier cohort were current smokers at recruitment, compared to 19% in the later cohort. We can also see somewhat higher average blood pressure and LDL cholesterol levels in the earlier cohort among individuals aged 60 and older. The 5-year cumulative CVD incidence differs being 5.7% in the earlier and 0.6% the later cohort. View this table: View inline View popup Table 1. Baseline participant characteristics. Data are mean (SD) unless noted otherwise. Cumulative incidence in different PRS percentiles The cumulative incidence curves on age scale ( Figure 1 ) show higher CVD incidence in men than women across all PRS groups and cohorts. By age 70, about 40% of men in the earlier cohort had experienced CVD, compared to about 25% in the later cohort. For women, the rates were 24% and 13%, respectively. Men in the later cohort had a similar CVD risk to women in the earlier cohort. Download figure Open in new tab Figure 1. Cumulative incidence of CVD with 95% CI in men and women aged 25–70 at recruitment by PRS percentiles and recruitment period using age as time scale. The effect of PRS is consistent across sexes and cohorts. Individuals in the highest PRS decile have significantly higher CVD risk compared to those in the middle or lowest percentiles. For example, in the earlier cohort, men in the highest PRS decile have over double the cumulative incidence of CVD by age 70 compared to those in the lowest decile. Men in the top PRS decile reach a 20% cumulative incidence of CVD six years earlier than average, while those in the lowest decile reach it five years later, creating a decade-long difference between extreme deciles. Similar patterns are observed in women. Compared to individuals in the 10th–90th PRS percentile range, the hazard ratio (HR) for those in the highest PRS decile is 1.7 (95% CI 1.5–1.9) for men and 1.5 (95% CI 1.3–1.7) for women in the earlier cohort, 1.9 (95% CI 1.6–2.4) for men and 1.6 (95% CI 1.3–2.0) for women in the later cohort. When the highest PRS decile is compared to the lowest, HRs range from 1.9 for women in the earlier cohort to 2.9 for men in the later cohort. Effect of PRS in a model with conventional risk factors HRs corresponding to PRS are shown in Figure 2 . These were obtained from Cox models with age as time scale, adjusted for current smoking, SBP, total cholesterol, HDL cholesterol and BMI. Models were fitted separately for the two cohorts, sexes, and age groups (in individuals aged 60+, sex-stratified models were fitted). Download figure Open in new tab Figure 2. HRs (95% CI) for one SD of PRS by sex, age group and recruitment period from Cox models fitted on the entire data using age as time scale and adjusted for all conventional risk factors. The effect of PRS is strongest in younger men, with similar HRs across cohorts. Younger women show slightly lower HRs, with a stronger effect in the later cohort. In individuals aged 60+, the effect is more pronounced in the later cohort, with no significant interaction between PRS and sex in this age group. Detailed parameter estimates are shown in Table S2. Discriminatory power of PRS and conventional risk factors: comparison of C-indices Cox models using age as time scale and stratified by cohort (and additionally by sex in the age group 60+) were fitted with each traditional risk factor and PRS as a single covariate to assess their relative importance. Harrell’s C-index estimates remained below 0.6 ( Figure 3 ), reflecting the model’s discriminative ability for individuals of the same baseline age, as age was used as time scale. Download figure Open in new tab Figure 3. C-indices for individual risk factors, the conventional model, and the PRS added to the conventional model. In younger men, PRS is the strongest single predictor, with a C-index of 0.60 (95% CI 0.58– 0.61), nearly equivalent to that of the model with traditional risk factors. Combining traditional risk factors and PRS leads to a C-index of 0.63 (95% CI 0.61–0.64), representing an improvement of 0.028 (p<0.0001) over the model without PRS. In younger women, the effect of PRS alone (C-index 0.56, 95% CI 0.54–0.57) is not stronger than that of SBP or BMI, but it still leads to a slight improvement in the combined risk factor model (C-index 0.61, 95% CI 0.60–0.63, increase 0.004, p=0.085). In individuals aged 60+, PRS is the strongest single predictor (C-index 0.55, 95% CI 0.53– 0.56). Adding PRS to the traditional risk factor model increases the C-index to 0.57 (95% CI 0.55–0.58, increase 0.016, p=0.0002). In a separate analysis of the two cohorts, there were only minor differences in C-index values across cohorts, with the exception of 60+ age category, where the improvement was clearly higher in the new cohort (adding PRS increased the C-index by 0.008 in the 2002–2017 cohort and by 0.03 in the 2018–2022 cohort; details in Table S3). The performance of model-based predictions in an independent sample To assess the performance of the predictive algorithm on independent data, models including all traditional risk factors, with and without PRS, were refitted using the training set of 42 827 individuals. The parameter estimates were used to calculate the linear predictor values for 85 382 individuals in the validation set. As seen in Figure S1, calibration of the model with PRS is adequate in the validation set when the predicted risk of CVD in groups defined by quintiles of predicted risk is compared to the number of CVD events in these groups. Net reclassification analysis comparing the conventional model to the conventional model with PRS were done separately for the two age groups (25–59 and 60+) but jointly for the two cohorts ( Table 2 ). NRI was calculated according to predicted 5-year risk of CVD. For the categorical NRI analyses, risk categories (low, intermediate, and high) were defined as 5% within 5 years for the younger age group (25–59), and as 10% within 5 years for the older age group (60+) [ 32 ]. A detailed description of the NRI analysis is in the Supplementary file. View this table: View inline View popup Download powerpoint Table 2. Results of NRI analysis. There is a significant improvement in reclassification for both events and non-events in both age groups. In the age group 25–59, the overall NRI was 19.1% (95% CI 13.3%–24.9%). The categorical NRI in this group showed a modest but significant improvement of 3.0% (95% CI 1.2%– 4.8%). In the age group 60+, the overall NRI was slightly lower but still significant at 13.9% (95% CI 8.1%–19.6%), and the categorical NRI in this age group was slightly higher at 3.1% (95% CI 1.1%–5.0%). Detailed results of categorical NRI analysis are in Table S4. Figure 4 displays results of the categorical NRI analysis, focusing on the reclassification of individuals initially classified in the intermediate risk group. In the age group 25–59, the initial intermediate risk category included 19 871 individuals, with a 5-year CVD incidence of 2.6%. After PRS adjustment, 8.1% of these individuals were reclassified to low risk, where the incidence was 0.9%, and 3.7% were moved to high risk, where the incidence increased to 5.8%. In the age group 60+, the initial intermediate risk category included 2 931 individuals, with a 5-year CVD incidence of 8.8%. After PRS incorporation, 15.1% were reclassified to low risk, where the incidence was 5.4%, and 11.4% to high risk, where the incidence was 12.9%. Download figure Open in new tab Figure 4. Reclassification of individuals initially categorized as intermediate risk for 5-year CVD incidence using the conventional model. Arrows indicate the movement of individuals between categories, with corresponding percentages representing the proportion of individuals reclassified. Potential for using PRS-based risk prediction in practice: an illustration The practical use of the proposed risk prediction algorithm involves communication of the effect of unmodifiable genetic component (PRS) alongside potentially modifiable risk factors. We illustrate the message that could be delivered to men and women of age 50. Figure 5 shows the predicted 10-year risk for non-smoking individuals with average PRS, as well as predictions for individuals whose PRS exceeds the mean by two SDs and/or who are current smokers. The values of all other risk factors were fixed at their mean values for men or women aged 25–59 in the 2002–2017 cohort. Download figure Open in new tab Figure 5. Predicted risk of CVD for men and women aged 50 at recruitment by PRS value and current smoking status. The predictions are calculated from the Cox model fitted using the full data of the age group 25–59 recruited in 2002–2017 and using time on study as time scale. The high and average PRS correspond to PRS values of 2 (mean + 2SD) and 0 (mean/median), respectively. The plot indicates that the 10-year CVD risk for a non-smoking man with a high PRS (20.2%, 95% CI 17.0%–23.5%) exceeds that of a smoker with an average PRS (15.6%, 95% CI 14.1%–17.3%). Among women, a current smoker with an average PRS faces a comparable risk (8.4%, 95% CI 7.4%–9.5%) to a non-smoker with a high PRS (8.8%, 95% CI 7.6%–10.2%). Communicating this information to healthcare professionals enhances their understanding of the importance of PRS and supports patient discussions, encouraging high-risk individuals to adjust their behaviour while considering the individual differences in genetic predisposition. DISCUSSION To the best of our knowledge this is the first tailored model for a high CVD-risk population combining polygenic and traditional risk factors for estimating cardiovascular disease risk. This study found a significant improvement in risk discrimination for both men and women when incorporating the CAD PRS to the prediction model. Our model used Estonian Biobank data and demonstrated a significant rise in the risk of CVD events and/or mortality with higher CAD PRS. Moreover, individuals with high PRS may experience the onset of the disease up to a decade earlier compared to those with medium or low PRS. This suggests that a high PRS (top 10%) constitutes a substantial risk factor, comparable to that of smoking or high cholesterol, with elevated risks apparent as early as in one’s 30s. Our study demonstrates that the PRS-inclusive risk model performed best in younger cohorts, which is crucial given the pressing need to develop high-quality risk assessment tools for younger populations, especially since existing models like SCORE-2 can only be used starting at the age of 40 [ 17 ]. Similar results have also been shown in previous PRS population-based models created with Finnish and UK Biobank data [ 6 ], [ 22 ], [ 33 ]. We constructed models utilizing two distinct cohorts from different time periods, yet the significance and consistency of the PRS effect remained consistent in both, underscoring its importance in clinical risk assessment, as also reported by Patel, et al [ 27 ]. The Estonian population is categorised as high risk for cardiovascular disease based on standardised cardiovascular disease mortality rates, thus serving as a proxy to multiple high-risk populations in Eastern Europe [ 17 ]. The baseline risk differs between the two Biobank cohorts, with the later cohort having lower levels of conventional risk factors resulting in lower CVD incidence rates compared to the earlier cohort (reflecting contemporary declining trends in CVD incidence in developed countries).Regardless of cohort differences, the inclusion of PRS in the model demonstrates its consistent effect on overall risk across both cohorts, highlighting the model’s robustness. As public health advancements and behavioural changes reduce the impact of traditional risk factors on disease risk, the relative importance of PRS grows [ 34 ]. Consequently, models like ours can help identify high-risk individuals within the population and guide targeted interventions. The current European Society of Cardiology (ESC) Cardiovascular Disease Prevention Guidelines do not advocate for the routine collection of genetic data in primary prevention [ 32 ]. However, a recent clinical consensus highlights the critical need to quantify the potential benefits of PRS in clinical practice [ 35 ]. In European countries, polygenic risk scores are not yet integrated into routinely collected administrative health data for risk prediction at either the population or individual level. This practice should be reassessed as the global focus on prevention and health promotion increases, highlighting the growing importance of genetics and polygenic risk in everyday healthcare. Clinicians often rely on risk prediction models like SCORE2 and QRISK3 for cardiovascular risk management, but the potential of genetic testing in primary prevention requires more rigorous evaluation. Before integrating PRS into clinical practice, the effectiveness of PRS-based primary CVD prevention strategies, such as lipid-lowering treatments, must be validated through randomized clinical trials. As our study has revealed that individuals with very high CAD PRS face a significantly elevated risk of CVD, we argue that such individuals should be directed toward more intensive primary preventive measures at an early age [ 36 ]. Strengths and limitations A key strength of our study is the use of the Estonian Biobank database, which uniquely integrates clinical and genetic data—unlike most biobanks where these datasets are typically separate [ 37 ]. The availability of CAD PRS for nearly all biobank participants allowed us to leverage a large and comprehensive population for model development. This analysis was further enhanced by the high quality of the genetic data utilized to compute the polygenic scores as all SNPs were measured consistently using the same genotyping arrays. Despite the EstBB cohort being a relatively large sample, a key limitation is the varied follow-up time for some individuals (7.3 years being the average follow-up time) and potential selection effects, due to volunteer-based sampling scheme. In addition, not all predictors had been uniformly measured for every individual, and the demographic composition, including ethnic and racial representation, reflects the Estonian population rather than Europe as a whole. To assure that the best possible model based on conventional predictors is used in both cohorts, accounting for possible differences from a random population-based sample, we did not rely on standardized risk-prediction algorithms such as SCORE2 but instead developed the models that provided the best fit for our data. To maintain simplicity, we chose not to include data on cardiovascular disease (CVD) treatment nor calculate the Charlson Comorbidity Index in the model [ 38 ]. Conclusions The results of this study emphasize the importance of using polygenic risk scores in combination with traditional risk factors for identifying individuals at high risk for atherosclerotic cardiovascular disease. From a primary prevention perspective, polygenic risk scores allow for the early assessment of risk, enabling the implementation of proactive prevention strategies aimed at reducing the burden of cardiovascular disease, particularly in a younger age. Author contributions K.F., L.M. and M.J. contributed to the conception of the study. T.P., K.F. contributed to the acquisition, analysis of the data included in the modelling. A.P and T.P. contributed to the interpretation of the data and drafted the manuscript. A.E., L.M., M.J. supported the interpretation of the data, and provided a scientific evaluation of the study making necessary adjustments. Funding This project has received funding from the European Union’s Horizon Europe research and innovation programme under grant agreement No 101060011. The research was conducted using the Estonian Center of Genomics/Roadmap II funded by the Estonian Research Council (project number TT17). The work of TP and KF was supported by Estonian Research Council (grant PRG1197). The work of AE was supported by Estonian Research Council (grant PRG2078). Data availability The data sets supporting the findings of this study are available in aggregated format upon reasonable request from the corresponding author, ensuring compliance with confidentiality and ethical considerations. Supplementary Materials View this table: View inline View popup Download powerpoint Table S1. Endpoint definitions Management of predictors All continuous predictors were centered and scaled before the analysis: age at 45 and by 5 years, BMI at 25 and by 5 kg/m2, SBP at 125 and by 20 mmHg, total cholesterol at 5 and by 1 mmol/L, HDL cholesterol at 1.5 and by 0.5 mmol/L. For BMI, values below 22 kg/m 2 or above 40 kg/m 2 were truncated at 22 and 40, respectively. For total cholesterol, values below 3 mmol/L or above 9 mmol/L were truncated at 3 and 9, respectively. Analogously, HDL cholesterol values were truncated at 0.7 mmol/L and 2.5 mmol/L. PRS for CAD was standardized to zero–mean and unit-variance. View this table: View inline View popup Table S2. Hazard ratios and p-values from CVD risk models with and without the PRS. The models are derived in the full data of the earlier and later cohort using sex-specific and sex-stratified analysis for age groups 25–59 and 60+, respectively. View this table: View inline View popup Table S3. Model discrimination in age and recruitment groups. The models are derived and the C-indices are calculated using the full data from the earlier and later cohort, with sex-specific and sex-stratified analyses for age groups 25–59 and 60+, respectively. NRI analysis Reclassification analyses comparing the conventional model to the conventional model with PRS were done separately for the two age groups (25–59 and 60+) but jointly for the two cohorts. Models were fitted separately for six subgroups based on cohort, age group, and sex. Specifically, there were separate models for the two cohorts, two age groups (25–59 and 60+), with separate models for men and women in the age group 25–59, and a sex-stratified model for the age group 60+. All models were developed using the training data. NRI was calculated according to predicted 5-year risk of CVD. The predictions were calculated for the individuals in the validation data. Individuals who died of other causes within 5 years were not included in these analyses. Before calculating the NRI, the predictions for men and women in the 25–59 age group were combined, as were the predictions for the two cohorts, so that the NRI analysis included only two groups. For the categorical NRI analyses, risk categories—low, intermediate, and high—were defined as 5% within 5 years for the younger age group (25–59), and as 10% within 5 years for the older age group (60+). The risk thresholds were slightly adjusted from the ones provided in the 2021 ESC Guidelines on cardiovascular disease prevention in clinical practice [ 1 ] to account for the different age groups used in this study and were halved to define 5-year risk thresholds based on the guideline-recommended 10-year CVD risk thresholds. View this table: View inline View popup Table S4. Categorical reclassification of 5-year CVD risk. The conventional model with and without PRS are compared separately in 25–59 and 60+ age group. Download figure Open in new tab Figure S1. Calibration of absolute CVD risk in validation set. The calibration is assessed by sex, age group, and recruitment period. 10– and 5-year CVD risks are used for the earlier and the later cohort, respectively. Acknowledgements The Estonian Biobank Research Team, including Andres Metspalu, Lili Milani, Tõnu Esko, Reedik Mägi, Mait Metspalu, Mari Nelis, and Georgi Hudjashov, contributed to data collection, genotyping, QC, and imputation, while Priit Palta, Nele Taba, Erik Abner, Jaanika Kronberg, and Urmo Võsa contributed to the generation, development, and QC of the NMR data. The activities of the EstBB are regulated by the Human Genes Research Act, which was adopted in 2000 specifically for the operations of the EstBB. Individual level data analysis in the EstBB was carried out under ethical approval 1.1-12/624 from the Estonian Committee on Bioethics and Human Research (Estonian Ministry of Social Affairs), using data from the Estonian Biobank. References [1]. ↵ M. Naghavi et al. , “ Global, regional, and national age-sex specific mortality for 264 causes of death, 1980–2016: a systematic analysis for the Global Burden of Disease Study 2016 ,” The Lancet , vol. 390 , no. 10100 , pp. 1151 – 1210 , Sep. 2017 , doi: 10.1016/S0140-6736(17)32152-9 . OpenUrl CrossRef [2]. ↵ M. E. Weale et al. , “ Validation of an Integrated Risk Tool, Including Polygenic Risk Score, for Atherosclerotic Cardiovascular Disease in Multiple Ethnicities and Ancestries ,” The American Journal of Cardiology , vol. 148 , pp. 157 – 164 , Jun. 2021 , doi: 10.1016/j.amjcard.2021.02.032 . OpenUrl CrossRef [3]. ↵ K. G. Aragam et al. , “ Limitations of Contemporary Guidelines for Managing Patients at High Genetic Risk of Coronary Artery Disease ,” Journal of the American College of Cardiology , vol. 75 , no. 22 , pp. 2769 – 2780 , Jun. 2020 , doi: 10.1016/j.jacc.2020.04.027 . OpenUrl FREE Full Text [4]. ↵ J. I. Rotter and H. J. Lin , “ An Outbreak of Polygenic Scores for Coronary Artery Disease ,” Journal of the American College of Cardiology , vol. 75 , no. 22 , pp. 2781 – 2784, Jun. 2020 , doi: 10.1016/j.jacc.2020.04.054 . OpenUrl FREE Full Text [5]. ↵ Ling Li , Shichao Pang , Fabian Starnecker , Heribert Schunkert , and Bertram Mueller- Myhsok , “ Integration of a polygenic score into guideline-recommended prediction of cardiovascular disease ,” European Heart Journal , vol. ehae048 , Mar. 29, 2024 . Accessed: Apr. 22, 2024. [Online]. Available: https://academic.oup.com/eurheartj/advance-article/doi/10.1093/eurheartj/ehae048/7637416?login=false [6]. ↵ J. Elliott et al. , “ Predictive Accuracy of a Polygenic Risk Score-Enhanced Prediction Model vs a Clinical Risk Score for Coronary Artery Disease .,” JAMA , vol. 323 , no. 7 , pp. 636 – 645 , Feb. 2020 , doi: 10.1001/jama.2019.22241 . OpenUrl CrossRef PubMed [7]. ↵ H. Tada et al. , “ Risk prediction by genetic risk scores for coronary heart disease is independent of self-reported family history ,” European Heart Journal , vol. 37 , no. 6 , pp. 561 – 567 , Feb. 2016 , doi: 10.1093/eurheartj/ehv462 . OpenUrl CrossRef PubMed [8]. ↵ K. Ding , K. R. Bailey , I. J. Kullo , and I. J. Kullo , “ Genotype-informed estimation of risk of coronary heart disease based on genome-wide association data linked to the electronic medical record ,” BMC Cardiovascular Disorders , vol. 11 , no. 1 , pp. 66 – 66 , Nov. 2011 , doi: 10.1186/1471-2261-11-66 . OpenUrl CrossRef PubMed [9]. ↵ J. St-Pierre et al. , “ Considering strategies for SNP selection in genetic and polygenic risk scores ,” Front Genet , vol. 13 , p. 900595 , 2022 , doi: 10.3389/fgene.2022.900595 . OpenUrl CrossRef PubMed [10]. ↵ A. Khera et al. , “ Genome-wide polygenic scores for common diseases identify individuals with risk equivalent to monogenic mutations ,” Nature Genetics , vol. 50 , no. 9 , pp. 1219 – 1224 , Aug. 2018 , doi: 10.1038/s41588-018-0183-z . OpenUrl CrossRef [11]. ↵ S. Thériault , R. Lali , M. Chong , J. L. Velianou , M. K. Natarajan , and G. Paré , “ Polygenic Contribution in Individuals With Early-Onset Coronary Artery Disease .,” Circulation-cardiovascular Genetics , vol. 11 , no. 1 , Jan. 2018 , doi: 10.1161/circgen.117.001849 . OpenUrl CrossRef [12]. ↵ Monica Isgut et al. , “ Highly elevated polygenic risk scores are better predictors of myocardial infarction risk early in life than later ,” Genome Medicine , vol. 13 , no. 1 , pp. 1 – 16 , 2021 , doi: 10.1186/s13073-021-00828-8 . OpenUrl CrossRef PubMed [13]. ↵ Derek Klarin and P. Natarajan , “ Clinical utility of polygenic risk scores for coronary artery disease .,” Nature Reviews Cardiology , pp. 1 – 11 , Nov. 2021 , doi: 10.1038/s41569-021-00638-w . OpenUrl CrossRef PubMed [14]. ↵ Jack W O’Sullivan et al. , “ Polygenic Risk Scores for Cardiovascular Disease: A Scientific Statement From the American Heart Association ,” Circulation , Jul. 2022 , doi: 10.1161/cir.0000000000001077 . OpenUrl CrossRef [15]. ↵ M. Viigimaa et al. , “ Effectiveness and feasibility of cardiovascular disease personalized prevention on high polygenic risk score subjects: a randomized controlled pilot study ,” European Heart Journal Open , 2022 , doi: 10.1093/ehjopen/oeac079 . OpenUrl CrossRef [16]. ↵ J. Kumuthini et al. , “ The clinical utility of polygenic risk scores in genomic medicine practices: a systematic review ,” Hum Genet , vol. 141 , no. 11 , pp. 1697 – 1704 , Nov. 2022 , doi: 10.1007/s00439-022-02452-x . OpenUrl CrossRef [17]. ↵ S. H. J. Hageman et al. , “ SCORE2 risk prediction algorithms: new models to estimate 10-year risk of cardiovascular disease in Europe ,” European Heart Journal , vol. 42 , no. 25 , pp. 2439 – 2454 , Jun. 2021 , doi: 10.1093/eurheartj/ehab309 . OpenUrl CrossRef [18]. ↵ David C. Goff et al. , “ 2013 ACC/AHA Guideline on the Assessment of Cardiovascular Risk: A Report of the American College of Cardiology/American Heart Association Task Force on Practice Guidelines,” Journal of the American College of Cardiology , vol. 63 , no. 25 , pp. 2935 – 2959 , Jun. 2014 , doi: 10.1016/j.jacc.2013.11.005 . OpenUrl FREE Full Text [19]. ↵ J. Hippisley-Cox , C. Coupland , and P. Brindle , “ Development and validation of QRISK3 risk prediction algorithms to estimate future risk of cardiovascular disease: prospective cohort study ,” BMJ , vol. 357 , May 2017 , doi: 10.1136/bmj.j2099 . OpenUrl Abstract / FREE Full Text [20]. ↵ Areti Sofogianni , N. Stalikas , Christina Antza , and K. Tziomalos , “ Cardiovascular Risk Prediction Models and Scores in the Era of Personalized Medicine ,” Journal of Personalized Medicine , vol. 12 , no. 7 , pp. 1180 – 1180 , Jul. 2022 , doi: 10.3390/jpm12071180 . OpenUrl CrossRef PubMed [21]. ↵ T. Tillmann et al. , “ Development and validation of two SCORE-based cardiovascular risk prediction models for Eastern Europe: a multicohort study ,” European Heart Journal , vol. 41 , no. 35 , pp. 3325 – 3333 , 2020 , doi: 10.1093/eurheartj/ehaa571 . OpenUrl CrossRef PubMed [22]. ↵ L. Sun et al. , “ Polygenic risk scores in cardiovascular risk prediction: A cohort study and modelling analyses ,” PLOS Medicine , vol. 18 , no. 1 , 2021 , doi: 10.1371/journal.pmed.1003498 . OpenUrl CrossRef PubMed [23]. ↵ F. Riveros-Mckay et al. , “ Integrated Polygenic Tool Substantially Enhances Coronary Artery Disease Prediction .,” Circulation Genomic and Precision Medicine , vol. 14 , no. 2 , pp. 192 – 200 , 2021 , doi: 10.1161/circgen.120.003304 . OpenUrl CrossRef [24]. ↵ European Association of Preventive Cardiology , “ European Risk Regions. Based on SCORE2 and SCORE2-OP risk regions .,” European Risk Regions . Accessed: Sep. 09, 2023 . [Online]. Available: https://www.heartscore.org/en_GB/heartscore-europe-risk-regions [25]. ↵ L. Leitsalu et al. , “ Cohort Profile: Estonian Biobank of the Estonian Genome Center, University of Tartu ,” Int. J. Epidemiol ., vol. 44 , no. 4 , pp. 1137 – 1147 , Aug. 2015 , doi: 10.1093/ije/dyt268 . OpenUrl CrossRef PubMed [26]. ↵ S. A. Lambert et al. , “ The Polygenic Score Catalog as an open database for reproducibility and systematic evaluation ,” Nat Genet , vol. 53 , no. 4 , pp. 420 – 425 , Apr. 2021 , doi: 10.1038/s41588-021-00783-5 . OpenUrl CrossRef PubMed [27]. ↵ A. P. Patel et al. , “ A multi-ancestry polygenic risk score improves risk prediction for coronary artery disease ,” Nat Med , vol. 29 , no. 7 , pp. 1793 – 1803 , Jul. 2023 , doi: 10.1038/s41591-023-02429-x . OpenUrl CrossRef [28]. ↵ H. Putter , M. Fiocco , and R. B. Geskus , “ Tutorial in biostatistics: competing risks and multilJstate models ,” Statistics in Medicine , vol. 26 , no. 11 , pp. 2389 – 2430 , May 2007 , doi: 10.1002/sim.2712 . OpenUrl CrossRef PubMed Web of Science [29]. ↵ Gerds T , Ohlendorff J , Ozenne B , riskRegression: Risk Regression Models and Prediction Scores for Survival Analysis with Competing Risks . ( Dec . 21 , 2023 ). [R]. Available: https://CRAN.R-project.org/package=riskRegression . [30]. ↵ J. Benichou and M. H. Gail , “ Estimates of Absolute Cause-Specific Risk in Cohort Studies ,” Biometrics , vol. 46 , no. 3 , p. 813, Sep. 1990 , doi: 10.2307/2532098 . OpenUrl CrossRef PubMed Web of Science [31]. ↵ R Core Team , R. [Online] . Available: https://www.r-project.org/ [32]. ↵ F. L. J. Visseren et al. , “ 2021 ESC Guidelines on cardiovascular disease prevention in clinical practice,” European Heart Journal , vol. 42 , no. 34 , pp. 3227 – 3337 , Sep. 2021 , doi: 10.1093/eurheartj/ehab484 . OpenUrl CrossRef PubMed [33]. ↵ E. Tikkanen et al. , “Genetic Risk Prediction and a 2-Stage Risk Screening Strategy for Coronary Heart Disease,” Arteriosclerosis , Thrombosis and Vascular Biology , Sep. 2013 , doi: 10.1161/atvbaha.112.301120 . OpenUrl CrossRef [34]. ↵ T. Yanes , A. M. McInerney-Leo , M. H. Law , and S. Cummings , “ The emerging field of polygenic risk scores and perspective for use in clinical care ,” Human Molecular Genetics , vol. 29 , no. R2 , pp. R165 – R176 , Oct. 2020 , doi: 10.1093/hmg/ddaa136 . OpenUrl CrossRef PubMed [35]. ↵ H. Schunkert et al. , “Clinical utility and implementation of polygenic risk scores for predicting cardiovascular disease,” European Heart Journal , p. ehae649 , Feb. 2025 , doi: 10.1093/eurheartj/ehae649 . OpenUrl CrossRef [36]. ↵ N. A. Marston et al. , “ Predictive Utility of a Coronary Artery Disease Polygenic Risk Score in Primary Prevention ,” JAMA Cardiol , vol. 8 , no. 2 , p. 130 , Feb. 2023 , doi: 10.1001/jamacardio.2022.4466 . OpenUrl CrossRef PubMed [37]. ↵ Ehealth – for continuity of care: proceedings of mie2014 . in Studies in health technology and informatics , no. v. 205. Washington, DC : IOS Press , 2014 . [38]. ↵ M. E. Charlson , P. Pompei , K. L. Ales , and C. R. MacKenzie , “ A new method of classifying prognostic comorbidity in longitudinal studies: Development and validation ,” Journal of Chronic Diseases , vol. 40 , no. 5 , pp. 373 – 383 , Jan. 1987 , doi: 10.1016/0021-9681(87)90171-8 . OpenUrl CrossRef PubMed Web of Science [39]. J. J. Stuart et al. , “ Cardiovascular Risk Factors Mediate the Long-Term Maternal Risk Associated With Hypertensive Disorders of Pregnancy ,” Journal of the American College of Cardiology , vol. 79 , no. 19 , pp. 1901 – 1913 , May 2022 , doi: 10.1016/j.jacc.2022.03.335 . OpenUrl CrossRef PubMed [40]. Y.-X. Wang et al. , “ Hypertensive Disorders of Pregnancy and Subsequent Risk of Premature Mortality ,” Journal of the American College of Cardiology , vol. 77 , no. 10 , pp. 1302 – 1312 , Mar. 2021 , doi: 10.1016/j.jacc.2021.01.018 . OpenUrl CrossRef PubMed View the discussion thread. Back to top Previous Next Posted April 04, 2025. Download PDF Email Thank you for your interest in spreading the word about medRxiv. NOTE: Your email address is requested solely to identify you as the sender of this article. Your Email * Your Name * Send To * Enter multiple addresses on separate lines or separate them with commas. You are going to email the following Development and validation of a risk prediction algorithm for high-risk populations combining genetic and conventional risk factors of cardiovascular disease Message Subject (Your Name) has forwarded a page to you from medRxiv Message Body (Your Name) thought you would like to see this page from the medRxiv website. Your Personal Message CAPTCHA This question is for testing whether or not you are a human visitor and to prevent automated spam submissions. Share Development and validation of a risk prediction algorithm for high-risk populations combining genetic and conventional risk factors of cardiovascular disease Tuuli Puusepp , Ave Põld , Lili Milani , Aet Elken , Estonian Biobank Research Team , Mikk Jürisson , Krista Fischer medRxiv 2025.04.02.25324383; doi: https://doi.org/10.1101/2025.04.02.25324383 Share This Article: Copy Citation Tools Development and validation of a risk prediction algorithm for high-risk populations combining genetic and conventional risk factors of cardiovascular disease Tuuli Puusepp , Ave Põld , Lili Milani , Aet Elken , Estonian Biobank Research Team , Mikk Jürisson , Krista Fischer medRxiv 2025.04.02.25324383; doi: https://doi.org/10.1101/2025.04.02.25324383 Citation Manager Formats BibTeX Bookends EasyBib EndNote (tagged) EndNote 8 (xml) Medlars Mendeley Papers RefWorks Tagged Ref Manager RIS Zotero Tweet Widget Facebook Like Google Plus One Subject Area Cardiovascular Medicine Subject Areas All Articles Addiction Medicine (568) Allergy and Immunology (863) Anesthesia (300) Cardiovascular Medicine (4436) Dentistry and Oral Medicine (444) Dermatology (382) Emergency Medicine (608) Endocrinology (including Diabetes Mellitus and Metabolic Disease) (1509) Epidemiology (15229) Forensic Medicine (30) Gastroenterology (1124) Genetic and Genomic Medicine (6600) Geriatric Medicine (668) Health Economics (997) Health Informatics (4538) Health Policy (1368) Health Systems and Quality Improvement (1613) Hematology (542) HIV/AIDS (1264) Infectious Diseases (except HIV/AIDS) (15916) Intensive Care and Critical Care Medicine (1103) Medical Education (623) Medical Ethics (146) Nephrology (667) Neurology (6599) Nursing (346) Nutrition (998) Obstetrics and Gynecology (1144) Occupational and Environmental Health (957) Oncology (3333) Ophthalmology (974) Orthopedics (369) Otolaryngology (420) Pain Medicine (436) Palliative Medicine (130) Pathology (663) Pediatrics (1693) Pharmacology and Therapeutics (691) Primary Care Research (711) Psychiatry and Clinical Psychology (5447) Public and Global Health (9232) Radiology and Imaging (2198) Rehabilitation Medicine and Physical Therapy (1370) Respiratory Medicine (1196) Rheumatology (593) Sexual and Reproductive Health (712) Sports Medicine (530) Surgery (712) Toxicology (99) Transplantation (289) Urology (265) (function(){function c(){var b=a.contentDocument||a.contentWindow.document;if(b){var d=b.createElement('script');d.innerHTML="window.__CF$cv$params={r:'a00e9a032bcb3fe2',t:'MTc3OTY0OTgzOA=='};var a=document.createElement('script');a.src='/cdn-cgi/challenge-platform/scripts/jsd/main.js';document.getElementsByTagName('head')[0].appendChild(a);";b.getElementsByTagName('head')[0].appendChild(d)}}if(document.body){var a=document.createElement('iframe');a.height=1;a.width=1;a.style.position='absolute';a.style.top=0;a.style.left=0;a.style.border='none';a.style.visibility='hidden';document.body.appendChild(a);if('loading'!==document.readyState)c();else if(window.addEventListener)document.addEventListener('DOMContentLoaded',c);else{var e=document.onreadystatechange||function(){};document.onreadystatechange=function(b){e(b);'loading'!==document.readyState&&(document.onreadystatechange=e,c())}}}})();
Text is read by the "Ask this paper" AI Q&A widget below.
Extraction quality varies by source — PMC NXML preserves structure
cleanly, OA-HTML may include some navigation residue, and OA-PDF can
have broken hyphenation. The publisher copy
(via DOI)
is the canonical version.