Creation of gene expression database on preeclampsia-affected human placenta

preprint OA: closed CC-BY-4.0
📄 Open PDF Full text JSON View at publisher
⚙ AI-generated summary by gemini-2.5-flash-lite, 2026-07-15 ⓘ

This study describes a new database of standardized gene expression data from 1000+ preeclampsia-affected human placenta samples, enabling more accurate cross-experiment analysis.

One-sentence paraphrase of the abstract; not a substitute for reading it. No clinical advice. How this works

⚙ AI-generated deep summary by qwen3.7-flash, 2026-09-10 · read from full text ⓘ

This paper describes the creation of a specialized gene expression database for preeclampsia-affected human placentas, integrating 33 datasets from ArrayExpress to standardize metadata and enable integrative analysis. The resulting collection includes over 1000 samples with sufficient clinical information, such as diagnosis and gestational age, to form robust case-control groups for identifying differentially expressed genes. While the study focuses on placental transcriptomics in hypertensive disorders of pregnancy, it does not explicitly discuss endometriosis or adenomyosis; it was included in the corpus via a keyword match in the upstream search index.

Read from the paper's body, not the abstract. Not a substitute for reading the paper. No clinical advice. How this works

Abstract

Publication of gene expression raw data in open access at online resources like NCBI or ArrayExpress made it possible to use these data for cross-experiment integrative analysis and make new insights into biological phenomena. However, most popular of the present online resources are meant to be archives rather than ready for immediate access and interpretation databases. Data uploaded by independent contributors is not standardized and sometimes incomplete and needs further processing before it is ready for the analysis. Hence, the need for a specialized database appears. Given in this article is the description of the database that was created after processing a collection of 33 relevant datasets on pre-eclampsia-affected human placenta. Data processing includes the choice of relevant experiments from ArrayExpress database, the experiment sample attributes standardization according to MeSH term dictionary and Experimental Factor Ontology and the completion of missing data using information from the corresponding articles and authors. A database of more than 1000 samples contains sufficient sample-wise metadata for them to be arranged into relevant case-control groups. Metadata includes information on biological specimen, donor’s diagnosis, gestational age, mode of delivery etc. The average size of these groups will be higher than it is in separate experiments. This will reduce experiment bias and enhance statistical accuracy of the subsequent analysis such as search for differentially expressed genes or inferring gene networks. The article concludes with the guidelines for the microarray experiment metadata uploading for future contributors.
Full text 34,035 characters · extracted from preprint-html · click to expand
Creation of gene expression database on preeclampsia-affected human placenta | bioRxiv /* */ /* */ <!-- <!-- /*! * yepnope1.5.4 * (c) WTFPL, GPLv2 */ (function(a,b,c){function d(a){return"[object Function]"==o.call(a)}function e(a){return"string"==typeof a}function f(){}function g(a){return!a||"loaded"==a||"complete"==a||"uninitialized"==a}function h(){var a=p.shift();q=1,a?a.t?m(function(){("c"==a.t?B.injectCss:B.injectJs)(a.s,0,a.a,a.x,a.e,1)},0):(a(),h()):q=0}function i(a,c,d,e,f,i,j){function k(b){if(!o&&g(l.readyState)&&(u.r=o=1,!q&&h(),l.onload=l.onreadystatechange=null,b)){"img"!=a&&m(function(){t.removeChild(l)},50);for(var d in y[c])y[c].hasOwnProperty(d)&&y[c][d].onload()}}var j=j||B.errorTimeout,l=b.createElement(a),o=0,r=0,u={t:d,s:c,e:f,a:i,x:j};1===y[c]&&(r=1,y[c]=[]),"object"==a?l.data=c:(l.src=c,l.type=a),l.width=l.height="0",l.onerror=l.onload=l.onreadystatechange=function(){k.call(this,r)},p.splice(e,0,u),"img"!=a&&(r||2===y[c]?(t.insertBefore(l,s?null:n),m(k,j)):y[c].push(l))}function j(a,b,c,d,f){return q=0,b=b||"j",e(a)?i("c"==b?v:u,a,b,this.i++,c,d,f):(p.splice(this.i++,0,a),1==p.length&&h()),this}function k(){var a=B;return a.loader={load:j,i:0},a}var l=b.documentElement,m=a.setTimeout,n=b.getElementsByTagName("script")[0],o={}.toString,p=[],q=0,r="MozAppearance"in l.style,s=r&&!!b.createRange().compareNode,t=s?l:n.parentNode,l=a.opera&&"[object Opera]"==o.call(a.opera),l=!!b.attachEvent&&!l,u=r?"object":l?"script":"img",v=l?"script":u,w=Array.isArray||function(a){return"[object Array]"==o.call(a)},x=[],y={},z={timeout:function(a,b){return b.length&&(a.timeout=b[0]),a}},A,B;B=function(a){function b(a){var a=a.split("!"),b=x.length,c=a.pop(),d=a.length,c={url:c,origUrl:c,prefixes:a},e,f,g;for(f=0;f<d;f++)g=a[f].split("="),(e=z[g.shift()])&&(c=e(c,g));for(f=0;f<b;f++)c=x[f](c);return c}function g(a,e,f,g,h){var i=b(a),j=i.autoCallback;i.url.split(".").pop().split("?").shift(),i.bypass||(e&&(e=d(e)?e:e[a]||e[g]||e[a.split("/").pop().split("?")[0]]),i.instead?i.instead(a,e,f,g,h):(y[i.url]?i.noexec=!0:y[i.url]=1,f.load(i.url,i.forceCSS||!i.forceJS&&"css"==i.url.split(".").pop().split("?").shift()?"c":c,i.noexec,i.attrs,i.timeout),(d(e)||d(j))&&f.load(function(){k(),e&&e(i.origUrl,h,g),j&&j(i.origUrl,h,g),y[i.url]=2})))}function h(a,b){function c(a,c){if(a){if(e(a))c||(j=function(){var a=[].slice.call(arguments);k.apply(this,a),l()}),g(a,j,b,0,h);else if(Object(a)===a)for(n in m=function(){var b=0,c;for(c in a)a.hasOwnProperty(c)&&b++;return b}(),a)a.hasOwnProperty(n)&&(!c&&!--m&&(d(j)?j=function(){var a=[].slice.call(arguments);k.apply(this,a),l()}:j[n]=function(a){return function(){var b=[].slice.call(arguments);a&&a.apply(this,b),l()}}(k[n])),g(a[n],j,b,n,h))}else!c&&l()}var h=!!a.test,i=a.load||a.both,j=a.callback||f,k=j,l=a.complete||f,m,n;c(h?a.yep:a.nope,!!i),i&&c(i)}var i,j,l=this.yepnope.loader;if(e(a))g(a,0,l,0);else if(w(a))for(i=0;i (function(w,d,s,l,i){w[l]=w[l]||[];w[l].push({'gtm.start':new Date().getTime(),event:'gtm.js'});var f=d.getElementsByTagName(s)[0];var j=d.createElement(s);var dl=l!='dataLayer'?'&l='+l:'';j.src='//www.googletagmanager.com/gtm.js?id='+i+dl;j.type='text/javascript';j.async=true;f.parentNode.insertBefore(j,f);})(window,document,'script','dataLayer','GTM-M677548'); Skip to main content Home About Submit ALERTS / RSS Search for this keyword Advanced Search New Results Creation of gene expression database on preeclampsia-affected human placenta View ORCID Profile Oleksandr Lykhenko , View ORCID Profile Alina Frolova , Maria Obolenska doi: https://doi.org/10.1101/102012 Oleksandr Lykhenko Find this author on Google Scholar Find this author on PubMed Search for this author on this site ORCID record for Oleksandr Lykhenko For correspondence: o.k.lykhenko{at}imbg.org.ua Alina Frolova Find this author on Google Scholar Find this author on PubMed Search for this author on this site ORCID record for Alina Frolova Maria Obolenska Find this author on Google Scholar Find this author on PubMed Search for this author on this site Abstract Full Text Info/History Metrics Supplementary material Preview PDF Abstract Publication of gene expression raw data in open access at online resources like NCBI or ArrayExpress made it possible to use these data for cross-experiment integrative analysis and make new insights into biological phenomena. However, most popular of the present online resources are meant to be archives rather than ready for immediate access and interpretation databases. Data uploaded by independent contributors is not standardized and sometimes incomplete and needs further processing before it is ready for the analysis. Hence, the need for a specialized database appears. Given in this article is the description of the database that was created after processing a collection of 33 relevant datasets on pre-eclampsia-affected human placenta. Data processing includes the choice of relevant experiments from ArrayExpress database, the experiment sample attributes standardization according to MeSH term dictionary and Experimental Factor Ontology and the completion of missing data using information from the corresponding articles and authors. A database of more than 1000 samples contains sufficient sample-wise metadata for them to be arranged into relevant case-control groups. Metadata includes information on biological specimen, donor’s diagnosis, gestational age, mode of delivery etc. The average size of these groups will be higher than it is in separate experiments. This will reduce experiment bias and enhance statistical accuracy of the subsequent analysis such as search for differentially expressed genes or inferring gene networks. The article concludes with the guidelines for the microarray experiment metadata uploading for future contributors. Background Open gene expression databases have over time acquired tremendous amounts of data [ 9 , 16 ]. It is now possible to make original discoveries just by analyzing these data. An integrative analysis is one of the approaches in this case. Unlike meta-analysis , which is essentially adding up results of different experiments, an integrative analysis implies merging or integration of raw data which, as studies show [ 11 , 12 ], identifies significantly more differentially expressed genes [ 19 ]. Besides, integrative analysis provides larger sample sets and thus increases statistical significance and reduces experimental bias [ 20 ] which makes it most useful in cases when individual experiments’ average sample set is small. To identify which pieces of data are similar enough for the integration to make sense each individual sample must be supplemented with at least minimal data. Here and later these sample clinical and biological information will be called sample metadata (do not confuse with meta-analysis, a method to unite results of different studies). Unfortunately, most popular public databases like NCBI or ArrayExpress not always contain suffcient sample-wise metadata. There are several reasons for that. The main one is that the integration is not a primary goal for these databases. As is stated in [ 17 ] ArrayExpress primary goal is to serve as an archive for microarray data associated with scientific publications and other research. Also, since the data was uploaded by independent contributors the biological sample metadata is not standardized and, consequently, not ready for immediate automated access by a search query. The choice of relevant characteristics of biological sample is also up to a contributor and those characteristics do not always match the needed ones for the data integration. To address these problems investigators develop specialized databases. Here we present a curated collection of publicly available datasets on preeclampsia-affected human placenta. Pre-eclampsia [ 1 ] is a disorder that occurs only during pregnancy and the post-partum period and affects both the mother and the unborn baby. Affecting at least 5-8% of all pregnancies, it is a rapidly progressive condition characterized by high blood pressure and the presence of protein in the urine. Typically, pre-eclampsia occurs after 20 weeks gestation (in the late 2nd or 3rd trimesters or middle to late pregnancy) and up to six weeks postpartum (after delivery), though in rare cases it can occur earlier than 20 weeks. Globally, pre-eclampsia and other hypertensive disorders of pregnancy are a leading cause of maternal and infant illness and death. Although etiology and pathogenesis of pre-eclampsia are still unknown numerous studies (listed in Supplement 1) point at the gene disregulation in placenta to be one of the possible prerequisites for the disease. Besides, there are also known cases when placenta develops without fetus, called molar pregnancy [ 3 ], which are associated with very early-onset pre-eclampsia. This is why we focused our study on gene expression in placenta. Finally, we chose cDNA microarray technology as most popular in pre-eclampsia studies among ones providing information on the whole transcriptome. Our database now contains suffcient metadata for the samples to be united into relevant case-control groups for the subsequent integrative analysis and further search for differentially expressed genes and inferring gene networks. Methods Our software is written in Python language using Django framework for web interface development and Postgres for relational database support. All experiment and sample metadata were automatically extracted from ArrayExpress database via Bioservices which is a Python interface to ArrayExpress. NCBI database was used to supplement the missing data along with the corresponding scientific articles and authors personally. Source code and database backup file are available at GitHub: https://github.com/Sashkow/placenta-preeclampsia Web interface for our database at its current stage can be accessed at: http://194.44.31.241:24173/ Results The database at its current stage contains a total of 32 experiment datasets and more than 1000 placenta samples, about 900 of which contain minimal metadata for relevant cross-experiment study groups to be constructed, which is diagnosis, gestational age and tissue type. A separate group of 11 experiment datasets and 300 samples are in vitro experiments with cell cultures and immortalized cell lines as biological samples. Apart from pre-eclampsia, some related complication were considered including fetal growth retardation, HELLP syndrome and some genetic disorders such as mosaicism and trisomy of chromosome 16. Database structure Figure 1 shows general structure of the database. Rectangles represent tables in database. Tables are linked with one-to-many and many-to-many relations. For example, Experiments table is in many-to-many relation with Microarrays table for each experiment may utilize multiple microarray platfoms and each platform can be used in multiple experiments. On the other hand, Experiments table is in one-to-many relation with Samples table as experiment may contain multiple samples while each sample can be a part of only one experiment. Experiments and Microarrays tables contain HStore field for storing mappings of strings to strings which are experiment/microarray attribute name-value pairs like “accession:E-GEOD-42424” in our case. SampleAttributes table containing all attributes for all samples is here to provide opportunity for different samples to have different attributes and to make new names reverse compatible with the originally downloaded ones. StandardSampleAttributeNames and StabdarSampleAttributeValues tables store standard terms for sample attribute names and values. The field named additional_info is HStore field with information about such as term source and short description. These tables for starndard terms also contain synonyms many-to-many fields with themselves providing lists of synonymous terms. Detailed explanation of the database design process can be found in Supplement 3. Download figure Open in new tab Figure 1: Database entity relation diagram. Download figure Open in new tab Figure 2: Sample core characteristics. Choice of datasets The initial list of 43 relevant datasets was obtained as a result of ArrayExpress search by the following query: “preeclampsia OR pre-eclampsia OR preeclamptic OR pre-eclamptic” with results filtered by organism “Homo sapiens”, experiment type “rna assay”, experiment type “array assay”. E-GEOD-25906 was excluded due to data retrieval failure. E-MTAB-3732 was excluded since it is a compilation of microarray experiments for different diseases taken from publicly available sources. E-GEOD-15787, E-GEOD-22526, E-MEXP-1050 were excluded due to old microarray design or failure to find probe nucleodide sequences for the array. The full list of included and excluded can be found in Supplement 1. Sample attribute names and values standardization Medical Subject Headings (MeSH) terms dictionary [ 4 ] was chosen as a standard for naming biological sample attribute names and values. We also used Ontology Lookup Service (OLS) as a secondary source in cases where no fitting MeSH term was found, since OLS utilizes multiple ontologies and has been recently updated to have more convenient interface than it used to. Furthermore, ArrayExpress itself uses one of OLS’s ontologies, namely Experimental Factor Ontology (EFO), to perform advanced search over genetic experiments and biological samples [ 5 ]. It means that our standardized samples metadata is potentially compatible with ArrayExpress and might be used to improve its search quality regarding datasets we have processed. Here comes the list of standard sample attribute names and values, which are MeSH or other ontology terms, with lists of mappings onto sample attribute names and values originally downloaded from ArrayExpress. The format is the following: Standard Name 1 (list of original names) Standard Value 1(list of original values) Standard Value2(list of original values) … Standard Name 2(list of original names) Standard Value 1(list of original values) Standard Value 2(list of original values) … … Note that mapping is ambiguous meaning that original name may map onto different standard names depending on the context of the experiment. For example, original name “placenta” will map onto standard name “Chorion” if the context of the experiment suggests that. Some standard names do not map on any original name since that information was absent in ArrayExpress data and obtained manually either from the corresponding articles or from the authors personally. Here is a list of mappings for Diagnosis sample attribute provided as an example. The entire list of mappings can be found in Supplement 2. Diagnosis ( diagnosis , subject status , disease , disease state , disease status , genotype , condition , group , phenotype , classification , Disease State , Disease State ) Healthy ( normotensive , normotensive control patient , normal , , healthy , control , Health , preterm labor , Normotensive control pregnancy , preterm ) Pre–Eclampsia (preeclamptic , pre–eclamptic patient , preeclampsia , Pre–Eclampsia , EOPET , PE , Preeclampsia after 20 weeks of gestation to 140/90 mmHg , excretion of 0.3 g protein in 2 urine samples , early on set preeclampsia , late onset preeclampsia , mild preeclampsia ) Pre–Eclampsia and HELLP Pre–Eclampsia and FGR ( EOPET, pre–eclampsia and fgr , pree clampsia=infant SGA) Severe Preeclampsia ( severe preeclampsia , preeclampsia , severe pre–eclampsia ) Fetal Growth Retardation ( IUGR , Fetal growth restriction (FGR) , infant SGA ) Sample metadata completeness While sample metadata from all considered datasets are now standardized , sample metadata remains far from complete. Yet it is already sufficient to tell whether any two of the samples are comparable i.e. whether they can be put to the same case-control group or not. We identified three key attributes that shall be our criteria for comparability: Diagnosis, Gestation Age, Biological Specimen - for only the samples of the same tissue at the same stage of the placental development under similar conditions can be put into the same study group whenever search for differentially expressed genes between control and deviation groups is to be performed. Gestational Age is for some reason often not mentioned directly in sample data. However, it is possible to infer a meaningful category for gestational age by looking, for instance, at whether delivery was performed at term (37 to 41 week of gestation) or prematurely (20 to 37 weeks) or at pregnancy trimester if mentioned. Also, authors often publish gestation age mean and deviation for a study group which as well can be used to determine gestation age category for the samples. We used Gestational Age distribution for samples with known exact value to identify five gestation categories: first trimester, second trimester, early preterm, late preterm, term ( Figure 3 ). Then we used approximate gestational age information for the samples with unknown exact Gestational Age value to put them into one of the five suggested gestation categories. Thus we determined gestation age for all samples to within gestational age category. Download figure Open in new tab Figure 3: Gestation age distribution for samples with known exact value. Five gestation categories may be considered: first trimester, second trimester early preterm (25 to 30 week), late preterm (30 to 37 week), term (37 to 42 week). It is worth noticing that the value of gestational age depends on measurement technique. If it is measured as first day of the woman’s last menstrual cycle it is roughly 2 weeks earlier than if measured since actual conception date [ 6 ]. Although Diagnosis is specified almost everywhere as the majority of the considered experiments are case-control in design, it is not always enough just to know value of this attribute. For example, if the sample is a “control” it can be placenta delivered from natural birth or after c-section, it can be term or preterm birth, medically induced or spontaneous. All these characteristics are to be considered when matching controls to cases during formation of cross-experiment study groups. Likewise, when the sample is “pre-eclampsia” there are also several additional attributes to look at: pre-eclampsia onset, severity, additional complications, formal pre-eclampsia criteria such as blood pressure and urine protein concentration. Biological Specimen is mostly placenta in our study. A biological sample tissue type is specified wherever mentioned to reduce the noise that would be caused by differential expression among cell types inside placenta. Sample attribute coverage summary for the rest of attributes can be seen in Figure 4 . Metadata may be further completed in future either through utilizing existing approximate data as was done here with gestational age or by inferring metadata directly from the expression data as was done, for instance, with fetal sex in [ 14 ]. Download figure Open in new tab Figure 4: Percentage of samples containing the attribute. Most of the considered samples are comparable and ready for integrative analysis at this moment. These biological samples can now constitute case-control groups of larger size than in original datasets. Discussion Numerous attempts to organize gene expression data for cross-experiment analysis have been taken by different investigators. Integrative approach for pre-eclampsia study has been taken in [ 14 ], though no database creation was intended. Authors merged 7 gene expression microarray datasets of case-control studies of preeclampsia-affected placenta and created study groups of 77 and 96 samples for pre-eclampsia and control groups respectively. Authors got interesting insights into the nature of the disease that could not have been made from mere observation of individual studies’ results. Namely, unsupervised clustering based on both sample expression and metadata revealed three clusters of samples. One of the clusters turned out to be largely composed of the healthiest term placentas along with supposedly pre-eclamptic ones. A hypothesis was made that unhealthy placentas of that cluster had been misdiagnosed as pre-eclamptic ones, and were really afflicted with another maternal hypertensive disorder, such as gestational hypertension or chronic hypertension. This is supported by the GSEA and gestational age comparisons, which indicate that cluster 1 is largely composed of the healthiest term placentas in this data set. GENEVESTIGATOR is a search engine for gene expression over a compendium of manually curated datasets [ 2 , 13 , 21 ]. It provides variety of tools addressing the questions of finding conditions under which the genes of interest are the most up- or down-regulated as well as finding the genes that are the most expressed under certain conditions. Biological samples are carefully annotated using custom ontology with a variety of characteristics allowing to filter for highly specific cases. Cross-experiment analysis is based on a concept of meta-profile. As explained in [ 10 ] meta-profiles summarize expression levels from many samples according to their biological context. Each sample is annotated with five attributes: anatomical parts, cell lines, cancers types, developmental stages, and perturbations - to generate meta-profiles. Hence each meta-profile is an expression profile of an “average” sample generated from all expression profiles of the samples of specific tissue or developmental stage. An exception is the Perturbation meta-profile, which consists of responses to various experimental conditions (drugs, chemicals, hormones, etc.), diseases, and genotypes. For Perturbation meta-profile results are created by comparing groups of samples from individual experiments. Data from multiple experiments are not mixed to create a single value. As a result, this tool contains large compendia of response types collected from many experiments. It also worth mentioning that samples of two different microarray platforms can not be considered simultaneously. This implies that GENEVESTIGATOR is closer to meta analysis than to integrative analysis. It integrates sample data for samples of the same platform and biological context to create meta-profiles but other tasks such as comparison of samples of different biological contexts or comparison of differentially expressed genes for different medical or biological cases (perturbations) are pure meta-analysis. The GENEVESTIGATOR’s does not currently contain datasets of our interest and does non allow user to upload data directly and, thus, can not satisfy our goals. Another resource enabling interactive query and navigation of transcriptome datasets is Gene Expression Browser [ 18 ] and its specific implementation for placental gene expression [ 8 ]. While its data is more relevant to our pursues and the interface is quite convenient and has some extended features in comparison to GEO or ArrayExpress, such as search for differentially expressed genes in a single experiment, building relevant plots, giving detailed information on found genes and and a bit of meta-analysis tools such as search for experiments that have given gene differentially expressed with a certain rate of difference, no tools for integrative analysis are provided. An Integrative Meta-Analysis of Expression Data ( INMEX ) service [ 7 ] is the one having options for both meta- and integrative analysis, not to mention a huge ensemble of tools for further data analysis. It is worth emphasizing that INMEX is not a database, but a pure tool with no pre-uploaded data except for several examples: uploading raw data and constructing proper study groups are the user’s responsibility. Opportunely, our newly developed database can make for a perfect data input to INMEX and it is quite possible that INMEX’es output will be a solution to our pursue of data integration and subsequent analysis. Finally, we can not but emphasize how important it is to upload well annotated sample-wise experiment metadata. In doing so we join the call from D.M. Nelson and G.J. Burton, editor of the Placenta journal, who published A technical note to improve the reporting of studies of the human placenta featuring those essential parameters attention to which “will lead to enhanced chances that comparisons will be of apples with apples, instead of comparisons of two fruits”[ 15 ]. While we notice the improvement in time of the data reporting quality we could not manage to gather even the half of what the note suggests from the publicly available datasets and articles ( Figure 4 ). Namely, the information regarding drugs, previous prenatal admissions, screened for diabetes, antibiotics, beta strep status, antenatal steroids, magnesium sulfate, anesthesia and cervical ripening agent was almost entirely absent. Despite that, most of the samples in our database are now provided with sufficient metadata for them to be comparable. These biological samples can now constitute case-control groups of larger size than in individual datasets. Described here data gathering and standardization are the first steps to further analysis, namely, the search for differentially expressed genes between different cross-experiment case and control groups. Also, our database is designed and expected to be enlarged with forthcoming gene expression experiments’ metadata and, potentially, with data from other omics (genome, metabolome) with relatively low effort. Competing interests The authors declare that they have no competing interests. Author’s contributions OL, AF designed the database. AF, MO made suggestions on database content. OL implemented database design and its admin web interface and performed metadata standardization. All authors read and approved the final version of the article. Acknowledgements We would like to express out gratitude to Sandra A. Founds for providing clinical data for E-GEOD-12767. Additional Files Supplement 1 — ArrayExpress Experiments’ Accession Numbers List of ArrayExpress accession numbers for experiments taken for the study. Supplement 2 — Sample Attribute Names and values List of all standardized names and values for sample attributes along with the corresponding names and values originally downloaded from ArrayExpress. References 1. ↵ About preeclampsia: Preeclampsia foundation . http://www.preeclampsia.org/health-information/about-preeclampsia . (Accessed on 12/12/2016 ). 2. ↵ Genevestigator , a high performance search engine for gene expression . https://genevestigator.com/gv/doc/intro_biomed.jsp . 3. ↵ Mesh browser, molar pregnancy , unique id: D006828 . https://meshb.nlm.nih.gov/#/record/ui?ui=D006828 . (Accessed on 12/12/2016 ). 4. ↵ Mesh term browser . https://meshb.nlm.nih.gov/#/fieldSearch . (Accessed on 12/12/2016 ). 26. ↵ Searching in arrayexpress . https://www.ebi.ac.uk/arrayexpress/help/how_to_search.html . (Accessed on 12/06/2016 . 26. ↵ Week 1 - Month 1 Fetal development information on baby growth in pregnancy over weeks months trimesters . http://www.baby2see.com/development/week1.html . 7. ↵ INMEX-a web-based tool for integrative meta-analysis of expression data . Nucleic acids research , 41 : W63 – 70 , jul 2013 . OpenUrl CrossRef PubMed Web of Science 8. ↵ A curated transcriptome dataset collection to investigate the development and differentiation of the human placenta and its associated pathologies[version 2; referees: 2 approved] . F1000Research , 2016 . 9. ↵ Tanya Barrett , Dennis B Troup , Stephen E Wilhite , Pierre Ledoux , Carlos Evangelista , Irene F Kim , Maxim Tomashevsky , Kimberly A Marshall , Katherine H Phillippy , Patti M Sherman , Rolf N Muertter , Michelle Holko , Oluwabukunmi Ayanbule , Andrey Yefanov , and Alexandra Soboleva . NCBI GEO: archive for functional genomics data sets-10 years on . Nucleic acids research , 39 (Database issue): 1005 – 10 , 1 2011 . OpenUrl 26. ↵ Celine Clauzel , Jean-christophe Foltete , Xavier Girardet , and Gilles Vuidel . User Manual. (May) : 0 – 37 , 2016 . 26. ↵ Alina Frolova , Vladyslav Bondarenko , and Maria Obolenska . Comparing alternative pipelines for cross-platform microarray gene expression data integration with RNA-seq data in breast cancer . bioRxiv , 2016 . 26. ↵ Alina Frolova and Maria Obolenska . Integrative approaches for data analysis in systems biology: Current advances. In Applied Physics and Engineering (YSF), 2016 II International Young Scientists Forum on , pages 194 – 198 . IEEE, 2016 . 26. ↵ Tomas Hruz , Oliver Laule , Gabor Szabo , Frans Wessendorp , Stefan Bleuler , Lukas Oertle , Peter Widmayer , Wilhelm Gruissem , and Philip Zimmermann . Genevestigator V3: A Reference Expression Database for the Meta-Analysis of Transcriptomes . Advances in Bioinformatics , 5 , 2008 . 14. ↵ Katherine Leavey , Shannon A Bainbridge , and Brian J Cox . Large scale aggregate microarray analysis reveals three distinct molecular subclasses of human preeclampsia . PloS one , 10 ( 2 ): e0116508 , 1 2015 . OpenUrl CrossRef PubMed 15. ↵ Michael Nelson D. and Graham J. Burton . A technical note to improve the reporting of studies of the human placenta . Placenta , 32 ( 2 ): 195 – 196 , 2011 . OpenUrl 26. ↵ Gabriella Rustici , Kolesnikov Nikolay , Brandizi Marco , Burdett Tony , Miroslaw Dylag , Ibrahim Emam , Anna Farne , Hastings Emma , Ison Jon , Maria Keays , Natalja Kurbatova , Malone James , Mani Roby , Annalisa Mupo , Rui Pedro Pereira , Ekaterina Pilicheva , Johan Rung , Anjan Sharma , Tang Y Amy , Tobias Ternent , Andrew Tikhonov , Danielle Welter , Eleanor Williams , Alvis Brazma , Parkinson Helen , and Sarkans Ugis . ArrayExpress update-trends in database growth and links to data analysis tools . Nucleic acids research , 41 (Database issue): 987 – 90 , 1 2013 . OpenUrl 26. ↵ Sarkans Ugis , Parkinson Helen , Gonzalo Garcia Lara , Ahmet Oezcimen , Anjan Sharma , Niran Abeygunawardena , Sergio Contrino , Ele Holloway , Philippe Rocca-Serra , Gaurab Mukherjee , Mohammadreza Shojatalab , Misha Kapushesky , Susanna Assunta Sansone , Anna Farne , Tim Rayner , and Alvis Brazma . The ArrayExpress gene expression database: A software engineering and implementation perspective . Bioinformatics , 21 ( 8 ): 1495 – 1501 , 2005 . OpenUrl CrossRef PubMed 26. ↵ Cate Speake , Scott Presnell , Kelly Domico , Brad Zeitner , Anna Bjork , David Anderson , Michael J Mason , Elizabeth Whalen , Olivia Vargas , Dimitry Popov , Darawan Rinchai , Noemie Jourde-Chiche , Laurent Chiche , Charlie Quinn , and Damien Chaussabel . An interactive web application for the dissemination of human systems immunology data . Journal of translational medicine , 13 : 196 , 6 2015 . OpenUrl 19. ↵ Jonatan Taminau , Cosmin Lazar , Stijn Meganck , Ann Nowé , Jonatan Taminau , Cosmin Lazar , Stijn Meganck , Ann Nowé Nowé, and Ann. Comparison of Merging and Meta-Analysis as Alternative Approaches for Integrative Gene Expression Analysis . ISRN Bioinformatics , 2014 : 1 – 7 , 2014 . OpenUrl 20. ↵ Christopher J Walsh , Pingzhao Hu , Jane Batt , and Claudia C Dos Santos . Microarray Meta-Analysis and Cross-Platform Normalization: Integrative Genomics for Robust Biomarker Discovery . Microarrays (Basel, Switzerland) , 4 ( 3 ): 389 – 406 , 8 2015 . OpenUrl 21. ↵ Philip Zimmermann , Matthias Hirsch-Hoffmann , Lars Hennig , and Wilhelm Gruissem . GENEVESTIGATOR . Arabidopsis Microarray Database and Analysis Toolbox 1[w] . Back to top Previous Next Posted January 21, 2017. Download PDF Supplementary Material Email Thank you for your interest in spreading the word about bioRxiv. NOTE: Your email address is requested solely to identify you as the sender of this article. Your Email * Your Name * Send To * Enter multiple addresses on separate lines or separate them with commas. You are going to email the following Creation of gene expression database on preeclampsia-affected human placenta Message Subject (Your Name) has forwarded a page to you from bioRxiv Message Body (Your Name) thought you would like to see this page from the bioRxiv website. Your Personal Message CAPTCHA This question is for testing whether or not you are a human visitor and to prevent automated spam submissions. Share Creation of gene expression database on preeclampsia-affected human placenta Oleksandr Lykhenko , Alina Frolova , Maria Obolenska bioRxiv 102012; doi: https://doi.org/10.1101/102012 Share This Article: Copy Citation Tools Creation of gene expression database on preeclampsia-affected human placenta Oleksandr Lykhenko , Alina Frolova , Maria Obolenska bioRxiv 102012; doi: https://doi.org/10.1101/102012 Citation Manager Formats BibTeX Bookends EasyBib EndNote (tagged) EndNote 8 (xml) Medlars Mendeley Papers RefWorks Tagged Ref Manager RIS Zotero Tweet Widget Facebook Like Google Plus One Subject Area Bioinformatics Subject Areas All Articles Animal Behavior and Cognition (7976) Biochemistry (18668) Bioengineering (14789) Bioinformatics (44225) Biophysics (22489) Cancer Biology (19616) Cell Biology (26774) Clinical Trials (138) Developmental Biology (13908) Ecology (20906) Epidemiology (2067) Evolutionary Biology (25347) Genetics (16112) Genomics (23419) Immunology (18631) Microbiology (42297) Molecular Biology (17968) Neuroscience (93036) Paleontology (694) Pathology (2973) Pharmacology and Toxicology (5071) Physiology (8080) Plant Biology (15918) Scientific Communication and Education (2093) Synthetic Biology (4539) Systems Biology (10194) Zoology (2376) window.__CF$cv$params={r:'a38ffaec5911750b',t:'MTc4OTA1OTUzNA==',u:'01a08c4217ae74a091be17288c3ac168',ut:'R2hkAZuTs.KSSBs_VU5AlLuExggkNODdGmliYo0eFbE-1789059536-1.2.1.1-Nk9L6vfqJBOaVAlC4jXIw.VcgCoMV356c3LDflMQIWhiRfFScF3ZeVpI8QLNrrqipnNbPHBlqldOu2eqX_V_Hf1m7H4OGKIHIgp4nPdz4_E',i:60};(function(){if(!document.body)return;var s=document.createElement('script');s.src='/cdn-cgi/challenge-platform/scripts/precursor/main.js';document.head.appendChild(s);})();

Text is read by the "Ask this paper" AI Q&A widget below. Extraction quality varies by source — PMC NXML preserves structure cleanly, OA-HTML may include some navigation residue, and OA-PDF can have broken hyphenation. The publisher copy (via DOI) is the canonical version.

My notes (saved in your browser only)

⚙ Ask this paper AI returns verbatim quotes from the full text · source: preprint-html ⓘ

Answers must be backed by verbatim quotes from this paper's full text. Hallucinated quotes are dropped automatically; if no verbatim passage answers the question, we say so. How this works

Citation neighborhood (no data yet)

We don't have any in-corpus citations linked to this paper yet. The paper's references may be in our DB but unresolved to ``paper_id`` (resolution happens at ingest when the cited DOI matches a row we already have). Run the cross-source citation reconcile pass to retry.

Source provenance

europepmc
last seen: 2026-05-19T01:45:01.086888+00:00
unpaywall
last seen: 2026-05-23T02:00:01.238055+00:00
License: CC-BY-4.0