Establishing the ELIXIR Microbiome Community

preprint OA: closed CC-BY-4.0

Abstract

Microbiome research has grown substantially over the past decade in terms of the range of biomes sampled, identified taxa, and the volume of data derived from the samples. In particular, experimental approaches such as metagenomics, metabarcoding, metatranscriptomics and metaproteomics have provided profound insights into the vast, hitherto unknown, microbial biodiversity. The ELIXIR Marine Metagenomics Community, initiated amongst researchers focusing on marine microbiomes, has concentrated on promoting standards around microbiome-derived sequence analysis, as well as understanding the gaps in methods and reference databases, and solutions to computational overheads of performing such analyses. Nevertheless, the methods used and the challenges faced are not confined to marine studies, but are broadly applicable to all other biomes. Thus, expanding this Community to a more inclusive ELIXIR Microbiome Community will enable it to encompass a broad range of biomes and link expertise across ‘omics technologies. Furthermore, engaging with a large number of researchers will improve the efficiency and sustainability of bioinformatics infrastructure and resources for microbiome research (standards, data, tools, workflows, training), which will enable a deeper understanding of the function and taxonomic composition of the different microbial communities.
Full text 396,463 characters · extracted from preprint-html · click to expand
Establishing the ELIXIR Microbiome Community | F1000Research "use strict";function _typeof(t){return(_typeof="function"==typeof Symbol&&"symbol"==typeof Symbol.iterator?function(t){return typeof t}:function(t){return t&&"function"==typeof Symbol&&t.constructor===Symbol&&t!==Symbol.prototype?"symbol":typeof t})(t)}!function(){var t=function(){var t,e,o=[],n=window,r=n;for(;r;){try{if(r.frames.__tcfapiLocator){t=r;break}}catch(t){}if(r===n.top)break;r=r.parent}t||(!function t(){var e=n.document,o=!!n.frames.__tcfapiLocator;if(!o)if(e.body){var r=e.createElement("iframe");r.style.cssText="display:none",r.name="__tcfapiLocator",e.body.appendChild(r)}else setTimeout(t,5);return!o}(),n.__tcfapi=function(){for(var t=arguments.length,n=new Array(t),r=0;r 3&&2===parseInt(n[1],10)&&"boolean"==typeof n[3]&&(e=n[3],"function"==typeof n[2]&&n[2]("set",!0)):"ping"===n[0]?"function"==typeof n[2]&&n[2]({gdprApplies:e,cmpLoaded:!1,cmpStatus:"stub"}):o.push(n)},n.addEventListener("message",(function(t){var e="string"==typeof t.data,o={};if(e)try{o=JSON.parse(t.data)}catch(t){}else o=t.data;var n="object"===_typeof(o)&&null!==o?o.__tcfapiCall:null;n&&window.__tcfapi(n.command,n.version,(function(o,r){var a={__tcfapiReturn:{returnValue:o,success:r,callId:n.callId}};t&&t.source&&t.source.postMessage&&t.source.postMessage(e?JSON.stringify(a):a,"*")}),n.parameter)}),!1))};"undefined"!=typeof module?module.exports=t:t()}(); dataLayer = dataLayer || []; // Standard GTM initialization - Google Consent Mode handles consent automatically (function(w,d,s,l,i){w[l]=w[l]||[];w[l].push({'gtm.start': new Date().getTime(),event:'gtm.js'});var f=d.getElementsByTagName(s)[0], j=d.createElement(s),dl=l!='dataLayer'?'&l='+l:'';j.async=true;j.src= 'https://www.googletagmanager.com/gtm.js?id='+i+dl+ '>m_auth=hzk0Vc3qFsQYhCrIoHz68A>m_preview=env-1>m_cookies_win=x';f.parentNode.insertBefore(j,f); })(window,document,'script','dataLayer','GTM-MWFK8L5J'); ;window.NREUM||(NREUM={});NREUM.init={distributed_tracing:{enabled:true},privacy:{cookies_enabled:true},ajax:{deny_list:["bam.nr-data.net"]}}; ;NREUM.loader_config={accountID:"438030",trustKey:"438030",agentID:"772317073",licenseKey:"97f8f67f26",applicationID:"772317073"} ;NREUM.info={beacon:"bam.nr-data.net",errorBeacon:"bam.nr-data.net",licenseKey:"97f8f67f26",applicationID:"772317073",sa:1} ;/*! For license information please see nr-loader-spa-1.236.0.min.js.LICENSE.txt */ (()=>{"use strict";var e,t,r={5763:(e,t,r)=>{r.d(t,{P_:()=>l,Mt:()=>g,C5:()=>s,DL:()=>v,OP:()=>T,lF:()=>D,Yu:()=>y,Dg:()=>h,CX:()=>c,GE:()=>b,sU:()=>_});var n=r(8632),i=r(9567);const o={beacon:n.ce.beacon,errorBeacon:n.ce.errorBeacon,licenseKey:void 0,applicationID:void 0,sa:void 0,queueTime:void 0,applicationTime:void 0,ttGuid:void 0,user:void 0,account:void 0,product:void 0,extra:void 0,jsAttributes:{},userAttributes:void 0,atts:void 0,transactionName:void 0,tNamePlain:void 0},a={};function s(e){if(!e)throw new Error("All info objects require an agent identifier!");if(!a[e])throw new Error("Info for ".concat(e," was never set"));return a[e]}function c(e,t){if(!e)throw new Error("All info objects require an agent identifier!");a[e]=(0,i.D)(t,o),(0,n.Qy)(e,a[e],"info")}var u=r(7056);const d=()=>{const e={blockSelector:"[data-nr-block]",maskInputOptions:{password:!0}};return{allow_bfcache:!0,privacy:{cookies_enabled:!0},ajax:{deny_list:void 0,enabled:!0,harvestTimeSeconds:10},distributed_tracing:{enabled:void 0,exclude_newrelic_header:void 0,cors_use_newrelic_header:void 0,cors_use_tracecontext_headers:void 0,allowed_origins:void 0},session:{domain:void 0,expiresMs:u.oD,inactiveMs:u.Hb},ssl:void 0,obfuscate:void 0,jserrors:{enabled:!0,harvestTimeSeconds:10},metrics:{enabled:!0},page_action:{enabled:!0,harvestTimeSeconds:30},page_view_event:{enabled:!0},page_view_timing:{enabled:!0,harvestTimeSeconds:30,long_task:!1},session_trace:{enabled:!0,harvestTimeSeconds:10},harvest:{tooManyRequestsDelay:60},session_replay:{enabled:!1,harvestTimeSeconds:60,sampleRate:.1,errorSampleRate:.1,maskTextSelector:"*",maskAllInputs:!0,get blockClass(){return"nr-block"},get ignoreClass(){return"nr-ignore"},get maskTextClass(){return"nr-mask"},get blockSelector(){return e.blockSelector},set blockSelector(t){e.blockSelector+=",".concat(t)},get maskInputOptions(){return e.maskInputOptions},set maskInputOptions(t){e.maskInputOptions={...t,password:!0}}},spa:{enabled:!0,harvestTimeSeconds:10}}},f={};function l(e){if(!e)throw new Error("All configuration objects require an agent identifier!");if(!f[e])throw new Error("Configuration for ".concat(e," was never set"));return f[e]}function h(e,t){if(!e)throw new Error("All configuration objects require an agent identifier!");f[e]=(0,i.D)(t,d()),(0,n.Qy)(e,f[e],"config")}function g(e,t){if(!e)throw new Error("All configuration objects require an agent identifier!");var r=l(e);if(r){for(var n=t.split("."),i=0;i {r.d(t,{D:()=>i});var n=r(50);function i(e,t){try{if(!e||"object"!=typeof e)return(0,n.Z)("Setting a Configurable requires an object as input");if(!t||"object"!=typeof t)return(0,n.Z)("Setting a Configurable requires a model to set its initial properties");const r=Object.create(Object.getPrototypeOf(t),Object.getOwnPropertyDescriptors(t)),o=0===Object.keys(r).length?e:r;for(let a in o)if(void 0!==e[a])try{"object"==typeof e[a]&&"object"==typeof t[a]?r[a]=i(e[a],t[a]):r[a]=e[a]}catch(e){(0,n.Z)("An error occurred while setting a property of a Configurable",e)}return r}catch(e){(0,n.Z)("An error occured while setting a Configurable",e)}}},6818:(e,t,r)=>{r.d(t,{Re:()=>i,gF:()=>o,q4:()=>n});const n="1.236.0",i="PROD",o="CDN"},385:(e,t,r)=>{r.d(t,{FN:()=>a,IF:()=>u,Nk:()=>f,Tt:()=>s,_A:()=>o,il:()=>n,pL:()=>c,v6:()=>i,w1:()=>d});const n="undefined"!=typeof window&&!!window.document,i="undefined"!=typeof WorkerGlobalScope&&("undefined"!=typeof self&&self instanceof WorkerGlobalScope&&self.navigator instanceof WorkerNavigator||"undefined"!=typeof globalThis&&globalThis instanceof WorkerGlobalScope&&globalThis.navigator instanceof WorkerNavigator),o=n?window:"undefined"!=typeof WorkerGlobalScope&&("undefined"!=typeof self&&self instanceof WorkerGlobalScope&&self||"undefined"!=typeof globalThis&&globalThis instanceof WorkerGlobalScope&&globalThis),a=""+o?.location,s=/iPad|iPhone|iPod/.test(navigator.userAgent),c=s&&"undefined"==typeof SharedWorker,u=(()=>{const e=navigator.userAgent.match(/Firefox[/\s](\d+\.\d+)/);return Array.isArray(e)&&e.length>=2?+e[1]:0})(),d=Boolean(n&&window.document.documentMode),f=!!navigator.sendBeacon},1117:(e,t,r)=>{r.d(t,{w:()=>o});var n=r(50);const i={agentIdentifier:"",ee:void 0};class o{constructor(e){try{if("object"!=typeof e)return(0,n.Z)("shared context requires an object as input");this.sharedContext={},Object.assign(this.sharedContext,i),Object.entries(e).forEach((e=>{let[t,r]=e;Object.keys(i).includes(t)&&(this.sharedContext[t]=r)}))}catch(e){(0,n.Z)("An error occured while setting SharedContext",e)}}}},8e3:(e,t,r)=>{r.d(t,{L:()=>d,R:()=>c});var n=r(2177),i=r(1284),o=r(4322),a=r(3325);const s={};function c(e,t){const r={staged:!1,priority:a.p[t]||0};u(e),s[e].get(t)||s[e].set(t,r)}function u(e){e&&(s[e]||(s[e]=new Map))}function d(){let e=arguments.length>0&&void 0!==arguments[0]?arguments[0]:"",t=arguments.length>1&&void 0!==arguments[1]?arguments[1]:"feature";if(u(e),!e||!s[e].get(t))return a(t);s[e].get(t).staged=!0;const r=[...s[e]];function a(t){const r=e?n.ee.get(e):n.ee,a=o.X.handlers;if(r.backlog&&a){var s=r.backlog[t],c=a[t];if(c){for(var u=0;s&&u {let[t,r]=e;return r.staged}))&&(r.sort(((e,t)=>e[1].priority-t[1].priority)),r.forEach((e=>{let[t]=e;a(t)})))}function f(e,t){var r=e[1];(0,i.D)(t[r],(function(t,r){var n=e[0];if(r[0]===n){var i=r[1],o=e[3],a=e[2];i.apply(o,a)}}))}},2177:(e,t,r)=>{r.d(t,{c:()=>f,ee:()=>u});var n=r(8632),i=r(2210),o=r(1284),a=r(5763),s="nr@context";let c=(0,n.fP)();var u;function d(){}function f(e){return(0,i.X)(e,s,l)}function l(){return new d}function h(){u.aborted=!0,u.backlog={}}c.ee?u=c.ee:(u=function e(t,r){var n={},c={},f={},g=!1;try{g=16===r.length&&(0,a.OP)(r).isolatedBacklog}catch(e){}var p={on:b,addEventListener:b,removeEventListener:y,emit:v,get:x,listeners:w,context:m,buffer:A,abort:h,aborted:!1,isBuffering:E,debugId:r,backlog:g?{}:t&&"object"==typeof t.backlog?t.backlog:{}};return p;function m(e){return e&&e instanceof d?e:e?(0,i.X)(e,s,l):l()}function v(e,r,n,i,o){if(!1!==o&&(o=!0),!u.aborted||i){t&&o&&t.emit(e,r,n);for(var a=m(n),s=w(e),d=s.length,f=0;fn,p:()=>i});var n=r(2177).ee.get("handle");function i(e,t,r,i,o){o?(o.buffer([e],i),o.emit(e,t,r)):(n.buffer([e],i),n.emit(e,t,r))}},4322:(e,t,r)=>{r.d(t,{X:()=>o});var n=r(5546);o.on=a;var i=o.handlers={};function o(e,t,r,o){a(o||n.E,i,e,t,r)}function a(e,t,r,i,o){o||(o="feature"),e||(e=n.E);var a=t[o]=t[o]||{};(a[r]=a[r]||[]).push([e,i])}},3239:(e,t,r)=>{r.d(t,{bP:()=>s,iz:()=>c,m$:()=>a});var n=r(385);let i=!1,o=!1;try{const e={get passive(){return i=!0,!1},get signal(){return o=!0,!1}};n._A.addEventListener("test",null,e),n._A.removeEventListener("test",null,e)}catch(e){}function a(e,t){return i||o?{capture:!!e,passive:i,signal:t}:!!e}function s(e,t){let r=arguments.length>2&&void 0!==arguments[2]&&arguments[2],n=arguments.length>3?arguments[3]:void 0;window.addEventListener(e,t,a(r,n))}function c(e,t){let r=arguments.length>2&&void 0!==arguments[2]&&arguments[2],n=arguments.length>3?arguments[3]:void 0;document.addEventListener(e,t,a(r,n))}},4402:(e,t,r)=>{r.d(t,{Ht:()=>u,M:()=>c,Rl:()=>a,ky:()=>s});var n=r(385);const i="xxxxxxxx-xxxx-4xxx-yxxx-xxxxxxxxxxxx";function o(e,t){return e?15&e[t]:16*Math.random()|0}function a(){const e=n._A?.crypto||n._A?.msCrypto;let t,r=0;return e&&e.getRandomValues&&(t=e.getRandomValues(new Uint8Array(31))),i.split("").map((e=>"x"===e?o(t,++r).toString(16):"y"===e?(3&o()|8).toString(16):e)).join("")}function s(e){const t=n._A?.crypto||n._A?.msCrypto;let r,i=0;t&&t.getRandomValues&&(r=t.getRandomValues(new Uint8Array(31)));const a=[];for(var s=0;s {r.d(t,{Bq:()=>n,Hb:()=>o,oD:()=>i});const n="NRBA",i=144e5,o=18e5},7894:(e,t,r)=>{function n(){return Math.round(performance.now())}r.d(t,{z:()=>n})},7243:(e,t,r)=>{r.d(t,{e:()=>o});var n=r(385),i={};function o(e){if(e in i)return i[e];if(0===(e||"").indexOf("data:"))return{protocol:"data"};let t;var r=n._A?.location,o={};if(n.il)t=document.createElement("a"),t.href=e;else try{t=new URL(e,r.href)}catch(e){return o}o.port=t.port;var a=t.href.split("://");!o.port&&a[1]&&(o.port=a[1].split("/")[0].split("@").pop().split(":")[1]),o.port&&"0"!==o.port||(o.port="https"===a[0]?"443":"80"),o.hostname=t.hostname||r.hostname,o.pathname=t.pathname,o.protocol=a[0],"/"!==o.pathname.charAt(0)&&(o.pathname="/"+o.pathname);var s=!t.protocol||":"===t.protocol||t.protocol===r.protocol,c=t.hostname===r.hostname&&t.port===r.port;return o.sameOrigin=s&&(!t.hostname||c),"/"===o.pathname&&(i[e]=o),o}},50:(e,t,r)=>{function n(e,t){"function"==typeof console.warn&&(console.warn("New Relic: ".concat(e)),t&&console.warn(t))}r.d(t,{Z:()=>n})},2587:(e,t,r)=>{r.d(t,{N:()=>c,T:()=>u});var n=r(2177),i=r(5546),o=r(8e3),a=r(3325);const s={stn:[a.D.sessionTrace],err:[a.D.jserrors,a.D.metrics],ins:[a.D.pageAction],spa:[a.D.spa],sr:[a.D.sessionReplay,a.D.sessionTrace]};function c(e,t){const r=n.ee.get(t);e&&"object"==typeof e&&(Object.entries(e).forEach((e=>{let[t,n]=e;void 0===u[t]&&(s[t]?s[t].forEach((e=>{n?(0,i.p)("feat-"+t,[],void 0,e,r):(0,i.p)("block-"+t,[],void 0,e,r),(0,i.p)("rumresp-"+t,[Boolean(n)],void 0,e,r)})):n&&(0,i.p)("feat-"+t,[],void 0,void 0,r),u[t]=Boolean(n))})),Object.keys(s).forEach((e=>{void 0===u[e]&&(s[e]?.forEach((t=>(0,i.p)("rumresp-"+e,[!1],void 0,t,r))),u[e]=!1)})),(0,o.L)(t,a.D.pageViewEvent))}const u={}},2210:(e,t,r)=>{r.d(t,{X:()=>i});var n=Object.prototype.hasOwnProperty;function i(e,t,r){if(n.call(e,t))return e[t];var i=r();if(Object.defineProperty&&Object.keys)try{return Object.defineProperty(e,t,{value:i,writable:!0,enumerable:!1}),i}catch(e){}return e[t]=i,i}},1284:(e,t,r)=>{r.d(t,{D:()=>n});const n=(e,t)=>Object.entries(e||{}).map((e=>{let[r,n]=e;return t(r,n)}))},4351:(e,t,r)=>{r.d(t,{P:()=>o});var n=r(2177);const i=()=>{const e=new WeakSet;return(t,r)=>{if("object"==typeof r&&null!==r){if(e.has(r))return;e.add(r)}return r}};function o(e){try{return JSON.stringify(e,i())}catch(e){try{n.ee.emit("internal-error",[e])}catch(e){}}}},3960:(e,t,r)=>{r.d(t,{K:()=>a,b:()=>o});var n=r(3239);function i(){return"undefined"==typeof document||"complete"===document.readyState}function o(e,t){if(i())return e();(0,n.bP)("load",e,t)}function a(e){if(i())return e();(0,n.iz)("DOMContentLoaded",e)}},8632:(e,t,r)=>{r.d(t,{EZ:()=>u,Qy:()=>c,ce:()=>o,fP:()=>a,gG:()=>d,mF:()=>s});var n=r(7894),i=r(385);const o={beacon:"bam.nr-data.net",errorBeacon:"bam.nr-data.net"};function a(){return i._A.NREUM||(i._A.NREUM={}),void 0===i._A.newrelic&&(i._A.newrelic=i._A.NREUM),i._A.NREUM}function s(){let e=a();return e.o||(e.o={ST:i._A.setTimeout,SI:i._A.setImmediate,CT:i._A.clearTimeout,XHR:i._A.XMLHttpRequest,REQ:i._A.Request,EV:i._A.Event,PR:i._A.Promise,MO:i._A.MutationObserver,FETCH:i._A.fetch}),e}function c(e,t,r){let i=a();const o=i.initializedAgents||{},s=o[e]||{};return Object.keys(s).length||(s.initializedAt={ms:(0,n.z)(),date:new Date}),i.initializedAgents={...o,[e]:{...s,[r]:t}},i}function u(e,t){a()[e]=t}function d(){return function(){let e=a();const t=e.info||{};e.info={beacon:o.beacon,errorBeacon:o.errorBeacon,...t}}(),function(){let e=a();const t=e.init||{};e.init={...t}}(),s(),function(){let e=a();const t=e.loader_config||{};e.loader_config={...t}}(),a()}},7956:(e,t,r)=>{r.d(t,{N:()=>i});var n=r(3239);function i(e){let t=arguments.length>1&&void 0!==arguments[1]&&arguments[1],r=arguments.length>2?arguments[2]:void 0,i=arguments.length>3?arguments[3]:void 0;return void(0,n.iz)("visibilitychange",(function(){if(t)return void("hidden"==document.visibilityState&&e());e(document.visibilityState)}),r,i)}},1214:(e,t,r)=>{r.d(t,{em:()=>v,u5:()=>N,QU:()=>S,_L:()=>I,Gm:()=>L,Lg:()=>M,gy:()=>U,BV:()=>Q,Kf:()=>ee});var n=r(2177);const i="nr@original";var o=Object.prototype.hasOwnProperty,a=!1;function s(e,t){return e||(e=n.ee),r.inPlace=function(e,t,n,i,o){n||(n="");var a,s,c,u="-"===n.charAt(0);for(c=0;c 2?n-2:0),o=2;o {r(A[T],e,w),r(E[T],e,w)})),r(l._A,"fetch",y),t.on(y+"end",(function(e,r){var n=this;if(r){var i=r.headers.get("content-length");null!==i&&(n.rxSize=i),t.emit(y+"done",[null,r],n)}else t.emit(y+"done",[e],n)})),t}const O={},j=["pushState","replaceState"];function S(e){const t=function(e){return(e||n.ee).get("history")}(e);return!l.il||O[t.debugId]++||(O[t.debugId]=1,s(t).inPlace(window.history,j,"-")),t}var P=r(3239);const C={},R=["appendChild","insertBefore","replaceChild"];function I(e){const t=function(e){return(e||n.ee).get("jsonp")}(e);if(!l.il||C[t.debugId])return t;C[t.debugId]=!0;var r=s(t),i=/[?&](?:callback|cb)=([^&#]+)/,o=/(.*)\.([^.]+)/,a=/^(\w+)(\.|$)(.*)$/;function c(e,t){var r=e.match(a),n=r[1],i=r[3];return i?c(i,t[n]):t[n]}return r.inPlace(Node.prototype,R,"dom-"),t.on("dom-start",(function(e){!function(e){if(!e||"string"!=typeof e.nodeName||"script"!==e.nodeName.toLowerCase())return;if("function"!=typeof e.addEventListener)return;var n=(a=e.src,s=a.match(i),s?s[1]:null);var a,s;if(!n)return;var u=function(e){var t=e.match(o);if(t&&t.length>=3)return{key:t[2],parent:c(t[1],window)};return{key:e,parent:window}}(n);if("function"!=typeof u.parent[u.key])return;var d={};function f(){t.emit("jsonp-end",[],d),e.removeEventListener("load",f,(0,P.m$)(!1)),e.removeEventListener("error",l,(0,P.m$)(!1))}function l(){t.emit("jsonp-error",[],d),t.emit("jsonp-end",[],d),e.removeEventListener("load",f,(0,P.m$)(!1)),e.removeEventListener("error",l,(0,P.m$)(!1))}r.inPlace(u.parent,[u.key],"cb-",d),e.addEventListener("load",f,(0,P.m$)(!1)),e.addEventListener("error",l,(0,P.m$)(!1)),t.emit("new-jsonp",[e.src],d)}(e[0])})),t}var k=r(5763);const H={};function L(e){const t=function(e){return(e||n.ee).get("mutation")}(e);if(!l.il||H[t.debugId])return t;H[t.debugId]=!0;var r=s(t),i=k.Yu.MO;return i&&(window.MutationObserver=function(e){return this instanceof i?new i(r(e,"fn-")):i.apply(this,arguments)},MutationObserver.prototype=i.prototype),t}const z={};function M(e){const t=function(e){return(e||n.ee).get("promise")}(e);if(z[t.debugId])return t;z[t.debugId]=!0;var r=n.c,o=s(t),a=k.Yu.PR;return a&&function(){function e(r){var n=t.context(),i=o(r,"executor-",n,null,!1);const s=Reflect.construct(a,[i],e);return t.context(s).getCtx=function(){return n},s}l._A.Promise=e,Object.defineProperty(e,"name",{value:"Promise"}),e.toString=function(){return a.toString()},Object.setPrototypeOf(e,a),["all","race"].forEach((function(r){const n=a[r];e[r]=function(e){let i=!1;[...e||[]].forEach((e=>{this.resolve(e).then(a("all"===r),a(!1))}));const o=n.apply(this,arguments);return o;function a(e){return function(){t.emit("propagate",[null,!i],o,!1,!1),i=i||!e}}}})),["resolve","reject"].forEach((function(r){const n=a[r];e[r]=function(e){const r=n.apply(this,arguments);return e!==r&&t.emit("propagate",[e,!0],r,!1,!1),r}})),e.prototype=a.prototype;const n=a.prototype.then;a.prototype.then=function(){var e=this,i=r(e);i.promise=e;for(var a=arguments.length,s=new Array(a),c=0;c e())),t};function m(e,t){i.inPlace(t,["onreadystatechange"],"fn-",E)}function b(){var e=this,t=r.context(e);e.readyState>3&&!t.resolved&&(t.resolved=!0,r.emit("xhr-resolved",[],e)),i.inPlace(e,f,"fn-",E)}if(function(e,t){for(var r in e)t[r]=e[r]}(o,p),p.prototype=o.prototype,i.inPlace(p.prototype,J,"-xhr-",E),r.on("send-xhr-start",(function(e,t){m(e,t),function(e){h.push(e),a&&(y?y.then(A):u?u(A):(w=-w,x.data=w))}(t)})),r.on("open-xhr-start",m),a){var y=c&&c.resolve();if(!u&&!c){var w=1,x=document.createTextNode(w);new a(A).observe(x,{characterData:!0})}}else t.on("fn-end",(function(e){e[0]&&e[0].type===d||A()}));function A(){for(var e=0;e {r.d(t,{t:()=>n});const n=r(3325).D.ajax},6660:(e,t,r)=>{r.d(t,{A:()=>i,t:()=>n});const n=r(3325).D.jserrors,i="nr@seenError"},3081:(e,t,r)=>{r.d(t,{gF:()=>o,mY:()=>i,t9:()=>n,vz:()=>s,xS:()=>a});const n=r(3325).D.metrics,i="sm",o="cm",a="storeSupportabilityMetrics",s="storeEventMetrics"},4649:(e,t,r)=>{r.d(t,{t:()=>n});const n=r(3325).D.pageAction},7633:(e,t,r)=>{r.d(t,{Dz:()=>i,OJ:()=>a,qw:()=>o,t9:()=>n});const n=r(3325).D.pageViewEvent,i="firstbyte",o="domcontent",a="windowload"},9251:(e,t,r)=>{r.d(t,{t:()=>n});const n=r(3325).D.pageViewTiming},3614:(e,t,r)=>{r.d(t,{BST_RESOURCE:()=>i,END:()=>s,FEATURE_NAME:()=>n,FN_END:()=>u,FN_START:()=>c,PUSH_STATE:()=>d,RESOURCE:()=>o,START:()=>a});const n=r(3325).D.sessionTrace,i="bstResource",o="resource",a="-start",s="-end",c="fn"+a,u="fn"+s,d="pushState"},7836:(e,t,r)=>{r.d(t,{BODY:()=>A,CB_END:()=>E,CB_START:()=>u,END:()=>x,FEATURE_NAME:()=>i,FETCH:()=>_,FETCH_BODY:()=>v,FETCH_DONE:()=>m,FETCH_START:()=>p,FN_END:()=>c,FN_START:()=>s,INTERACTION:()=>l,INTERACTION_API:()=>d,INTERACTION_EVENTS:()=>o,JSONP_END:()=>b,JSONP_NODE:()=>g,JS_TIME:()=>T,MAX_TIMER_BUDGET:()=>a,REMAINING:()=>f,SPA_NODE:()=>h,START:()=>w,originalSetTimeout:()=>y});var n=r(5763);const i=r(3325).D.spa,o=["click","submit","keypress","keydown","keyup","change"],a=999,s="fn-start",c="fn-end",u="cb-start",d="api-ixn-",f="remaining",l="interaction",h="spaNode",g="jsonpNode",p="fetch-start",m="fetch-done",v="fetch-body-",b="jsonp-end",y=n.Yu.ST,w="-start",x="-end",A="-body",E="cb"+x,T="jsTime",_="fetch"},5938:(e,t,r)=>{r.d(t,{W:()=>o});var n=r(5763),i=r(2177);class o{constructor(e,t,r){this.agentIdentifier=e,this.aggregator=t,this.ee=i.ee.get(e,(0,n.OP)(this.agentIdentifier).isolatedBacklog),this.featureName=r,this.blocked=!1}}},9144:(e,t,r)=>{r.d(t,{j:()=>m});var n=r(3325),i=r(5763),o=r(5546),a=r(2177),s=r(7894),c=r(8e3),u=r(3960),d=r(385),f=r(50),l=r(3081),h=r(8632);function g(){const e=(0,h.gG)();["setErrorHandler","finished","addToTrace","inlineHit","addRelease","addPageAction","setCurrentRouteName","setPageViewName","setCustomAttribute","interaction","noticeError","setUserId"].forEach((t=>{e[t]=function(){for(var r=arguments.length,n=new Array(r),i=0;i 1?r-1:0),i=1;i {e.exposed&&e.api[t]&&o.push(e.api[t](...n))})),o.length>1?o:o[0]}(t,...n)}}))}var p=r(2587);function m(e){let t=arguments.length>1&&void 0!==arguments[1]?arguments[1]:{},m=arguments.length>2?arguments[2]:void 0,v=arguments.length>3?arguments[3]:void 0,{init:b,info:y,loader_config:w,runtime:x={loaderType:m},exposed:A=!0}=t;const E=(0,h.gG)();y||(b=E.init,y=E.info,w=E.loader_config),(0,i.Dg)(e,b||{}),(0,i.GE)(e,w||{}),(0,i.sU)(e,x),y.jsAttributes??={},d.v6&&(y.jsAttributes.isWorker=!0),(0,i.CX)(e,y),g();const T=function(e,t){t||(0,c.R)(e,"api");const h={};var g=a.ee.get(e),p=g.get("tracer"),m="api-",v=m+"ixn-";function b(t,r,n,o){const a=(0,i.C5)(e);return null===r?delete a.jsAttributes[t]:(0,i.CX)(e,{...a,jsAttributes:{...a.jsAttributes,[t]:r}}),x(m,n,!0,o||null===r?"session":void 0)(t,r)}function y(){}["setErrorHandler","finished","addToTrace","inlineHit","addRelease"].forEach((e=>h[e]=x(m,e,!0,"api"))),h.addPageAction=x(m,"addPageAction",!0,n.D.pageAction),h.setCurrentRouteName=x(m,"routeName",!0,n.D.spa),h.setPageViewName=function(t,r){if("string"==typeof t)return"/"!==t.charAt(0)&&(t="/"+t),(0,i.OP)(e).customTransaction=(r||"http://custom.transaction")+t,x(m,"setPageViewName",!0)()},h.setCustomAttribute=function(e,t){let r=arguments.length>2&&void 0!==arguments[2]&&arguments[2];if("string"==typeof e){if(["string","number"].includes(typeof t)||null===t)return b(e,t,"setCustomAttribute",r);(0,f.Z)("Failed to execute setCustomAttribute.\nNon-null value must be a string or number type, but a type of was provided."))}else(0,f.Z)("Failed to execute setCustomAttribute.\nName must be a string type, but a type of was provided."))},h.setUserId=function(e){if("string"==typeof e||null===e)return b("enduser.id",e,"setUserId",!0);(0,f.Z)("Failed to execute setUserId.\nNon-null value must be a string type, but a type of was provided."))},h.interaction=function(){return(new y).get()};var w=y.prototype={createTracer:function(e,t){var r={},i=this,a="function"==typeof t;return(0,o.p)(v+"tracer",[(0,s.z)(),e,r],i,n.D.spa,g),function(){if(p.emit((a?"":"no-")+"fn-start",[(0,s.z)(),i,a],r),a)try{return t.apply(this,arguments)}catch(e){throw p.emit("fn-err",[arguments,this,"string"==typeof e?new Error(e):e],r),e}finally{p.emit("fn-end",[(0,s.z)()],r)}}}};function x(e,t,r,i){return function(){return(0,o.p)(l.xS,["API/"+t+"/called"],void 0,n.D.metrics,g),i&&(0,o.p)(e+t,[(0,s.z)(),...arguments],r?null:this,i,g),r?void 0:this}}function A(){r.e(439).then(r.bind(r,7438)).then((t=>{let{setAPI:r}=t;r(e),(0,c.L)(e,"api")})).catch((()=>(0,f.Z)("Downloading runtime APIs failed...")))}return["actionText","setName","setAttribute","save","ignore","onEnd","getContext","end","get"].forEach((e=>{w[e]=x(v,e,void 0,n.D.spa)})),h.noticeError=function(e,t){"string"==typeof e&&(e=new Error(e)),(0,o.p)(l.xS,["API/noticeError/called"],void 0,n.D.metrics,g),(0,o.p)("err",[e,(0,s.z)(),!1,t],void 0,n.D.jserrors,g)},d.il?(0,u.b)((()=>A()),!0):A(),h}(e,v);return(0,h.Qy)(e,T,"api"),(0,h.Qy)(e,A,"exposed"),(0,h.EZ)("activatedFeatures",p.T),T}},3325:(e,t,r)=>{r.d(t,{D:()=>n,p:()=>i});const n={ajax:"ajax",jserrors:"jserrors",metrics:"metrics",pageAction:"page_action",pageViewEvent:"page_view_event",pageViewTiming:"page_view_timing",sessionReplay:"session_replay",sessionTrace:"session_trace",spa:"spa"},i={[n.pageViewEvent]:1,[n.pageViewTiming]:2,[n.metrics]:3,[n.jserrors]:4,[n.ajax]:5,[n.sessionTrace]:6,[n.pageAction]:7,[n.spa]:8,[n.sessionReplay]:9}}},n={};function i(e){var t=n[e];if(void 0!==t)return t.exports;var o=n[e]={exports:{}};return r[e](o,o.exports,i),o.exports}i.m=r,i.d=(e,t)=>{for(var r in t)i.o(t,r)&&!i.o(e,r)&&Object.defineProperty(e,r,{enumerable:!0,get:t[r]})},i.f={},i.e=e=>Promise.all(Object.keys(i.f).reduce(((t,r)=>(i.f[r](e,t),t)),[])),i.u=e=>(({78:"page_action-aggregate",147:"metrics-aggregate",242:"session-manager",317:"jserrors-aggregate",348:"page_view_timing-aggregate",412:"lazy-feature-loader",439:"async-api",538:"recorder",590:"session_replay-aggregate",675:"compressor",733:"session_trace-aggregate",786:"page_view_event-aggregate",873:"spa-aggregate",898:"ajax-aggregate"}[e]||e)+"."+{78:"ac76d497",147:"3dc53903",148:"1a20d5fe",242:"2a64278a",317:"49e41428",348:"bd6de33a",412:"2f55ce66",439:"30bd804e",538:"1b18459f",590:"cf0efb30",675:"ae9f91a8",733:"83105561",786:"06482edd",860:"03a8b7a5",873:"e6b09d52",898:"998ef92b"}[e]+"-1.236.0.min.js"),i.o=(e,t)=>Object.prototype.hasOwnProperty.call(e,t),e={},t="NRBA:",i.l=(r,n,o,a)=>{if(e[r])e[r].push(n);else{var s,c;if(void 0!==o)for(var u=document.getElementsByTagName("script"),d=0;d {s.onerror=s.onload=null,clearTimeout(h);var i=e[r];if(delete e[r],s.parentNode&&s.parentNode.removeChild(s),i&&i.forEach((e=>e(n))),t)return t(n)},h=setTimeout(l.bind(null,void 0,{type:"timeout",target:s}),12e4);s.onerror=l.bind(null,s.onerror),s.onload=l.bind(null,s.onload),c&&document.head.appendChild(s)}},i.r=e=>{"undefined"!=typeof Symbol&&Symbol.toStringTag&&Object.defineProperty(e,Symbol.toStringTag,{value:"Module"}),Object.defineProperty(e,"__esModule",{value:!0})},i.j=364,i.p="https://js-agent.newrelic.com/",(()=>{var e={364:0,953:0};i.f.j=(t,r)=>{var n=i.o(e,t)?e[t]:void 0;if(0!==n)if(n)r.push(n[2]);else{var o=new Promise(((r,i)=>n=e[t]=[r,i]));r.push(n[2]=o);var a=i.p+i.u(t),s=new Error;i.l(a,(r=>{if(i.o(e,t)&&(0!==(n=e[t])&&(e[t]=void 0),n)){var o=r&&("load"===r.type?"missing":r.type),a=r&&r.target&&r.target.src;s.message="Loading chunk "+t+" failed.\n("+o+": "+a+")",s.name="ChunkLoadError",s.type=o,s.request=a,n[1](s)}}),"chunk-"+t,t)}};var t=(t,r)=>{var n,o,[a,s,c]=r,u=0;if(a.some((t=>0!==e[t]))){for(n in s)i.o(s,n)&&(i.m[n]=s[n]);if(c)c(i)}for(t&&t(r);u {i.r(o);var e=i(3325),t=i(5763);const r=Object.values(e.D);function n(e){const n={};return r.forEach((r=>{n[r]=function(e,r){return!1!==(0,t.Mt)(r,"".concat(e,".enabled"))}(r,e)})),n}var a=i(9144);var s=i(5546),c=i(385),u=i(8e3),d=i(5938),f=i(3960),l=i(50);class h extends d.W{constructor(e,t,r){let n=!(arguments.length>3&&void 0!==arguments[3])||arguments[3];super(e,t,r),this.auto=n,this.abortHandler,this.featAggregate,this.onAggregateImported,n&&(0,u.R)(e,r)}importAggregator(){let e=arguments.length>0&&void 0!==arguments[0]?arguments[0]:{};if(this.featAggregate||!this.auto)return;const r=c.il&&!0===(0,t.Mt)(this.agentIdentifier,"privacy.cookies_enabled");let n;this.onAggregateImported=new Promise((e=>{n=e}));const o=async()=>{let t;try{if(r){const{setupAgentSession:e}=await Promise.all([i.e(860),i.e(242)]).then(i.bind(i,3228));t=e(this.agentIdentifier)}}catch(e){(0,l.Z)("A problem occurred when starting up session manager. This page will not start or extend any session.",e)}try{if(!this.shouldImportAgg(this.featureName,t))return void(0,u.L)(this.agentIdentifier,this.featureName);const{lazyFeatureLoader:r}=await i.e(412).then(i.bind(i,8582)),{Aggregate:o}=await r(this.featureName,"aggregate");this.featAggregate=new o(this.agentIdentifier,this.aggregator,e),n(!0)}catch(e){(0,l.Z)("Downloading and initializing ".concat(this.featureName," failed..."),e),this.abortHandler?.(),n(!1)}};c.il?(0,f.b)((()=>o()),!0):o()}shouldImportAgg(r,n){return r!==e.D.sessionReplay||!1!==(0,t.Mt)(this.agentIdentifier,"session_trace.enabled")&&(!!n?.isNew||!!n?.state.sessionReplay)}}var g=i(7633),p=i(7894);class m extends h{static featureName=g.t9;constructor(r,n){let i=!(arguments.length>2&&void 0!==arguments[2])||arguments[2];if(super(r,n,g.t9,i),("undefined"==typeof PerformanceNavigationTiming||c.Tt)&&"undefined"!=typeof PerformanceTiming){const n=(0,t.OP)(r);n[g.Dz]=Math.max(Date.now()-n.offset,0),(0,f.K)((()=>n[g.qw]=Math.max((0,p.z)()-n[g.Dz],0))),(0,f.b)((()=>{const t=(0,p.z)();n[g.OJ]=Math.max(t-n[g.Dz],0),(0,s.p)("timing",["load",t],void 0,e.D.pageViewTiming,this.ee)}))}this.importAggregator()}}var v=i(1117),b=i(1284);class y extends v.w{constructor(e){super(e),this.aggregatedData={}}store(e,t,r,n,i){var o=this.getBucket(e,t,r,i);return o.metrics=function(e,t){t||(t={count:0});return t.count+=1,(0,b.D)(e,(function(e,r){t[e]=w(r,t[e])})),t}(n,o.metrics),o}merge(e,t,r,n,i){var o=this.getBucket(e,t,n,i);if(o.metrics){var a=o.metrics;a.count+=r.count,(0,b.D)(r,(function(e,t){if("count"!==e){var n=a[e],i=r[e];i&&!i.c?a[e]=w(i.t,n):a[e]=function(e,t){if(!t)return e;t.c||(t=x(t.t));return t.min=Math.min(e.min,t.min),t.max=Math.max(e.max,t.max),t.t+=e.t,t.sos+=e.sos,t.c+=e.c,t}(i,a[e])}}))}else o.metrics=r}storeMetric(e,t,r,n){var i=this.getBucket(e,t,r);return i.stats=w(n,i.stats),i}getBucket(e,t,r,n){this.aggregatedData[e]||(this.aggregatedData[e]={});var i=this.aggregatedData[e][t];return i||(i=this.aggregatedData[e][t]={params:r||{}},n&&(i.custom=n)),i}get(e,t){return t?this.aggregatedData[e]&&this.aggregatedData[e][t]:this.aggregatedData[e]}take(e){for(var t={},r="",n=!1,i=0;i t.max&&(t.max=e),e 2&&void 0!==arguments[2])||arguments[2];super(e,r,j.t,n),c.il&&((0,t.OP)(e).initHidden=Boolean("hidden"===document.visibilityState),(0,N.N)((()=>(0,s.p)("docHidden",[(0,p.z)()],void 0,j.t,this.ee)),!0),(0,O.bP)("pagehide",(()=>(0,s.p)("winPagehide",[(0,p.z)()],void 0,j.t,this.ee))),this.importAggregator())}}var P=i(3081);class C extends h{static featureName=P.t9;constructor(e,t){let r=!(arguments.length>2&&void 0!==arguments[2])||arguments[2];super(e,t,P.t9,r),this.importAggregator()}}var R,I=i(2210),k=i(1214),H=i(2177),L={};try{R=localStorage.getItem("__nr_flags").split(","),console&&"function"==typeof console.log&&(L.console=!0,-1!==R.indexOf("dev")&&(L.dev=!0),-1!==R.indexOf("nr_dev")&&(L.nrDev=!0))}catch(e){}function z(e){try{L.console&&z(e)}catch(e){}}L.nrDev&&H.ee.on("internal-error",(function(e){z(e.stack)})),L.dev&&H.ee.on("fn-err",(function(e,t,r){z(r.stack)})),L.dev&&(z("NR AGENT IN DEVELOPMENT MODE"),z("flags: "+(0,b.D)(L,(function(e,t){return e})).join(", ")));var M=i(6660);class B extends h{static featureName=M.t;constructor(r,n){let i=!(arguments.length>2&&void 0!==arguments[2])||arguments[2];super(r,n,M.t,i),this.skipNext=0;try{this.removeOnAbort=new AbortController}catch(e){}const o=this;o.ee.on("fn-start",(function(e,t,r){o.abortHandler&&(o.skipNext+=1)})),o.ee.on("fn-err",(function(t,r,n){o.abortHandler&&!n[M.A]&&((0,I.X)(n,M.A,(function(){return!0})),this.thrown=!0,(0,s.p)("err",[n,(0,p.z)()],void 0,e.D.jserrors,o.ee))})),o.ee.on("fn-end",(function(){o.abortHandler&&!this.thrown&&o.skipNext>0&&(o.skipNext-=1)})),o.ee.on("internal-error",(function(t){(0,s.p)("ierr",[t,(0,p.z)(),!0],void 0,e.D.jserrors,o.ee)})),this.origOnerror=c._A.onerror,c._A.onerror=this.onerrorHandler.bind(this),c._A.addEventListener("unhandledrejection",(t=>{const r=function(e){let t="Unhandled Promise Rejection: ";if(e instanceof Error)try{return e.message=t+e.message,e}catch(t){return e}if(void 0===e)return new Error(t);try{return new Error(t+(0,D.P)(e))}catch(e){return new Error(t)}}(t.reason);(0,s.p)("err",[r,(0,p.z)(),!1,{unhandledPromiseRejection:1}],void 0,e.D.jserrors,this.ee)}),(0,O.m$)(!1,this.removeOnAbort?.signal)),(0,k.gy)(this.ee),(0,k.BV)(this.ee),(0,k.em)(this.ee),(0,t.OP)(r).xhrWrappable&&(0,k.Kf)(this.ee),this.abortHandler=this.#e,this.importAggregator()}#e(){this.removeOnAbort?.abort(),this.abortHandler=void 0}onerrorHandler(t,r,n,i,o){"function"==typeof this.origOnerror&&this.origOnerror(...arguments);try{this.skipNext?this.skipNext-=1:(0,s.p)("err",[o||new F(t,r,n),(0,p.z)()],void 0,e.D.jserrors,this.ee)}catch(t){try{(0,s.p)("ierr",[t,(0,p.z)(),!0],void 0,e.D.jserrors,this.ee)}catch(e){}}return!1}}function F(e,t,r){this.message=e||"Uncaught error with no additional information",this.sourceURL=t,this.line=r}let U=1;const q="nr@id";function G(e){const t=typeof e;return!e||"object"!==t&&"function"!==t?-1:e===c._A?0:(0,I.X)(e,q,(function(){return U++}))}function V(e){if("string"==typeof e&&e.length)return e.length;if("object"==typeof e){if("undefined"!=typeof ArrayBuffer&&e instanceof ArrayBuffer&&e.byteLength)return e.byteLength;if("undefined"!=typeof Blob&&e instanceof Blob&&e.size)return e.size;if(!("undefined"!=typeof FormData&&e instanceof FormData))try{return(0,D.P)(e).length}catch(e){return}}}var X=i(7243);class W{constructor(e){this.agentIdentifier=e,this.generateTracePayload=this.generateTracePayload.bind(this),this.shouldGenerateTrace=this.shouldGenerateTrace.bind(this)}generateTracePayload(e){if(!this.shouldGenerateTrace(e))return null;var r=(0,t.DL)(this.agentIdentifier);if(!r)return null;var n=(r.accountID||"").toString()||null,i=(r.agentID||"").toString()||null,o=(r.trustKey||"").toString()||null;if(!n||!i)return null;var a=(0,_.M)(),s=(0,_.Ht)(),c=Date.now(),u={spanId:a,traceId:s,timestamp:c};return(e.sameOrigin||this.isAllowedOrigin(e)&&this.useTraceContextHeadersForCors())&&(u.traceContextParentHeader=this.generateTraceContextParentHeader(a,s),u.traceContextStateHeader=this.generateTraceContextStateHeader(a,c,n,i,o)),(e.sameOrigin&&!this.excludeNewrelicHeader()||!e.sameOrigin&&this.isAllowedOrigin(e)&&this.useNewrelicHeaderForCors())&&(u.newrelicHeader=this.generateTraceHeader(a,s,c,n,i,o)),u}generateTraceContextParentHeader(e,t){return"00-"+t+"-"+e+"-01"}generateTraceContextStateHeader(e,t,r,n,i){return i+"@nr=0-1-"+r+"-"+n+"-"+e+"----"+t}generateTraceHeader(e,t,r,n,i,o){if(!("function"==typeof c._A?.btoa))return null;var a={v:[0,1],d:{ty:"Browser",ac:n,ap:i,id:e,tr:t,ti:r}};return o&&n!==o&&(a.d.tk=o),btoa((0,D.P)(a))}shouldGenerateTrace(e){return this.isDtEnabled()&&this.isAllowedOrigin(e)}isAllowedOrigin(e){var r=!1,n={};if((0,t.Mt)(this.agentIdentifier,"distributed_tracing")&&(n=(0,t.P_)(this.agentIdentifier).distributed_tracing),e.sameOrigin)r=!0;else if(n.allowed_origins instanceof Array)for(var i=0;i 2&&void 0!==arguments[2])||arguments[2];super(r,n,Z.t,i),(0,t.OP)(r).xhrWrappable&&(this.dt=new W(r),this.handler=(e,t,r,n)=>(0,s.p)(e,t,r,n,this.ee),(0,k.u5)(this.ee),(0,k.Kf)(this.ee),function(r,n,i,o){function a(e){var t=this;t.totalCbs=0,t.called=0,t.cbTime=0,t.end=E,t.ended=!1,t.xhrGuids={},t.lastSize=null,t.loadCaptureCalled=!1,t.params=this.params||{},t.metrics=this.metrics||{},e.addEventListener("load",(function(r){_(t,e)}),(0,O.m$)(!1)),c.IF||e.addEventListener("progress",(function(e){t.lastSize=e.loaded}),(0,O.m$)(!1))}function s(e){this.params={method:e[0]},T(this,e[1]),this.metrics={}}function u(e,n){var i=(0,t.DL)(r);i.xpid&&this.sameOrigin&&n.setRequestHeader("X-NewRelic-ID",i.xpid);var a=o.generateTracePayload(this.parsedOrigin);if(a){var s=!1;a.newrelicHeader&&(n.setRequestHeader("newrelic",a.newrelicHeader),s=!0),a.traceContextParentHeader&&(n.setRequestHeader("traceparent",a.traceContextParentHeader),a.traceContextStateHeader&&n.setRequestHeader("tracestate",a.traceContextStateHeader),s=!0),s&&(this.dt=a)}}function d(e,t){var r=this.metrics,i=e[0],o=this;if(r&&i){var a=V(i);a&&(r.txSize=a)}this.startTime=(0,p.z)(),this.listener=function(e){try{"abort"!==e.type||o.loadCaptureCalled||(o.params.aborted=!0),("load"!==e.type||o.called===o.totalCbs&&(o.onloadCalled||"function"!=typeof t.onload)&&"function"==typeof o.end)&&o.end(t)}catch(e){try{n.emit("internal-error",[e])}catch(e){}}};for(var s=0;s 1?e[1]=i:e.push(i)}else e[0]&&e[0].headers&&s(e[0].headers,n)&&(this.dt=n);function s(e,t){var r=!1;return t.newrelicHeader&&(e.set("newrelic",t.newrelicHeader),r=!0),t.traceContextParentHeader&&(e.set("traceparent",t.traceContextParentHeader),t.traceContextStateHeader&&e.set("tracestate",t.traceContextStateHeader),r=!0),r}}function x(e,t){this.params={},this.metrics={},this.startTime=(0,p.z)(),this.dt=t,e.length>=1&&(this.target=e[0]),e.length>=2&&(this.opts=e[1]);var r,n=this.opts||{},i=this.target;"string"==typeof i?r=i:"object"==typeof i&&i instanceof Y?r=i.url:c._A?.URL&&"object"==typeof i&&i instanceof URL&&(r=i.href),T(this,r);var o=(""+(i&&i instanceof Y&&i.method||n.method||"GET")).toUpperCase();this.params.method=o,this.txSize=V(n.body)||0}function A(t,r){var n;this.endTime=(0,p.z)(),this.params||(this.params={}),this.params.status=r?r.status:0,"string"==typeof this.rxSize&&this.rxSize.length>0&&(n=+this.rxSize);var o={txSize:this.txSize,rxSize:n,duration:(0,p.z)()-this.startTime};i("xhr",[this.params,o,this.startTime,this.endTime,"fetch"],this,e.D.ajax)}function E(t){var r=this.params,n=this.metrics;if(!this.ended){this.ended=!0;for(var o=0;o 2&&void 0!==arguments[2])||arguments[2];super(e,t,we.t,r),this.importAggregator()}}new class{constructor(e){let t=arguments.length>1&&void 0!==arguments[1]?arguments[1]:(0,_.ky)(16);c._A?(this.agentIdentifier=t,this.sharedAggregator=new y({agentIdentifier:this.agentIdentifier}),this.features={},this.desiredFeatures=new Set(e.features||[]),this.desiredFeatures.add(m),Object.assign(this,(0,a.j)(this.agentIdentifier,e,e.loaderType||"agent")),this.start()):(0,l.Z)("Failed to initial the agent. Could not determine the runtime environment.")}get config(){return{info:(0,t.C5)(this.agentIdentifier),init:(0,t.P_)(this.agentIdentifier),loader_config:(0,t.DL)(this.agentIdentifier),runtime:(0,t.OP)(this.agentIdentifier)}}start(){const t="features";try{const r=n(this.agentIdentifier),i=[...this.desiredFeatures];i.sort(((t,r)=>e.p[t.featureName]-e.p[r.featureName])),i.forEach((t=>{if(r[t.featureName]||t.featureName===e.D.pageViewEvent){const n=function(t){switch(t){case e.D.ajax:return[e.D.jserrors];case e.D.sessionTrace:return[e.D.ajax,e.D.pageViewEvent];case e.D.sessionReplay:return[e.D.sessionTrace];case e.D.pageViewTiming:return[e.D.pageViewEvent];default:return[]}}(t.featureName);n.every((e=>r[e]))||(0,l.Z)("".concat(t.featureName," is enabled but one or more dependent features has been disabled (").concat((0,D.P)(n),"). This may cause unintended consequences or missing data...")),this.features[t.featureName]=new t(this.agentIdentifier,this.sharedAggregator)}})),(0,T.Qy)(this.agentIdentifier,this.features,t)}catch(e){(0,l.Z)("Failed to initialize all enabled instrument classes (agent aborted) -",e);for(const e in this.features)this.features[e].abortHandler?.();const r=(0,T.fP)();return delete r.initializedAgents[this.agentIdentifier]?.api,delete r.initializedAgents[this.agentIdentifier]?.[t],delete this.sharedAggregator,r.ee?.abort(),delete r.ee?.get(this.agentIdentifier),!1}}}({features:[J,m,S,class extends h{static featureName=oe;constructor(t,r){if(super(t,r,oe,!(arguments.length>2&&void 0!==arguments[2])||arguments[2]),!c.il)return;const n=this.ee;let i;(0,k.QU)(n),this.eventsEE=(0,k.em)(n),this.eventsEE.on(se,(function(e,t){this.bstStart=(0,p.z)()})),this.eventsEE.on(ae,(function(t,r){(0,s.p)("bst",[t[0],r,this.bstStart,(0,p.z)()],void 0,e.D.sessionTrace,n)})),n.on(ce+ne,(function(e){this.time=(0,p.z)(),this.startPath=location.pathname+location.hash})),n.on(ce+ie,(function(t){(0,s.p)("bstHist",[location.pathname+location.hash,this.startPath,this.time],void 0,e.D.sessionTrace,n)}));try{i=new PerformanceObserver((t=>{const r=t.getEntries();(0,s.p)(te,[r],void 0,e.D.sessionTrace,n)})),i.observe({type:re,buffered:!0})}catch(e){}this.importAggregator({resourceObserver:i})}},C,xe,B,class extends h{static featureName=de;constructor(e,r){if(super(e,r,de,!(arguments.length>2&&void 0!==arguments[2])||arguments[2]),!c.il)return;if(!(0,t.OP)(e).xhrWrappable)return;try{this.removeOnAbort=new AbortController}catch(e){}let n,i=0;const o=this.ee.get("tracer"),a=(0,k._L)(this.ee),s=(0,k.Lg)(this.ee),u=(0,k.BV)(this.ee),d=(0,k.Kf)(this.ee),f=this.ee.get("events"),l=(0,k.u5)(this.ee),h=(0,k.QU)(this.ee),g=(0,k.Gm)(this.ee);function m(e,t){h.emit("newURL",[""+window.location,t])}function v(){i++,n=window.location.hash,this[ve]=(0,p.z)()}function b(){i--,window.location.hash!==n&&m(0,!0);var e=(0,p.z)();this[pe]=~~this[pe]+e-this[ve],this[ye]=e}function y(e,t){e.on(t,(function(){this[t]=(0,p.z)()}))}this.ee.on(ve,v),s.on(be,v),a.on(be,v),this.ee.on(ye,b),s.on(ge,b),a.on(ge,b),this.ee.buffer([ve,ye,"xhr-resolved"],this.featureName),f.buffer([ve],this.featureName),u.buffer(["setTimeout"+le,"clearTimeout"+fe,ve],this.featureName),d.buffer([ve,"new-xhr","send-xhr"+fe],this.featureName),l.buffer([me+fe,me+"-done",me+he+fe,me+he+le],this.featureName),h.buffer(["newURL"],this.featureName),g.buffer([ve],this.featureName),s.buffer(["propagate",be,ge,"executor-err","resolve"+fe],this.featureName),o.buffer([ve,"no-"+ve],this.featureName),a.buffer(["new-jsonp","cb-start","jsonp-error","jsonp-end"],this.featureName),y(l,me+fe),y(l,me+"-done"),y(a,"new-jsonp"),y(a,"jsonp-end"),y(a,"cb-start"),h.on("pushState-end",m),h.on("replaceState-end",m),window.addEventListener("hashchange",m,(0,O.m$)(!0,this.removeOnAbort?.signal)),window.addEventListener("load",m,(0,O.m$)(!0,this.removeOnAbort?.signal)),window.addEventListener("popstate",(function(){m(0,i>1)}),(0,O.m$)(!0,this.removeOnAbort?.signal)),this.abortHandler=this.#e,this.importAggregator()}#e(){this.removeOnAbort?.abort(),this.abortHandler=void 0}}],loaderType:"spa"})})(),window.NRBA=o})(); window.jQuery || document.write(' ') CKEDITOR_BASEPATH='https://f1000research.com/js/vendor/ckeditor/' window.reactTheme = 'research'; window.MathJax = { CommonHTML: { linebreaks: { automatic: true } }, 'HTML-CSS': { linebreaks: { automatic: true } }, SVG: { linebreaks: { automatic: true } }, AuthorInit: function() { MathJax.Hub.Register.MessageHook('End Process', function () { let timeout = false; // holder for timeout id const delay = 250; // delay after event is "complete" to run callback const reflowMath = function() { const dispFormulas = document.querySelectorAll('.disp-formula.panel'); if (!dispFormulas) { return; } for (const dispFormula of dispFormulas) { const child = dispFormula.querySelector('.MathJax_Preview').nextSibling.firstChild; const isMultiline = MathJax.Hub.getAllJax(dispFormula)[0].root.isMultiline; if (dispFormula.offsetWidth < child.offsetWidth || isMultiline) { MathJax.Hub.Queue(['Rerender', MathJax.Hub, dispFormula]); } } }; window.addEventListener('resize', function() { clearTimeout(timeout); // clear the timeout timeout = setTimeout(reflowMath, delay); // start timing for event "completion" }); }); }, }; if (window.location.hash == '#_=_'){ window.location = window.location.href.split('#')[0] } !function(f,b,e,v,n,t,s){if(f.fbq)return;n=f.fbq=function() {n.callMethod? n.callMethod.apply(n,arguments):n.queue.push(arguments)} ;if(!f._fbq)f._fbq=n; n.push=n;n.loaded=!0;n.version='2.0';n.queue=[];t=b.createElement(e);t.async=!0; t.src=v;s=b.getElementsByTagName(e)[0];s.parentNode.insertBefore(t,s)}(window, document,'script','https://connect.facebook.net/en_US/fbevents.js'); fbq('init', '1641728616063202'); fbq('track', "PixelInitialized", {}); (function(h,o,t,j,a,r){ h.hj=h.hj||function(){(h.hj.q=h.hj.q||[]).push(arguments)}; h._hjSettings={hjid:2318163,hjsv:6}; a=o.getElementsByTagName('head')[0]; r=o.createElement('script');r.async=1; r.src=t+h._hjSettings.hjid+j+h._hjSettings.hjsv; a.appendChild(r); })(window,document,'https://static.hotjar.com/c/hotjar-','.js?sv='); search file_upload Submit your research search menu close search Browse Gateways & Collections How to Publish Submit your Research My Submissions Article Guidelines Article Guidelines (New Versions) Open Data, Software and Code Guidelines Open Data and Accessible Source Materials Guidelines (HSS) Open Data, Software and Code Guidelines (PSE) Prepublication Checks Production Process Posters and Slides Guidelines Document Guidelines Article Processing Charges Peer Review Finding Article Reviewers About How it Works For Reviewers Our Advisors Policies Glossary FAQs For Developers Newsroom Contact My Research Submissions Content and Tracking Alerts My Details Sign In file_upload Submit your research { "@context": "https://schema.org", "@type": "ScholarlyArticle", "mainEntityOfPage": { "@type": "WebPage", "@id": "https://f1000research.com/articles/13-50" }, "headline": "Establishing the ELIXIR Microbiome Community", "datePublished": "2024-01-08T11:34:53", "dateModified": "2025-09-08T15:23:03", "author": [ { "@type": "Person", "name": "Robert D. Finn" }, { "@type": "Person", "name": "Bachir Balech" }, { "@type": "Person", "name": "Josephine Burgin" }, { "@type": "Person", "name": "Physilia Chua" }, { "@type": "Person", "name": "Erwan Corre" }, { "@type": "Person", "name": "Cymon J. Cox" }, { "@type": "Person", "name": "Claudio Donati" }, { "@type": "Person", "name": "Vitor Martins dos Santos" }, { "@type": "Person", "name": "Bruno Fosso" }, { "@type": "Person", "name": "John Hancock" }, { "@type": "Person", "name": "Katharina F. Heil" }, { "@type": "Person", "name": "Naveed Ishaque" }, { "@type": "Person", "name": "Varsha Kale" }, { "@type": "Person", "name": "Benoit J. Kunath" }, { "@type": "Person", "name": "Claudine Médigue" }, { "@type": "Person", "name": "Evangelos Pafilis" }, { "@type": "Person", "name": "Graziano Pesole" }, { "@type": "Person", "name": "Lorna Richardson" }, { "@type": "Person", "name": "Monica Santamaria" }, { "@type": "Person", "name": "Tim Van Den Bossche" }, { "@type": "Person", "name": "Juan Antonio Vizcaíno" }, { "@type": "Person", "name": "Haris Zafeiropoulos" }, { "@type": "Person", "name": "Nils P. Willassen" }, { "@type": "Person", "name": "Eric Pelletier" }, { "@type": "Person", "name": "Bérénice Batut" } ], "publisher": { "@type": "Organization", "name": "F1000Research", "logo": { "@type": "ImageObject", "url": "https://f1000research.com/img/AMP/F1000Research_image.png", "height": 480, "width": 60 } }, "image": { "@type": "ImageObject", "url": "https://f1000research.com/img/AMP/F1000Research_image.png", "height": 1200, "width": 150 }, "description": "Microbiome research has grown substantially over the past decade in terms of the range of biomes sampled, identified taxa, and the volume of data derived from the samples. In particular, experimental approaches such as metagenomics, metabarcoding, metatranscriptomics and metaproteomics have provided profound insights into the vast, hitherto unknown, microbial biodiversity. The ELIXIR Marine Metagenomics Community, initiated amongst researchers focusing on marine microbiomes, has concentrated on promoting standards around microbiome-derived sequence analysis, as well as understanding the gaps in methods and reference databases, and solutions to computational overheads of performing such analyses. Nevertheless, the methods used and the challenges faced are not confined to marine studies, but are broadly applicable to all other biomes. Thus, expanding this Community to a more inclusive ELIXIR Microbiome Community will enable it to encompass a broad range of biomes and link expertise across ‘omics technologies. Furthermore, engaging with a large number of researchers will improve the efficiency and sustainability of bioinformatics infrastructure and resources for microbiome research (standards, data, tools, workflows, training), which will enable a deeper understanding of the function and taxonomic composition of the different microbial communities." } { "@context": "http://schema.org", "@type": "BreadcrumbList", "itemListElement": [ { "@type": "ListItem", "position": "1", "item": { "@id": "https://f1000research.com/", "name": "Home" } }, { "@type": "ListItem", "position": "2", "item": { "@id": "https://f1000research.com/browse/articles", "name": "Browse" } }, { "@type": "ListItem", "position": "3", "item": { "@id": "https://f1000research.com/articles/13-50/v1", "name": "Establishing the ELIXIR Microbiome Community" } } ] } Home Browse Establishing the ELIXIR Microbiome Community ALL Metrics - Views Downloads Get PDF Get XML Cite How to cite this article Finn RD, Balech B, Burgin J et al. Establishing the ELIXIR Microbiome Community [version 1; peer review: 1 approved, 1 approved with reservations] . F1000Research 2024, 13 (ELIXIR):50 ( https://doi.org/10.12688/f1000research.144515.1 ) NOTE: If applicable, it is important to ensure the information in square brackets after the title is included in all citations of this article. Close Copy Citation Details Export Export Citation Sciwheel EndNote Ref. Manager Bibtex ProCite Sente EXPORT Select a format first Track Share ▬ ✚ Opinion Article Establishing the ELIXIR Microbiome Community [version 1; peer review: 1 approved, 1 approved with reservations] Robert D. Finn https://orcid.org/0000-0001-8626-2148 1 , Bachir Balech https://orcid.org/0000-0002-4419-0729 2 , Josephine Burgin 1 , [...] Physilia Chua https://orcid.org/0000-0001-7229-4480 3 , Erwan Corre https://orcid.org/0000-0001-6354-2278 4 , Cymon J. Cox 5 , Claudio Donati 6 , Vitor Martins dos Santos 7 , Bruno Fosso 8 , John Hancock 9 , Katharina F. Heil https://orcid.org/0000-0003-3341-3736 3 , Naveed Ishaque https://orcid.org/0000-0002-8426-901X 10 , Varsha Kale 1 , Benoit J. Kunath 11 , Claudine Médigue 12,13 , Evangelos Pafilis 14 , Graziano Pesole https://orcid.org/0000-0003-3663-0859 2,8 , Lorna Richardson https://orcid.org/0000-0002-3655-5660 1 , Monica Santamaria 15 , Tim Van Den Bossche https://orcid.org/0000-0002-5916-2587 16,17 , Juan Antonio Vizcaíno https://orcid.org/0000-0002-3905-4335 1 , Haris Zafeiropoulos https://orcid.org/0000-0002-4405-6802 14 , Nils P. Willassen 18 , Eric Pelletier https://orcid.org/0000-0003-4228-1712 13,19 , Bérénice Batut https://orcid.org/0000-0001-9852-1987 12,20 Robert D. Finn https://orcid.org/0000-0001-8626-2148 1 , Bachir Balech https://orcid.org/0000-0002-4419-0729 2 , [...] Josephine Burgin 1 , Physilia Chua https://orcid.org/0000-0001-7229-4480 3 , Erwan Corre https://orcid.org/0000-0001-6354-2278 4 , Cymon J. Cox 5 , Claudio Donati 6 , Vitor Martins dos Santos 7 , Bruno Fosso 8 , John Hancock 9 , Katharina F. Heil https://orcid.org/0000-0003-3341-3736 3 , Naveed Ishaque https://orcid.org/0000-0002-8426-901X 10 , Varsha Kale 1 , Benoit J. Kunath 11 , Claudine Médigue 12,13 , Evangelos Pafilis 14 , Graziano Pesole https://orcid.org/0000-0003-3663-0859 2,8 , Lorna Richardson https://orcid.org/0000-0002-3655-5660 1 , Monica Santamaria 15 , Tim Van Den Bossche https://orcid.org/0000-0002-5916-2587 16,17 , Juan Antonio Vizcaíno https://orcid.org/0000-0002-3905-4335 1 , Haris Zafeiropoulos https://orcid.org/0000-0002-4405-6802 14 , Nils P. Willassen 18 , Eric Pelletier https://orcid.org/0000-0003-4228-1712 13,19 , Bérénice Batut https://orcid.org/0000-0001-9852-1987 12,20 PUBLISHED 08 Jan 2024 Author details Author details 1 European Bioinformatics Institute, European Molecular Biology Laboratory, Hinxton, UK 2 Institute of Biomembranes, Bioenergetics and Molecular Biotechnologies, Bari, Italy 3 ELIXIR Hub, Hixton, UK 4 Station Biologique de Roscoff, CNRS/Sorbonne Universite, Roscoff, France 5 Centro de Ciências do Mar, Universidade do Algarve, Faro, Portugal 6 Edmund Mach Foundation Research and Innovation Centre, San Michele all'Adige, Trentino-South Tyrol, Italy 7 Systems and Synthetic Biology, Wageningen University & Research, Wageningen, Gelderland, The Netherlands 8 Department of Biosciences, Biotechnologies and Biopharmaceutics, University of Bari, Bari, Italy 9 Faculty of Medicine, University of Ljubljana, Ljubljana, Slovenia 10 Berlin Institute of Health Charité, Universitätsmedizin Berlin, Berlin, Germany 11 Luxembourg Centre for Systems Biomedicine, University of Luxembourg, Esch-sur-Alzette, Luxembourg 12 Institut Français de Bioinformatique, CNRS, Evry, France 13 Genomics Metabolics, Genoscope, Institut François-Jacob / CEA / CNRS / Université Evry / Université Paris-Saclay, Evry, France 14 Institute of Marine Biology, Biotechnology and Aquaculture, Hellenic Centre for Marine Research, Heraklion, Greece 15 Department of Soil, Plant and Food Sciences (Di.S.S.P.A.), University of Bari, Bari, Italy 16 VIB, UGent Center for Medical Biotechnology, Ghent, Belgium 17 Department of Biomolecular Medicine, Faculty of Medicine and Health Sciences, Ghent, Belgium 18 UiT The Arctic University of Norway, Tromsø, Norway 19 Research Federation for the study of Global Ocean Systems Ecology and Evolution, Paris, France 20 Bioinformatics Group, Department of Computer Science, Albert-Ludwigs-University Freiburg, Freiburg, Germany Robert D. Finn Roles: Writing – Original Draft Preparation, Writing – Review & Editing Bachir Balech Roles: Writing – Original Draft Preparation, Writing – Review & Editing Josephine Burgin Roles: Writing – Original Draft Preparation, Writing – Review & Editing Physilia Chua Roles: Writing – Original Draft Preparation, Writing – Review & Editing Erwan Corre Roles: Writing – Original Draft Preparation, Writing – Review & Editing Cymon J. Cox Roles: Writing – Original Draft Preparation, Writing – Review & Editing Claudio Donati Roles: Writing – Original Draft Preparation, Writing – Review & Editing Vitor Martins dos Santos Roles: Writing – Original Draft Preparation, Writing – Review & Editing Bruno Fosso Roles: Writing – Original Draft Preparation, Writing – Review & Editing John Hancock Roles: Writing – Original Draft Preparation, Writing – Review & Editing Katharina F. Heil Roles: Writing – Original Draft Preparation, Writing – Review & Editing Naveed Ishaque Roles: Writing – Original Draft Preparation, Writing – Review & Editing Varsha Kale Roles: Writing – Original Draft Preparation, Writing – Review & Editing Benoit J. Kunath Roles: Writing – Original Draft Preparation, Writing – Review & Editing Claudine Médigue Roles: Writing – Original Draft Preparation, Writing – Review & Editing Evangelos Pafilis Roles: Writing – Original Draft Preparation, Writing – Review & Editing Graziano Pesole Roles: Writing – Original Draft Preparation, Writing – Review & Editing Lorna Richardson Roles: Writing – Original Draft Preparation, Writing – Review & Editing Monica Santamaria Roles: Writing – Original Draft Preparation, Writing – Review & Editing Tim Van Den Bossche Roles: Writing – Original Draft Preparation, Writing – Review & Editing Juan Antonio Vizcaíno Roles: Writing – Original Draft Preparation, Writing – Review & Editing Haris Zafeiropoulos Roles: Writing – Original Draft Preparation, Writing – Review & Editing Nils P. Willassen Roles: Writing – Original Draft Preparation, Writing – Review & Editing Eric Pelletier Roles: Writing – Original Draft Preparation, Writing – Review & Editing Bérénice Batut Roles: Writing – Original Draft Preparation, Writing – Review & Editing OPEN PEER REVIEW DETAILS REVIEWER STATUS This article is included in the ELIXIR gateway. This article is included in the EMBL-EBI collection. Abstract Microbiome research has grown substantially over the past decade in terms of the range of biomes sampled, identified taxa, and the volume of data derived from the samples. In particular, experimental approaches such as metagenomics, metabarcoding, metatranscriptomics and metaproteomics have provided profound insights into the vast, hitherto unknown, microbial biodiversity. The ELIXIR Marine Metagenomics Community, initiated amongst researchers focusing on marine microbiomes, has concentrated on promoting standards around microbiome-derived sequence analysis, as well as understanding the gaps in methods and reference databases, and solutions to computational overheads of performing such analyses. Nevertheless, the methods used and the challenges faced are not confined to marine studies, but are broadly applicable to all other biomes. Thus, expanding this Community to a more inclusive ELIXIR Microbiome Community will enable it to encompass a broad range of biomes and link expertise across ‘omics technologies. Furthermore, engaging with a large number of researchers will improve the efficiency and sustainability of bioinformatics infrastructure and resources for microbiome research (standards, data, tools, workflows, training), which will enable a deeper understanding of the function and taxonomic composition of the different microbial communities. READ ALL READ LESS Keywords Microbiome, ELIXIR Community, White Paper Corresponding Author(s) Robert D. Finn ( [email protected] ) Eric Pelletier ( [email protected] ) Bérénice Batut ( [email protected] ) Close Corresponding authors: Robert D. Finn, Eric Pelletier, Bérénice Batut Competing interests: No competing interests were disclosed. Grant information: CJC received Portuguese national funds from the Foundation for Science and Technology (FCT) through projects UIDB/04326/2020, UIDP/04326/2020, and LA/P/0101/2020. T.V.D.B. acknowledges funding from the Research Foundation Flanders (FWO) [1286824N]. GP acknowledges funding from MUR (Italy), CnrBiomics (grant number PIR01_00017) and ELIXIRxNextGenIT (grant number IR0000010). JUV received the C19/BM/13684739 grant, funded by National Research Fund Luxembourg (FNR). VK was supported by a Biotechnology and Biological Sciences Research Council [BB/V01868X/1]. LR and RDF were supported by EMBL core funds. The funders had no role in study design, data collection and analysis, decision to publish, or preparation of the manuscript. Copyright: © 2024 Finn RD et al . This is an open access article distributed under the terms of the Creative Commons Attribution License , which permits unrestricted use, distribution, and reproduction in any medium, provided the original work is properly cited. How to cite: Finn RD, Balech B, Burgin J et al. Establishing the ELIXIR Microbiome Community [version 1; peer review: 1 approved, 1 approved with reservations] . F1000Research 2024, 13 (ELIXIR):50 ( https://doi.org/10.12688/f1000research.144515.1 ) First published: 08 Jan 2024, 13 (ELIXIR):50 ( https://doi.org/10.12688/f1000research.144515.1 ) Latest published: 08 Sep 2025, 13 (ELIXIR):50 ( https://doi.org/10.12688/f1000research.144515.2 )  There is a newer version of this article available. Suppress this message for one day. Introduction The term “microbiome” is a description of an entire habitat that encompasses all the microbes (bacteria, archaea, eukaryotes, and viruses), their genomes, and the environment they are found in Ref. 1 . The microbiome is experimentally characterised by the application of one or more ‘omics techniques, especially metabarcoding, metagenomics and metatranscriptomics, but also metaproteomics and metabolomics, combined with contextual metadata about the surrounding environment, be it a geographic location (e.g. ocean), host-associated (e.g. human gut) or engineered (e.g. wastewater treatment plant). Over the past decade, scientists have become increasingly aware of the role performed by microbes in the health (or maintenance) of the environment, and that dysbiosis of the microbial community can lead to dysregulation and/or negative outcomes. Microbial communities can be very diverse and heterogeneous in composition across geospatial and temporal scales, and the culture-independent methods for identifying species with the microbiome often reveal hitherto unknown microbes. Despite methodological difficulties, understanding the taxonomic and functional composition of a microbiome, how compositional differences relate to phenotypes, and how these communities may be manipulated to restore a community to a natural or normal composition are key, current research questions. Due to the diminishing costs of nucleic acid sequencing and high-availability of sequencing platforms there are now millions of microbiome-derived sequence datasets, many of which are large (gigabytes to terabytes) and complex (thousands of related samples). Additionally, datasets from other ‘omics techniques such as metaproteomics are being increasingly generated, alone or in combination with metagenomics and/or metatranscriptomics data coming from the same samples. A key challenge facing the microbiome research community is how to: appropriately store the data; informatically process, integrate, compare and interpret microbiome-derived data; and how to make the data findable, accessible, interoperable and reproducible, i.e. FAIR. 2 ELIXIR 3 is a distributed infrastructure bringing together experts from across Europe to enable life science researchers throughout the world to access and analyse life science data. ELIXIR is formed by member states each with a national Node composed of one or more centres of excellence in bioinformatics. Each Node coordinates services, standards and resources, and collaborates with experts in other Nodes to create a sustainable Europe-wide infrastructure for biological data. ELIXIR Platforms bring together experts from Nodes to develop ELIXIR’s vision and coordinate activities in defined areas. The five Platforms are Data, Tools, Interoperability, Compute and Training. ELIXIR Communities bring together experts across ELIXIR Nodes and external partners to coordinate activities within specific life science domains. The ELIXIR Marine Metagenomics Community acted as a biome-specific network of researchers for the identification and organisation of domain-specific reference resources, development of reproducible workflows and the proposal of best practices. However, there is no underlying reason to restrict these activities to just the marine environment, with most of the aforementioned efforts broadly applicable to analysis of microbiome-derived sequence data from any environment. Furthermore, there is the need to extend the activities of the Community to integrate expertise and knowledge about other ‘omics technologies, such as metaproteomics and metabolomics, which are increasingly used in microbiome studies. Thus, this white paper outlines some of the historical aspects of the Community and the aims of the broader community, especially in the context of the other ELIXIR Communities and infrastructure platforms. From marine metagenomics to a more inclusive Community The ELIXIR Marine Metagenomics Community, established in 2015 as part of the European Commission funded ELIXIR EXCELERATE project (grant number 676559), was one of the first four ELIXIR Communities created as “Use Cases”. 4 During the EXCELERATE project, these ELIXIR “Use Cases” were expanded and renamed to Communities, with a unified aim of bringing European specialists together to provide sustainable data resources, benchmark different tools and workflows, provide access to computing and storage, improve interoperability, and develop training resources within their research domains. These activities were conducted in collaboration with the ELIXIR Platforms, to ensure harmonisation of the outputs. As such, the Marine Metagenomics Community focused on metagenomics analysis pipelines, addressing the lack of reference databases and promoting the best practices for the research community. Highlights include the incorporation of new tools and resources into the MGnify 5 and MetaPIPE 6 analytical pipelines (e.g. MAPseq, ITSOneDB), 7 , 8 the formal description of the MGnify pipeline using the common workflow language (CWL 9 ) to promote interoperability, the establishment of the Marine Metagenomics Portal and MAR databases, 10 and a community paper (beyond ELIXIR) promoting best practices advocating the use of community standards for contextual provenance and metadata at all stages of the research data life cycle. 11 Capacity building has also been an important activity since the establishment of the Community, and many hands-on workshops and training courses have been developed and completed to build competence and expertise in a broader marine academic community. However, the popularity of metagenomics has continued to grow, with current approaches providing greater genome-resolved insights into the community composition and the functions performed by the microbial constituents, with annotations spanning viruses, bacteria, archaea and microbial eukaryotes. 12 – 16 Furthermore, metagenomic-like approaches are increasingly being applied to untangle complex holobiont genomes such as lichens, where both the primary symbionts and secondary non-obligate microbes are captured. 17 Finally, multi ‘omics datasets are now being more routinely produced to understand not only the genetic potential, but also the actively produced transcripts, proteins and/or metabolites, with a view to establishing the links between genotype and phenotype. When a host organism is involved, such datasets can also be augmented with genetic data from the hosts, such as genome, single nucleotide polymorphisms and transcriptomic data. The collective data facilitate a hologenomic approach 18 to understanding host phenotypes, in the context of their environment and microbiome. Given this increasing complexity of study designs, and the broad applicability of microbiome research, we advocate expanding the Marine Metagenomics Community to include other areas of microbiome research. In particular, we highlight the need for an ELIXIR Microbiome Community to develop and promote standards and research infrastructures that enable the sharing of efforts, concepts, and best practices, while benefiting from the synergistic interplay with other ELIXIR Communities. The scope of the ELIXIR Microbiome Community The term metagenomics is often colloquially applied to many different areas of microbiome research (see Table 1 ), regularly (incorrectly) used to encompass both shotgun metagenomics (indiscriminate sequencing of DNA from an environmental sample) and metabarcoding approaches (the sequencing of a specific amplified marker gene). Depending on the nature of the scientific question being addressed and/or the environment, metagenomic analysis may also involve assembly, and potentially the generation of metagenome assembled genomes (MAGs). 19 Equally applicable is the analysis of unassembled raw-read data sets that can be used for taxonomic classification (e.g. Kraken, 20 MetaPhlan, 21 mOTUs 22 ) and functional profiling approaches that are especially effective when extensive reference databases are available. Emerging sequencing technologies such as long-read sequencing methodologies and the associated adaptive sequencing techniques, 23 together with changing protocols such as host material depletion protocols (e.g. Ref. 24 ), are facilitating the analysis of a wide-range of differing communities. However, the applicability of certain downstream processing and/or analysis tools changes fundamentally in these different contexts. Similarly with metagenomic data, metatranscriptomic data can be processed in different ways, and with an associated metagenomic dataset from the same sample, enables the estimation of both the genetic potential and actively transcribed fraction. Additionally, metaproteomics, an emerging technology in microbiome research, involves the study of the entire complement of proteins expressed by microbial communities in a given environment. Fundamentally, the ELIXIR Microbiome Community is about providing the necessary infrastructures required to perform analysis of nucleotide sequence data derived from a microbiome, especially the reproducibility of the results, the archiving and discovery of analyses and the interoperability of tools and data. Given the heterogeneity of such nucleotide data, this Community will work with other ELIXIR communities to determine how microbiome-derived data coming from different ‘omics approaches, may be processed and integrated. Table 1. Overview of the terms and techniques used to study microbiome samples. Term Definition Metabarcoding Amplification and sequencing of diagnostic marker gene(s) found in a microbial community Metagenomics Random sequencing of the total DNA found in a microbial community Metatranscriptomics As with metagenomics, but the sequencing of the total RNA Metabolomics (non-targeted) Indiscriminate study of small molecules and products of metabolism Metaproteomics Identification and quantification of proteins found in a microbial community A fundamental challenge for this Community to address will be the provision of infrastructures that are sufficiently adaptable to permit the most appropriate informatics analysis, depending on the environment sampled and the experiments conducted. Moreover, when wishing to contextualise the results with similar experiments, the way a dataset has been produced and processed must be transparent to establish whether it is comparable (e.g. amplified sequence variants can only be compared when the same amplified regions are compared). Furthermore, when different methods are applied, best practices in data stewardship are required to ensure that the connectivity of the derived sequence data products, together with functional and taxonomic assertions are kept in context of the original sample/sequencing effort and associated contextual metadata. Finally, and possible unique to this ELIXIR Community, is the variety of researchers wishing to undertake microbiome research, spanning clinicians aiming to understand the role of the human microbiome in disease aetiology, ecologists wanting to understand the changing landscape of biodiversity, the agritechnology sector wishing to enhance animal and crop production, to biotechnology scientists looking for novel enzymes, among others. The context within ELIXIR Given the breadth of aforementioned applications of microbiome research, it is unsurprising that there are many links to other current and future ELIXIR activities. Figure 1 presents a schematic layout of the experimental design of a multi ‘omic analysis of a microbiome sample. Even in this very high-level representation, it can be easily observed that the new ELIXIR Microbiome Community has many potential interactions with other ELIXIR Communities and Platforms along the experimental workflow. Thus, the microbiome Community represents a showcase of the essence of ELIXIR by bringing together diverse informatics infrastructures that can be coupled together (interoperate) to achieve complex data analyses (on compute infrastructures) that have the appropriate provenance, with data adequately archived in the relevant ELIXIR core data resources. At all levels in ELIXIR, it will be essential to coordinate activities to ensure functional harmony between ELIXIR Communities using Platform-devised solutions. Figure 1. A schematic of how a microbiome sample (i.e. community in the environment) may be analysed using different ‘omics approaches, with the main steps indicated in green. Underpinning these analyses will be the metagenomic and metatranscriptomic data, which will be used as a framework for the metaproteomic and metabolomic interpretation. Highlights in this figure are connections with the ELIXIR platforms (orange boxes) and other ELIXIR communities (dark blue boxes). Interactions with other ELIXIR Communities Some of the key areas of interactions, both ongoing and foreseen, with other Communities are listed in Table 2 . As indicated in Figure 1 , the interaction with other ELIXIR Communities, specifically those concerned with environmental sampling, begins at the start of the data lifecycle, concerning the sample acquisition and characterisation of the microbial communities. For example, the Food and Nutrition Community aims to understand the relationship between food choices and human health. While microbiome analysis forms part of this Community’s activities, the aim of the Food and Nutrition Community is to integrate microbiome data within the context of food and nutrition data, host genotype and phenotype information, and develop interventions that may impact disease. 24 , 29 Thus, in the case of the Food and Nutrition Community the microbiome is only a small part of the overall research program, and restricted to human microbiome research. Members of the existing ELIXIR Marine Metagenomics Community are already engaged with the Food & Nutrition Community, and have helped to provide microbiome sequence analysis services. Similarly the Biodiversity Community has multiple overlapping activities, but with a distinct remit. For example, computational infrastructures and tools borne out of metagenomics research are now being applied for pathogen and biodiversity surveillance. Furthermore, taxonomic inventories resulting from analysis of metagenomic/metabarcoding data are commonly accepted as biodiversity resources and biodiversity resources such as GBIF and OBIS 30 routinely incorporate data from both MGnify and International Nucleotide Sequence Database Collaboration ( INSDC ). Similarly, many of the biodiversity approaches use marker gene amplification for studying environmental DNA (eDNA). While this can overlap with the metabarcoding approaches used in the Microbiome Community, eDNA analysis extends to marker genes such as Cytochrome c oxidase subunit I (Cox1) that is specific to macro-organisms and, thus, out of scope for the Microbiome Community and falls into the realm of the ELIXIR Biodiversity Community. Isolation of genomes, sequencing and their annotation is another area under the Biodiversity Community, but will encompass both microbes and macro-organisms. Nevertheless there will be common approaches used by this Community and the Microbiome Community for de novo annotation of novel genomes. Table 2. Overview of existing and planned interactions with ELIXIR Communities. Community Existing and planned interaction 3D BioInfo Increasing numbers of structural models from metagenome-derived proteins are being produced (e.g. ESMAtlas ). 25 Improving the organisation, quality control and presentation of these models will be essential to understanding structure-function relationships and improving the functional annotations of microbiomes. Biodiversity Connecting biodiversity/observation resources with ‘omics data/analysis. The metabarcoding pipelines have many similarities with those used for eDNA analysis. Accessing the genomes of novel biodiversity presents similar annotation issues as faced by the Microbiome Community. Federated human data The field hotly debates whether human microbiomes should be considered as personal and sensitive medical data. Currently, the opinion and legislation varies from country to country. The use of solutions from the Federated human data Community for sharing sensitive data may become essential to the Microbiome Community. Food & Nutrition How metagenomics techniques can be used to understand the role of the gut microbiome in unlocking nutrients in food. We will provide standard analysis workflows to this Community. Galaxy The Galaxy Community operates across Platforms and Communities. Galaxy and its community provides a graphical interface to tools, a strong workflow manager to build workflows, computational resources with the European Galaxy server, and a powerful infrastructure and resources for training. Together with the Galaxy Community, we will continue to work to tailor and expand these resources to the Microbiome Community, given the needs identified by the ongoing evaluation study. Metabolomics Develop methods and tools to connect metabolic capabilities (metagenomics) and activities (metatranscriptomics and metaproteomics) of organisms/communities with effective presence/absence of the resulting products (metabolites). Microbial Biotechnology Improving the identification of valuable enzymatic activities from environmental genomics data to identify bioactives (e.g. enzyme, small molecule) of interest for the bioeconomy, for applications in food preservation, agriculture, chemistry or medicine. Plant Science Microbes play a key role in plant health and disease. Developing a greater understanding of the needs of the Plant Science Community for microbiome-based solutions to improve plant resilience and combating pests, as well as understanding how plants maintain their microbial communities across generations. Proteomics The interpretation of (meta-)proteomics spectral data requires sample-specific reference databases. Enable the production of tailored reference databases (e.g. biome and/or other contextual metadata) and the integration of metagenomic, metatranscriptomics and metaproteomics results. Single Cell Omics Single amplified genomes (SAGs) present an alternative strategy for understanding microbial communities. There are key areas of overlap in data standards, 26 with common issues on taxonomy, gene calling and protein functional annotation. Spatial single cell data is also improving the quality of MAGs/SAGs 27 and enabling the identification of microbes in tissues and tumours. 28 Systems Biology Empowering a better integration of multi-omics environmental approaches at the community- and environment-level to describe and understand how different community members interoperate to achieve processes. With the growing number of multi ‘omics datasets, establishing strong ties with the ELIXIR Metabolomics and Proteomics communities 31 will be essential for understanding how metagenomic and metatranscriptomic data may be utilised by these Communities (e.g. the production of reference databases), and the nature of the data types produced by these other ‘omics technologies, their limitations and how the data could be integrated. For example, overlaying metabolomics results on metagenomic data is currently non-trivial due to the scarcity of small molecule annotations that can be linked to functional annotations. Ongoing work with the Microbial Biotechnology and Systems Biology Communities has identified the need to augment the functional annotation of metagenomic and metatranscriptomic data with chemical reaction information from resources such as Rhea. 32 While this will improve the discovery of new industrial applications, there is still the need to expand the protein functional annotations of the, so-called, microbial dark matter. The advent of new structural modelling software 25 , 33 and data resources 25 , 34 means that there are now structural models for millions of proteins that currently lack functional annotations, yet appear structurally related to functionally characterised proteins. Connections to the 3D BioInfo Community will aid how we store and organise this structural model information, reuse software components for visualisation and leverage their training materials on how to interpret structural model data. This will allow the Microbiome Community to assess the merits and limitations of this data type. In summary, there are many synergies and connections between the Microbiome Community and the other existing ELIXIR Communities, but none of these Communities are focused on the core issues concerning microbiome-derived sequence analysis, infrastructure provision, data standards and best practices. Moreover, there are key societal challenges such as food security, climate changes, antimicrobial resistance (and new therapeutics) and pandemic preparedness where microbiome research plays a role, yet each one of these areas is far greater in scientific scope and therefore requires the collective outputs from more than one ELIXIR Community, and reach far beyond informatics research ( Table 3 ). Table 3. Description of current national and pan-European efforts aimed at microbiome research of relevance to the ELIXIR Microbiome Community. Initiative Acronym Country Aim Mutualised Digital Spaces For Life Sciences MuDIS4LS FR The main goal of MuDiS4LS is to develop a framework that will rely on the national and regional data centres to enable scientists controlling the flow of biological data, from their origin (data-producing national infrastructures) to their public release in national or international repositories, while ensuring their mid-term security during the intermediate phases of analysis and exploitation. microGalaxy Global microGalaxy is a community of practice to (i) develop and sustain microbial data analysis in Galaxy, (ii) implement standardised “best practices”, (iii) expand documentation and training, (iv) coordinate efforts in tools, workflows and training development NFDI4Microbiota DE The mission of this consortium, part of the German NFDI (National research Data Infrastructure), is to be the central hub in Germany for supporting the microbiology community with access to data, analysis services, data/metadata standards and training. The main aims and objectives of NFDI4Microbiota are: (i) generate a broad awareness of the importance of the FAIR principles, open science and reproducible research in the microbiological community and drive a cultural change toward their widespread adaptation; (ii) equip the community with the required skills and literacy for efficient and data-driven microbial research by providing a comprehensive training program; (iii) improve the research process by mobilising, structuring and linking available data, information and knowledge related to microorganisms; (iv) support high-quality research data management by introducing professional data stewards into the microbiological research process; (v) increase the value of data by standardising and systematically collecting rich metadata and building tools for querying; (vi) make research more reproducible by standardising data processing and analysis; (vii) provide computational tools and infrastructure for the translation of data into new knowledge. European Reference Genomes Atlas ERGA Europe The European Reference Genome Atlas (ERGA) initiative is a pan-European scientific response to current threats to biodiversity. Reference genomes provide the most complete insight into the genetic basis that forms each species and represent a powerful resource in understanding how biodiversity functions. With approximately one fifth of the ~200,000 European species at risk of extinction, we need to act fast and together to generate high-quality complete genome resources on a large scale. Metaproteomics Initiative 35 Promoting dissemination of metaproteomics fundamentals, advancements, and applications through collaborative networking in microbiome research. They aim to be the central information hub and open meeting place where newcomers and experts interact to communicate, standardise and accelerate experimental and bioinformatic methodologies in this field. Secured computing spaces for the data access and analysis project of the France 2030 programme « Food Systems, Microbiome and Health » Cloud4SAMS FR The Cloud4SAMS targeted project aims to deploy a distributed digital infrastructure enabling researchers to exploit microbiome and health data in a secure computing environment. It relies on the federation of computing resources operated by different institutions and spread over different sites: datasets produced by microbiome projects (in particular those to be funded by the France 2030 programme), software tools and workflows for processing these data, computing and storage platforms suitable for processing microbiome data and matching them with health data. These resources will be indexed in the Cloud4SAMS catalogue, and will serve as building blocks to define deployment recipes describing all the procedures to instantiate a virtual machine in a secure cloud, to install the whole software environment and to transfer the datasets needed for the project. Access to these data is facilitated by an interface that manages the requests to the access committees of each project, the validation of authorizations, the ad hoc extraction of the data and their transfer to the secure spaces. Consolidation of the Italian Infrastructure for Omics Data and Bioinformatics ELIXIRxNextGenIT IT ELIXIRxNextGenIT is a national project, in continuity with CNRBiOmics, aimed at consolidating the ELIXIR-IT Infrastructure for Omics and Bioinformatics. The project is focused on data production, computational analysis, facilities improvement and human resources recruitment and training, with a view to increasing the national ELIXIR Infrastructure potential, including the capability to host new resources such as the Federated European Genome-phenome Archive (FEGA). European e-Science Infrastructure for biodiversity and ecosystem research LifeWatch ERIC EU/IT The project aims to accelerate the sharing, integration and analysis of open-data and its Virtual Research Environments (VREs) to enable studies on biodiversity structure and conservation related to multiple drivers. The LifeWatch Italy Joint Research Unit (JRU) coordinates the Italian contribution to LifeWatch ERIC, the national activities of the LifeWatch Service Centre and the LifeWatch-ITA distributed e-Biodiversity Research Institute that includes the Biomolecular, Collections, Interactions and Mediterranean thematic centers. National Research Center in Bioinformatics for Omics Sciences CNRBiOmics IT The project aims to enhance the ELIXIR Italian node infrastructure mainly in the southern regions. With its headquarters in Bari (Apulia region), it is engaged in the establishment of a “centre of excellence” for ‘omics data production, management, and analysis. The most advanced laboratory platforms for second and third generation sequencing, proteomics, metabolomics and transcriptomics are integrated with computing and storage high power platforms situated in the Bari hub and interconnected with the existing ELIXIR infrastructure. The establishment of an higher education training platform to provide the necessary skills for the infrastructure optimal use is also envisaged. Similarly, microbiome research has many translation aspects, ranging from the discovery of biomarkers associated with health and disease to industrial applications such as using enzymes from microbes or the microbes themselves for performing bioremediation and/or replacing chemical processes. One topic that is an area of intensive research is the discovery of enzymes capable of degrading plastics, typically polyethylene terephthalate (PET). 36 While metagenomic assembly and analysis is providing a rich source of new enzymes, the informatics at the core of the Microbiome Community will not provide the information why one enzyme should be assayed in preference to another, how these alpha-beta hydrolases have adapted to utilising PET, or why one enzyme performs better than another. Such answers will come from the collaborative efforts that bridge across Communities, such as microbial biotechnology and 3D BioInfo and, of course, the wider research community. Interaction with ELIXIR Platforms Similar to the collaborations with the ELIXIR Communities, there are multiple ongoing and future interactions with the ELIXIR Platforms. In the following sections the connections between the past Marine Metagenomics Community or the future Microbiome Community and each of the Platforms will be highlighted. Data The aim of the ELIXIR Data Platform is to promote the use, re-use and value of life science data. A key part of this activity has been the establishment of the Core Data Resources (CDR). Underpinning sequenced-based microbiome research is the European Nucleotide Archive (ENA), which is a recognised CDR and part of the INSDC, which in collaboration with the National Institute of Genetics DNA DataBank of Japan (DDBJ) and the United States National Center for Biotechnology’s (NCBI) GenBank and Sequence Read Archive (SRA), facilitate the deposition and global exchange of sequence data. Alongside the archived sequence data, users can access comprehensive metadata that is important to contextualise where the data originated. Throughout the lifetime of the ELIXIR Marine Metagenomics Community there have been extensive efforts to increase the standardisation of derived sequence products from metagenomic short-read datasets, particularly increasing the availability of assemblies 5 and the introduction of the deposition layers to support the increase in the numbers of MAGs being generated. 37 In the new Microbiome Community we will continue to promote and develop these layers to accommodate Eukaryotic MAGs (see below), viral sequences and complex coassembly, as well as incorporating the latest community standards as they are approved by authoritative bodies. The work undertaken to generate the MAR databases highlighted that many marine samples in ENA lack key metadata fields. Through extensive curation efforts, using literature as well as contacting the original data submitters, much of this missing data was retrieved and added to the MAR database. While ENA (or any of the INSDC partners) can not add this metadata to the original sequence record, an ELIXIR sponsored initiative led to the establishment of the Contextual Data Clearinghouse ( CDCH ). The CDCH facilitates the capture of additional metadata using controlled vocabularies including a description of how this data was generated (e.g. manual assertion, computationally derived), so that they can be associated with an INSDC record. Longer term, this data will be incorporated into BioSamples. In other non-sequenced based ‘omics fields, microbiome data archiving and analysis is supported by data-type specific resources. In the case of metaproteomics, the PRIDE database repository (also an ELIXIR CDR) enables archiving and re-analysis of (meta) proteomics data, and now also encourages researchers to upload their metadata in SDRF-format. 38 , 39 PRIDE is the leading resource of the International ProteomeXchange Consortium of proteomics data resources, involving additional databases in the USA, Japan and China, in addition to PRIDE. Similarly, in the case of metabolites the data can be deposited in the MetaboLights repository 40 or similar resources. A current challenge facing the field is connecting different multi ‘omics data that have been derived from the same sample. The Data Platform also promotes the linkage between Europe PMC 41 and other CDR databases. This is critical for the Microbiome Community as additional contextual metadata can often be found in the literature, 42 , 43 providing crucial overarching context to the experiment, which can be important for meta-analyses. We will continue to promote such approaches, enriching metadata wherever possible. Last but not least, new activities will be promoted aimed at the integration of microbiome data coming from different ‘omics approaches. In this context, recently, the PRIDE and MGnify teams developed and implemented new pipelines in both platforms for the re-analysis and integration of metagenomic and metaproteomic data, allowing the re-analysis of metaproteomics datasets from PRIDE using sequence databases generated from MGnify, and contextualising the results back into the MGnify web interface in terms of assembly annotations ( https://github.com/PRIDE-reanalysis/MetaPUF ). The ELIXIR Microbiome Community will also work to move the Marine Metagenomics domain in the RDMKit towards a more general Microbiome domain. Tools Microbiome data analysis employs a large number of tools which are used to perform basic quality control on the sequence data, with separate tools (and reference databases) typically used for taxonomic and functional profiling. Installing and managing dependencies has been eased by the use of package management systems such as Conda, or through the use of containers, e.g. Singularity. The ELIXIR Microbiome Community will increase their use of BioContainers 44 to promote the packaging, containerisation and deployment of tools relevant to microbiome research. In order to make tools findable the Community will work on improving their annotation by (i) expanding the EDAM ontology 45 to include microbiome-specific keywords, (ii) performing periodic reviews of tools and their associated annotations in the bio.tools 46 catalogue. These annotations will subsequently be used to build a catalogue of tools for microbiome data analysis and their availability for different platforms, e.g. Galaxy, or as workflow descriptions (e.g. Snakemake, CWL, Nextflow), which can be readily combined to make new annotation workflows. Additionally, the Community will develop and maintain cloud-deployable and FAIR analysis pipelines using state of the art tools and following best open science practices by: (i) using workflow descriptions; (ii) documenting the workflows and depositing them in WorkflowHub 47 for easy discovery, re-use and assessment; (iii) making them available for the Community via platforms such as MGnify and Galaxy. As an integral part of the Tools platform, Galaxy has integration with OpenEBench, WorkflowHub EDAM, bio.tools and follows all Software Best Practices. A joint effort between the Microbiome and Galaxy Communities is running an evaluation of tool requirements for microbiome data analysis in the Galaxy ecosystem. This evaluation will lead to a shared roadmap between both Communities for tool integration and standardised workflow development for microbiome data analysis. Benchmarking Very few analyses in microbiome research employ a single tool, with the norm being the coupling of multiple tools and reference databases to achieve a comprehensive analysis that includes both taxonomic and functional results. Even relatively simple workflows that perform metagenomics assembly are computationally heavy. This combination of workflow complexity and typical computational overheads has always made the routine benchmarking tools for microbiome informatics research burdensome. Nevertheless, where two or more tools perform equivalent tasks, it can be relatively simple to modify existing formally described workflows to evaluate their respective performances, but that ease often depends on where they occur in the overall workflow and the metrics used to evaluate the tool. Many efforts have tried to compare the outputs of tools and workflows (e.g. Refs. 48 – 52 ), with the Critical Assessment of Microbiome Interpretations (CAMI) having become an internationally recognised benchmarking effort. 53 – 56 The CAMI challenges have established a range of benchmarking datasets for evaluating different categories of tools. Importantly, the organisers of CAMI have engaged data generators to provide data, such that truly independent benchmarking can be undertaken. However, these benchmark datasets can become outdated over time, as the underlying data enters the reference database. The Galaxy Community has already investigated implementing benchmarking infrastructure using CAMI datasets, and increasing the awareness of this infrastructure will be a key effort across Communities and Platforms. As the Microbiome Community establishes, we will develop a broader understanding of the requirements of the Community, feed this to the Tools Platform, as well as seek opportunities to interact with the Tools Platform to capture the diversity of tools and their utility via such benchmarking activities. Compute Given that most academic institutions have access to dedicated sequencing facilities or equivalent commercial facilities, coupled with the diminishing costs of DNA sequencing and other ‘omics technologies, it is relatively easy to generate large datasets, but significant computational resources are required to store and analyse the data. Depending on the analysis being performed, the computational requirements can be very different. For example, metagenomic assembly typically requires small numbers of cores on a large memory machine, whereas some forms of raw-read analysis require many cores (hundreds) with a small memory footprint. As such, microbiome researchers need to understand the likely computational costs, and their options for deploying them on high performance computing (HPC) and cloud environments. Efforts such as Blue Cloud have helped reduce some of the barriers to using the European Open Science Cloud (EOSC) for marine research through the delivery of a collaborative virtual environment, but the range of services is limited. While such efforts help, there are still many barriers to accessing compute resources and deploying complex metagenomic pipelines in a distributed or even hybrid fashion. Working with the Compute Platform, the ELIXIR Microbiome Community will continue to investigate solutions that facilitate the execution of workflows within such distributed and/or hybrid environments, e.g. using Pulsar network, the distributed compute network offered by the Galaxy Community, and provide guidance of the likely costs of using compute infrastructures. Interoperability Previous work by the Marine Metagenomics Community has leveraged many of the ELIXIR Interoperability Platform solutions, especially the use of workflow languages for the formal description of pipelines, improving the provenance of the data outputs. As such, both the MetaPIPE and MGnify pipelines have been described using the Common Workflow Language (CWL). This effort was paralleled by MG-RAST, which also allowed MGnify and MG-RAST to exchange pipelines 57 and establish that the biological signatures reported by the respective pipelines were very similar, yet confounded by different reference databases and methodologies for assigning function. Since then, MGnify has published their workflows in WorkflowHub, 58 further promoting their discovery and reuse. As an example of reuse, the MGnify pipeline has been used as the basis for the newly developed metaGOflow pipeline, 59 to be used by the Marine Genomic Observatories. Moreover, this work also employed Research Object Crate ( RO-crate ) 60 to package relevant metadata about the sample and the bioinformatics analysis applied and the data products. RO-crate offers new opportunities for sharing or federating the metagenomics analysis workload. In parallel, Galaxy, which supports the Tool Registry Service (TRS) protocol to exchange and run workflows between the WorkflowHub and Galaxy, gained support for RO-Crate (version 23.0) to export complete data analysis as a structured and FAIR digital object, supporting the GA4GH standards, and is in the process of applying to be a Recommended Interoperability Resource. The Microbiome Community will continue to work with the Interoperability Platform to make wider use of RO-crate, with a view to federate data analysis between resources. For example, future work by the new Community will enable the MGnify workflows to be made deployable on Galaxy, with the RO-crate to be transferred, verified and ingested into MGnify. Additional work needs to be undertaken to understand how universal this approach is, so that MGnify could become a hub for a range of additional analyses, thereby reducing the duplication of effort that currently exists in the community. Finally, we will work on the development of novel mechanisms to integrate and link data coming from multi-omic approaches using different tools and data resources. This will require the development of new data Interoperability layers for data resources that are not normally focused in Microbiome data, such as the PRIDE database in the case of metaproteomics data. Training One of the key areas commonly highlighted by national and international reports on the potential of microbiome research is the need for training, especially in the area of informatics. As already highlighted, microbiome analysis is an emerging and evolving research field by itself, with plenty of challenges still to be addressed. Combined with this complexity, the increasing number of researchers using such methods makes the need for continuous training and re-training a challenge on its own. Researchers need to become familiar with modern computing technologies, such as HPC and cloud computing, and follow the constant updates on experimental approaches, algorithms (new and updates) and pipeline developments. As new pipelines are established and existing pipelines improved through the incorporation of new tools and/or reference databases, this adds further complexity to the tool and data output landscape associated with microbiome research. Platforms such as MGnify support large-scale services for most, if not all, steps of a microbiome study, meaning the distribution of raw-data, production of assemblies, their analysis, and their potential use for meta-analysis, have proved of great benefit. Nevertheless, these analyses should be considered just the starting point for further downstream analysis, which requires the specific domain expertise of the researchers involved in undertaking the study. One approach can be the use of cloud-based initiatives such as Galaxy supporting graphical interfaces and allowing the users to choose more specific tools, while tuning their parameters and reference databases according to their environment being studied. Such infrastructures attempt to fill the gap between researchers without experience in computer science and their needs for FAIR and quality microbiome analysis. Despite both solutions being readily available, there remains knowledge gaps and/or reticence about using such resources, often due to a lack of training. To upskill microbiome scientists and keep them up-to-date in microbiome data analysis and standards, the ELIXIR Microbiome Community will work in coordination with the ELIXIR Training Platform to offer scalable and FAIR training. The Microbiome Community will continue to: (i) annotate training materials with appropriate metadata to create a comprehensive training portfolio; (ii) FAIRify the training content, making it open-access; (iii) register training material, national and international providers and events in ELIXIR’s Training Portal TeSS; 61 (iv) assist the Training platform in the development of annual training gap surveys; and (v) develop materials and design learning paths specific to different community needs (e.g. biomes or data types). To enable access to training resources and deliver this training, face-to-face and online workshops will be organised and videos will be recorded for “on demand” learning. The technical infrastructure for training, in particular the computational environment setup and software installation challenge will be addressed in coordination with the ELIXIR Tools and Compute Platforms, with the aim of promoting the use of Conda environments, containers, notebooks or platforms like Galaxy which mitigate many of the current obstacles. In order to make these aspirations possible, the Community will increase its training capacity by working with training communities on practices, organising Train the Trainers events and building a community of microbiome research trainers, with areas of expertise covering different environments, ‘omics approaches and data analysis pathways. Ensuring these trainers maintain their knowledge with the evolving informatics landscape is, arguably, a key challenge that is yet to be addressed and something this Community will strive to solve in collaboration with the ELIXIR Training Platform. Context with other international initiatives We have highlighted the need for promoting best practices and standards throughout this article. However, it is also important that the Microbiome Community continues to build upon engagement with organisations such as the Genome Standards Consortium (GSC 62 ). The GSC is an international organisation aimed at making genomic data discoverable through the establishment of standards, which are derived from community input. This consortium includes stakeholders from across the data life cycle, from research scientists producing the data, to data analysts to database providers. By engaging this range of stakeholders, the GSC have become critical for establishing many of the standards that underpin genomic research. Examples of GSC established standards that are particularly pertinent to the microbiome domain include: minimal information about any sequence (MIxS 63 ), the Biological Observation Matrix (BIOM) format 64 ; and the Minimum Information about a Metagenome-Assembled Genome (MIMAG 26 , 63 ). The GSC has many ongoing projects relevant to the ELIXIR Microbiome Community, especially the M5 project (Metagenomics, Metadata, MetaAnalysis, Models and MetaInfrastructure). Combining the activities on standards concerning workflows is critical for the global microbiome community to operate with a consistent and unified view on how to make microbiome analysis reproducible and stand-up to scientific interrogation. Note, such standards do not restrict what analysis can/should be performed, but rather provide the appropriate information, that given the same starting input data, exactly the same analysis, and hence result, can be achieved. There are also other ELIXIR Node-specific initiatives that the Microbiome Community connects with to ensure that the respective efforts are synergised. Examples of projects with ELIXIR Node involvement directly related to the ELIXIR Microbiome Community are presented in Table 3 , which cover a diverse range of topics. The engagement needs to be bi-directional to ensure that the needs of Nodes are well understood and that solutions developed at national levels can be spread across the ELIXIR Microbiome Community, and vice versa. In this context, the ELIXIR Microbiome Community leads will undertake coordinating roles, engaging with the project representatives, inviting them to relevant ELIXIR events and promoting active participation in relevant ELIXIR Communities. MicrobiomeSupport, formerly a European Commission funded Community Action Support program aimed at improving microbiome research and innovation, highlighted in their final report 65 that there was “limited connectedness” in microbiome research conducted on different environments/systems, and that during the course of this program the lack of connectedness did not improve. This independent finding reinforces the need for broadening the ELIXIR Marine Metagenomics Community to a more generalist Microbiome Community. It will also be important to showcase the ELIXIR Microbiome Community to European countries that are yet to join ELIXIR. For example, Romania has a thriving microbiome research community, but is faced with the same set of informatics challenges. Sharing knowledge beyond ELIXIR, will not be the primary goal, but will nevertheless be important to harmonise the activities internationally and promote the benefits of participation in ELIXIR. Beyond Europe, there are parallel organisations that strive to achieve similar goals to ELIXIR in other locations. For example, Australia BioCommons aims to promote bioinformatics and bioscience data infrastructures at a national level. Given the strength of microbiome research in Australia (see below), we will explore opportunities for international collaboration. In addition, it will be important to showcase the ELIXIR Microbiome Community to communities (within and outside Europe) that are not yet familiarised with ELIXIR activities. For example, the Metaproteomics Initiative is an international community that promotes dissemination of metaproteomics fundamentals, advancements, and applications through collaborative networking in microbiome research. 35 , 66 For example, recently, they benchmarked metaproteomics workflows and bioinformatics methods in the field in the first multi-lab benchmark study in metaproteomics (called CAMPI), showcasing the robustness of metaproteomics data analysis workflows. 66 Finally, the National Microbiome Data Collaborative (NMDC), 67 a US led initiative, is developing a unified data portal to support microbiome multi-omics data integration and analysis through an integrated, distributed framework. Many of the governing principles associated with this portal are common with those described here, especially with the desire to have containerised, reusable computational workflows, as well as trying to make the data compliant with the FAIR principles. Sharing experiences and best practices between NMDC and the ELIXIR Microbiome Community (and others) will improve the global standardisation of microbiome research. Interaction with other key data resources beyond ELIXIR Microbiome research is global, so it is also key that European microbiome research infrastructures are coordinated with other international resources. Below we highlight a small selection of widely used resources that are produced outside Europe, and place them in context of the ELIXIR Microbiome Community. Some of the most utilised tools and resources used by the current Microbiome Community are CheckM, 68 the Genome Taxonomy Database (GTDB) and the associated GTDB toolkit. 69 , 70 CheckM is widely used to assess the completeness and contamination of prokaryotic MAGs, and is part of the GSC reporting standard. The GTDB resources is a genome based taxonomy of prokaryotes, and the associated GTDB-tk facilitates the classification of other prokaryotic genomes against this framework, more often than not, to determine novelty. These are currently made available via the Australian research groups, who face similar challenges in maintaining resources. Other key resources are based in the US, with MG-RAST 71 produced by Argonne National Laboratories and a range of different resources produced by the Joint Genome Institute (JGI). MG-RAST facilitates the analysis of raw-reads and assemblies (metabarcoding, metagenomics and metatranscriptomics), but does not perform assembly nor offer any form or long-term archiving assurance. The JGI IMG/M resource 72 has many parallels with MGnify, offering a wide range of data analyses focused on assembly and MAG generation, but IMG/M does not deal with metabarcoding. Notably, JGI also produces IMG/VR, 73 a globally unique collection of viruses, many of which have been determined from metagenomic and metatranscriptomics. Any future effort in Europe focused on viruses must aim to minimise the duplication of effort and content with IMG/VR. Coordinating with these global initiatives is key to ensure the future availability of the tools and resources, ensuring interoperability between the resources, maintaining uniform standards and sharing of the informatics/computational burden. Specific challenges and objectives of the ELIXIR Microbiome Community A key early challenge in developing the ELIXIR Microbiome Community is to establish a detailed understanding of the current approaches and databases used for the analysis of different microbiomes. For example, it is widely accepted that current short-read assembly-based methods do not generally work as well for soil microbiomes due to the diversity of the microbial community typically present (the sequence depths are insufficient to build useful contigs or the datasets are so large, that they are computationally intractable). This current limitation, has led and will continue to lead to the development of new experimental methods, from sampling through to nucleic acid sequencing and informatics analysis. In this section, some of the key challenges associated with microbiome research are highlighted below, together with how these challenges will be addressed by the new ELIXIR Microbiome Community. Table 4 lists the key thematic areas and objectives that the Microbiome Community will address, split into short-term and longer-term objectives to provide a high-level overview of the proposed Community activities. Table 4. Objectives of the ELIXIR Microbiome Community. Area Objective Near-term (2 years timeframe) Community Expansion Survey of needs, key datasets, data analysis approaches, ‘omics data types and biome specific specialisation Identify key experts involved in viral, prokaryotic and eukaryotic analysis Establish and share a strategic technical roadmap with the Communities and Platforms, highlighting key contacts Identify relevant funding calls, with the aid of building microbiome research informatics capacities and connecting to key experts in other ‘omics (e.g. metaproteomics) Training Increase awareness of microbiome tools, resources, and their applicability to different microbiomes Address knowledge gaps in generating and adopting workflows Advanced containerisation and cloud deployment Co-ordinate Increase rates of data archival deposition, with rich contextual metadata. Establish a mapping between biome and checklists Data analysis through the use of services Sharing of ideas on the design and implementation of workflows for microbiome research, promoting the use of best practices Organise in-person and virtual meetings for the Microbiome Community Industry connection Microbiome research has many applications suitable for pharmaceutical and biotechnological applications. Use ELIXIR and Node forums to understand demands and current limitations impacting this sector. Longer-term (~3-5 years) Training Targeted training for different microbiome communities Addressing the issue of maintaining “Train the Trainer” Organise hackathon to improve integration of ELIXIR services providing microbiome data Establish a rich set of training materials, appropriately tagged to aid discoverability Federated data analysis Enable the execution of MGnify pipelines in Galaxy and/or other ata management workflows, and submission of results to MGnify Establish routine mechanisms for federating microbiome analysis (e.g. RO-Crates, resources). Demonstrate approaches to multi’omics integration, through collaborative, cross-Community initiatives Promoting new approaches In conjunction with GSC, establish new standards for microbiome research, particularly with respect to data analysis reporting and contextual metadata reporting Leverage new data-types and experimental approaches to improve the scope and/or quality of microbiome analysis Enhance existing or establish new reference databases in response to Community demand and capacity Established new methods for across study comparisons, mitigating against confounding factors to enhance discovery Provide a mechanism for estimating the cost/benefit of performing different types of analysis in the context of different microbiomes International harmonisation Represent the Microbiome Community at international conferences, promoting Community/ELIXIR outputs and solutions Foster international collaborations between other resources providers and databases to ensure global harmonisation of e-infrastructures for microbiome research Leverage the CAMI initiative to facilitate benchmarking of tools and workflows The Microbiome Community will also provide a mechanism for sharing knowledge about new approaches for microbiome research, be it experimental or informatics-based techniques. For example, there is an increasing number of metagenomics datasets that are produced using long-read sequence technologies. While long-read sequencing technologies can require larger quantities of DNA or may be more error prone compared to third-generation short-read sequencing technologies – which can limit their use – the long-reads can mitigate the computational burden of metagenomic assembly and increase the confidence in analysis results (e.g. MAGs produced by long-reads can have high contiguity and therefore less prone to contamination). The long-reads can be paired with short-read sequences, which can then be used in different ways (e.g. sequence error correction). Increasing the awareness of these long-read and hybrid-sequencing approaches, the workflows that support their analysis and when and where they could be applied will be a key output of the Microbiome Community. Similarly, there are other experimental approaches such as single amplified genomes (SAGs), which have increased in popularity. The Microbiome Community will also be important for assessing the utility of emerging sequencing approaches, such as adaptive sequencing approaches. In this case, the methods can access low abundance microbes, although such methods will not facilitate the generation of abundance profiles. Bringing these data types alongside the ubiquitous short-read datasets will require new standards and data integration approaches to be developed by the Microbiome Community. There has been a paradigm-shift in metagenomic analysis with a common goal now being the generation of prokaryotic MAGs, which has not only allowed the identification of thousands of specific functions, but facilitated them to be assigned to specific organisms. As such, this has started the development of specific MAG deposition layers, 74 and the development of MAG specific resources. The new Microbiome Community will promote the use of MAG deposition, and provide guidelines and software to aid their deposition. Workflows that encompass both MAG generation and quality verification will be developed that include the capture of both prokaryotic and eukaryotic MAGs. The Microbiome Community will help establish best practices for eukaryotic MAG discovery, as well as develop new standards for removing redundancy and methods for assigning taxonomy, which are recognised gaps in the area of eukaryotic MAG discovery. While prokaryotic MAG recovery methods are more mature and standardised, it is anticipated that there will be continuous improvements in both experimental and computational methods for generating longer contigs, and more datasets that enable different approaches to enhance the detection of contamination and/or misassembly. The ELIXIR Microbiome Community will also evaluate methods and establish best practices for the identification of sub-species/strains in metagenomic datasets. To do so, we will engage with efforts such as the Critical Assessment of Metagenome Interpretation (CAMI) 53 , 75 to identify tools that can scalably and accurately classify MAGs at a finer grain taxonomic level than species. Finally, the classification and naming of MAGs is going to be paramount, so that the novel biodiversity can be understood and more easily referenced by the scientific community. Currently, the Microbiome Community has widely adopted the GTDB 69 and the associated GTDB-tk 70 for classifying MAGs against a reference tree. However, the taxonomy of GTDB differs from the more widely-used NCBI taxonomy, and there is a need to increase the interoperability between these two taxonomies. The ELIXIR Microbiome Community will work on addressing the current issues associated with MAGs and taxonomy. Additionally, another key area of development will be increasing the linkage between genomic resources and marker genes, such as the ribosomal small subunit (SSU) RNA. In addition to cellular microbes, another area for the ELIXIR Microbiome Community to address is the development of the infrastructure and resources for identifying and cataloguing viruses in metagenomic and metatranscriptomic data. 76 – 79 Viral genomes are incredibly diverse in terms of composition and organisation. Viruses, particularly those that infect bacteria, are found ubiquitously in all environments and play critical roles in community dynamics. However, there are three challenges associated with viral microbiomes: (i) there is no universal marker gene covering all viruses; (ii) viral taxonomic frameworks are incomplete; (iii) there is no centralised database collecting the millions of viral sequences; and (iv) metagenomics informatics often only produces fragments of viruses, which causes ambiguities concerning their classification and functions. It will be critical for the ELIXIR Microbiome Community to engage with established viral infrastructures and organisations, such as the European Virus Bioinformatics Center, to establish methods, standards and resources for improving the analysis of viruses found in microbiome sequence data. The increase in metagenomic assemblies has resulted in a parallel increase in the number of protein sequences that have been identified, with sets of non-redundant proteins now in the billions. There is huge potential for discovery in these protein datasets, as well as de novo designs fit for purpose, e.g. carbonic anhydrases 80 and a key aim for the new ELIXIR Microbiome Community will be ensuring that these data are annotated, both as individual sequences or as higher order grouping (e.g. pathways, biosynthetic gene clusters). This will involve the evaluation of emerging tools, as well as harnessing structural models to allow the detection of relationships that are undetectable by current sequenced based methods. The Community will need to work together to shed light on the functions of the so-called ‘Dark Matter’, develop standards for functional labelling that encapsulate both the mechanisms and confidence of the annotation, and develop new infrastructural frameworks for accessing slices of the data based on the requirements. As identified by the ELIXIR Marine Metagenomics Community, experimental and contextual metadata is critical to comparative metagenomics. The absence of rich contextual and experimental metadata limits data reuse and the production of downstream data products, such as assemblies and MAGs. With the expanded Microbiome Community, we will identify areas where metadata standards need to be improved, with biome specific contextual metadata being the most likely source of specific metadata checklist. The Community will develop training promoting the need for metadata, checking compliance against standards, how the metadata can be captured and submitted to accompany the sequence data, and potentially other ‘omics data types. Within the Community, we will develop and promote standards around the analysis provenance (analytical metadata), and how the collective corpus of metadata can be used to improve meta-analysis and the identification of confounding factors when comparing different research projects. Another key challenge that the Microbiome Community needs to address is ensuring that compute resources are accessible for performing the different forms of data analysis that can be associated with microbiome derived sequence data. Previously, we have highlighted the need for interaction with the ELIXIR Compute, Interoperability and Training platforms, as well as ELIXIR Communities such as the Galaxy Community. This requires that analysis pipelines are readily discoverable and deployable, and that key issues regarding both compute processing and storage requirements are well understood. Additionally, given that microbiome associated data analysis has such computational overheads, it is vital that models for data archiving and/or sharing are developed by the Microbiome Community to increase the capacity of microbiome research within Europe. This may require the development of new or extensions to existing databases, but it requires an agreement from the research community to adopt them. Achieving this will involve both communication and training of the microbiome research community. While there are data resources such as MGnify that provide access to consistent analyses pertaining to different metabarcoding, metatranscriptomics and metagenomics datasets from a variety of biomes, it is fundamental to remember that these data outputs do not represent the end of the analysis pathway. Typically studies require comparison between different cohort groups (disease vs health, treatment vs non-treatment). Furthermore, as the biological signal from meta’omics datasets can be extremely noisy, there can often be the need to combine datasets to boost statistical significance of the biological signal. Similarly, the combination of studies can also be used to: (i) contextualise against previous studies (e.g. similar studies on the same diseases); (ii) understand the distribution of microbes or functional features (e.g. antimicrobial resistance genes) between different geographical locations; and/or (iii) study the relationship between biomes (e.g. studies adopting a OneHealth approach). To enable such large, complex studies there needs to be a greater understanding of the approaches suitable for cross study comparisons, and their limitations. Thus, a major objective for the Microbiome Community will be to include those researchers that are developing methods that can identify and mitigate experimental and informatic confounding factors, which currently limit data reuse. Existing approaches often rely on correlating contextual and experimental metadata with statistically significant factors identified in the datasets. There is also the need to develop and promote methods for performing robust statistical analysis of microbiome derived data, thereby enabling biological signals to be extracted from cross-sample/project datasets. Currently, there is a tendency to analyse the different ‘omics datasets independently, and then correlate the derived signals. However, statistical methods are being developed to facilitate the analysis of integrated multi-omics datasets, and it will be important that the Microbiome Community determines the applicability of these approaches for microbiome research. In the context of other ‘omics approaches, there are also some major challenges in metaproteomics. 81 One of the major challenges is the construction of tailored protein sequence databases which are needed to identify proteins in complex microbial communities. Metaproteomics aims to elucidate the functional and taxonomic interplay of proteins in microbiomes, but the diversity and vast number of unknown and uncharacterized proteins present in these communities makes database creation and accurate protein identification difficult. As microbial communities are highly dynamic and their protein expression can vary significantly, conventional protein sequence databases might not cover the entire diversity, leading to potential limitations in accurate protein identification (e.g. the use of de novo sequencing). Addressing this challenge is crucial for improving the reliability and confidence of metaproteomic analysis and obtaining comprehensive insights into the functional roles of proteins in complex microbiomes. As metagenomic methods have become a more routine method for studying microbial communities, metagenomics has been and will continue to be paired with more and more diverse sets of measurements of the microbiome. Examples of non-omics data collected alongside metagenomics data include geochemical (e.g. PANGAEA 82 ) measurements, meteorological, image data and even acoustics. While methods are already emerging for the integration of ‘omics datatypes (e.g. MOFA, 81 , 83 MIA ), integration of these additional non-omics data types will enable a broader understanding of microbiomes in context. For the new Microbiome Community, it will be essential to identify the appropriate archives for these data types, and establish the methods to facilitate navigating between datasets from the same samples. Only through achieving this, can new data visualisation schemas that enable the combination of environmental, geospatial and temporal data, in addition to biological data (taxonomy/function), be developed. Conclusions The overarching aim of the Microbiome Community is to develop a sustainable bioinformatics infrastructure for microbiome resources (data, tools, workflows, standards, training) which will enable a deep understanding of the function and taxonomy of the entire microbial fraction. We aim to be biome-agnostic, yet balanced in supporting the analysis and interpretation of data from different environments. We aim to highlight the very best approaches for the analysis and integration of different data types (e.g. sequences, metabolites, proteomics, and images) and their visualisation. By broadening the Community we will engage many more researchers and aspire to have a greater representation of scientists from different disciplines, such as ecologists and clinicians, complementing the strong molecular biology and genomics backgrounds already represented in the Community. The Microbiome Community will have key roles in engaging with policy makers (e.g. access and benefit sharing, climate change impact assessment), as well as the industrial sector, which is increasing the translation of basic research to microbiome-based products (e.g. UK Microbiome Strategic Roadmap for Innovation ). Such a strong microbiome infrastructure as envisaged by this Community is essential to maximise the impact that European research programs have in the field of microbiome research, and to facilitate the exploitation of microbiome-based solutions in a range of settings, from clinical to industrial processes, thereby addressing key societal challenges and needs. Data availability No data are associated with this article . References 1. Marchesi JR, Ravel J: The vocabulary of microbiome research: a proposal. Microbiome. 2015 Jul 30; 3 : 31. PubMed Abstract | Publisher Full Text | Free Full Text 2. Wilkinson MD, Dumontier M, Aalbersberg IJ, et al. : The FAIR Guiding Principles for scientific data management and stewardship. Scientific Data. 2016 Mar 15; 3 (1): 1–9. 3. Harrow J, Drysdale R, Smith A, et al. : ELIXIR: providing a sustainable infrastructure for life science data at European scale. Bioinformatics. 2021 Jun 27; 37 (16): 2506–2511. PubMed Abstract | Publisher Full Text | Free Full Text 4. Robertsen EM, Denise H, Mitchell A, et al. : ELIXIR pilot action: Marine metagenomics – towards a domain specific set of sustainable services. F1000Res. 2017 Jan 23; 6 (70): 70. PubMed Abstract | Publisher Full Text | Free Full Text 5. Richardson L, Allen B, Baldi G, et al. : MGnify: the microbiome sequence data analysis resource in 2023. Nucleic Acids Res. 2022 Dec 7; 51 (D1): D753–D759. Publisher Full Text 6. Agafonov A, Mattila K, Tuan CD, et al. : META-pipe cloud setup and execution. F1000Res. 2017 Nov 29; 6 : 2060. PubMed Abstract | Publisher Full Text | Free Full Text 7. Matias Rodrigues JF, Schmidt TSB, Tackmann J, et al. : MAPseq: highly efficient k-mer search with confidence estimates, for rRNA sequence analysis. Bioinformatics. 2017 Aug 14; 33 (23): 3808–3810. PubMed Abstract | Publisher Full Text | Free Full Text 8. Santamaria M, Fosso B, Licciulli F, et al. : ITSoneDB: a comprehensive collection of eukaryotic ribosomal RNA Internal Transcribed Spacer 1 (ITS1) sequences. Nucleic Acids Res. 2017 Sep 25; 46 (D1): D127–D132. Publisher Full Text 9. Nebojša T, Hervé M, Stian S-R, et al. : Methods included. Commun. ACM. 2022 May 20 [cited 2023 Oct 30]; 65 : 54–63. Publisher Full Text 10. Klemetsen T, Raknes IA, Fu J, et al. : The MAR databases: development and implementation of databases specific for marine metagenomics. Nucleic Acids Res. 2018 Jan 4; 46 (D1): D692–D699. PubMed Abstract | Publisher Full Text | Free Full Text 11. Ten Hoopen P, Finn RD, Bongo LA, et al. : The metagenomic data life-cycle: standards and best practices. Gigascience. 2017 Aug 1; 6 (8): 1–11. PubMed Abstract | Publisher Full Text 12. Jégousse C, Vannier P, Groben R, et al. : A total of 219 metagenome-assembled genomes of microorganisms from Icelandic marine waters. PeerJ. 2021 Apr 2; 9 : e11112. PubMed Abstract | Publisher Full Text | Free Full Text 13. Dávila-Ramos S, Castelán-Sánchez HG, Martínez-Ávila L, et al. : A Review on Viral Metagenomics in Extreme Environments. Front. Microbiol. 2019 Oct 18; 10 : 472040. Publisher Full Text 14. Wong HL, MacLeod FI, White RA, et al. : Microbial dark matter filling the niche in hypersaline microbial mats. Microbiome. 2020 Sep 16; 8 (1): 1–14. Publisher Full Text 15. Obiol A, Giner CR, Sánchez P, et al. : A metagenomic assessment of microbial eukaryotic diversity in the global ocean. Mol. Ecol. Resour. 2020 May 1; 20 (3): 718–731. PubMed Abstract | Publisher Full Text 16. Delmont TO, Gaia M, Hinsinger DD, et al. : Functional repertoire convergence of distantly related eukaryotic plankton lineages revealed by genome-resolved metagenomics. bioRxiv. 2021 [cited 2023 Oct 30]; p. 2020.10.15.341214. Publisher Full Text 17. Tagirdzhanova G, Saary P, Cameron ES, et al. : Evidence for a core set of microbial lichen symbionts from a global survey of metagenomes. bioRxiv. 2023 [cited 2023 Oct 30]; p. 2023.02.02.524463. Publisher Full Text 18. Alberdi A, Andersen SB, Limborg MT, et al. : Disentangling host-microbiota complexity through hologenomics. Nat. Rev. Genet. 2022 May; 23 (5): 281–297. PubMed Abstract | Publisher Full Text 19. Nielsen HB, Almeida M, Juncker AS, et al. : Identification and assembly of genomes and genetic elements in complex metagenomic samples without using reference genomes. Nat. Biotechnol. 2014 Aug; 32 (8): 822–828. PubMed Abstract | Publisher Full Text 20. Lu J, Rincon N, Wood DE, et al. : Metagenome analysis using the Kraken software suite. Nat. Protoc. 2022 Dec; 17 (12): 2815–2839. PubMed Abstract | Publisher Full Text | Free Full Text 21. Beghini F, McIver LJ, Blanco-Míguez A, et al. : Integrating taxonomic, functional, and strain-level profiling of diverse microbial communities with bioBakery 3. elife. 2021 May 4; 10 . PubMed Abstract | Publisher Full Text | Free Full Text 22. Ruscheweyh HJ, Milanese A, Paoli L, et al. : Cultivation-independent genomes greatly expand taxonomic-profiling capabilities of mOTUs across various environments. Microbiome. 2022 Dec 5; 10 (1): 212. PubMed Abstract | Publisher Full Text | Free Full Text 23. Martin S, Heavens D, Lan Y, et al. : Nanopore adaptive sampling: a tool for enrichment of low abundance species in metagenomic samples. Genome Biol. 2022 Jan 24; 23 (1): 1–27. Publisher Full Text 24. Nelson MT, Pope CE, Marsh RL, et al. : Human and Extracellular DNA Depletion for Metagenomic Analysis of Complex Clinical Infection Samples Yields Optimized Viable Microbiome Profiles. Cell Rep. 2019 Feb 19; 26 (8): 2227–40.e5. PubMed Abstract | Publisher Full Text | Free Full Text 25. Lin Z, Akin H, Rao R, et al. : Evolutionary-scale prediction of atomic-level protein structure with a language model. Science. 2023 Mar 17; 379 (6637): 1123–1130. PubMed Abstract | Publisher Full Text 26. Bowers RM, Kyrpides NC, Stepanauskas R, et al. : Minimum information about a single amplified genome (MISAG) and a metagenome-assembled genome (MIMAG) of bacteria and archaea. Nat. Biotechnol. 2017 Aug 1; 35 (8): 725–731. PubMed Abstract | Publisher Full Text | Free Full Text 27. Arikawa K, Ide K, Kogawa M, et al. : Recovery of strain-resolved genomes from human microbiome through an integration framework of single-cell genomics and metagenomics. Microbiome. 2021 Oct 12; 9 (1): 1–16. Publisher Full Text 28. Ghaddar B, Biswas A, Harris C, et al. : Tumor microbiome links cellular programs and immunity in pancreatic cancer. Cancer Cell. 2022 Oct 10; 40 (10): 1240–1253.e5. PubMed Abstract | Publisher Full Text | Free Full Text 29. Balech B, Brennan L, de Santa Pau EC , et al. : The future of food and nutrition in ELIXIR. F1000Res. 2022 Aug 25; 11 (978): 978. Publisher Full Text 30. Heberling JM, Miller JT, Noesgaard D, et al. : Data integration enables global biodiversity synthesis. Proc. Natl. Acad. Sci. U. S. A. 2021 Feb 9; 118 (6). PubMed Abstract | Publisher Full Text | Free Full Text 31. Vizcaíno JA, Walzer M, Jiménez RC, et al. : A community proposal to integrate proteomics activities in ELIXIR. F1000Res. 2017 Jun 13; 6 : 875. PubMed Abstract | Publisher Full Text | Free Full Text 32. Bansal P, Morgat A, Axelsen KB, et al. : Rhea, the reaction knowledgebase in 2022. Nucleic Acids Res. 2022 Jan 7; 50 (D1): D693–D700. PubMed Abstract | Publisher Full Text | Free Full Text 33. Jumper J, Evans R, Pritzel A, et al. : Highly accurate protein structure prediction with AlphaFold. Nature. 2021 Aug; 596 (7873): 583–589. PubMed Abstract | Publisher Full Text | Free Full Text 34. Varadi M, Anyango S, Deshpande M, et al. : AlphaFold Protein Structure Database: massively expanding the structural coverage of protein-sequence space with high-accuracy models. Nucleic Acids Res. 2021 Nov 17; 50 (D1): D439–D444. Publisher Full Text 35. Van Den Bossche T, Arntzen MØ, Becher D, et al. : The Metaproteomics Initiative: a coordinated approach for propelling the functional characterization of microbiomes. Microbiome. 2021 Dec 20; 9 (1): 243. PubMed Abstract | Publisher Full Text | Free Full Text 36. Yoshida S, Hiraga K, Takehana T, et al. : Response to Comment on “A bacterium that degrades and assimilates poly (ethylene terephthalate).”. Science. 2016 Aug 19; 353 (6301): 759. PubMed Abstract | Publisher Full Text 37. Gurbich TA, Almeida A, Beracochea M, et al. : MGnify Genomes: A Resource for Biome-specific Microbial Genome Catalogues. J. Mol. Biol. 2023 Jul 15; 435 (14): 168016. PubMed Abstract | Publisher Full Text | Free Full Text 38. Claeys T, Van Den Bossche T, Perez-Riverol Y, et al. : lesSDRF is more: maximizing the value of proteomics data through streamlined metadata annotation. Nat. Commun. 2023 Oct 24; 14 (1): 1–4. Publisher Full Text 39. Dai C, Füllgrabe A, Pfeuffer J, et al. : A proteomics sample metadata representation for multiomics integration and big data analysis. Nat. Commun. 2021 Oct 6; 12 (1): 1–8. Publisher Full Text 40. Haug K, Cochrane K, Nainala VC, et al. : MetaboLights: a resource evolving in response to the needs of its scientific community. Nucleic Acids Res. 2019 Nov 6; 48 (D1): D440–D444. Publisher Full Text 41. The Europe PMC Consortium: Europe PMC: a full-text literature database for the life sciences and platform for innovation. Nucleic Acids Res. 2014 Nov 6; 43 (D1): D1042–D1048. Publisher Full Text 42. Nassar M, Rogers AB, Talo’ F, et al. : A machine learning framework for discovery and enrichment of metagenomics metadata from open access publications. Gigascience. 2022 Aug 11; 11 . PubMed Abstract | Publisher Full Text | Free Full Text 43. Zafeiropoulos H, Paragkamian S, Ninidakis S, et al. : PREGO: A Literature and Data-Mining Resource to Associate Microorganisms, Biological Processes, and Environment Types. Microorganisms. 2022 Jan 26; 10 (2). PubMed Abstract | Publisher Full Text | Free Full Text 44. Gruening B, Sallou O, Moreno P, et al. : Recommendations for the packaging and containerizing of bioinformatics software. F1000Research. 2018: 7. 45. Ison J, Kalas M, Jonassen I, et al. : EDAM: an ontology of bioinformatics operations, types of data and identifiers, topics and formats. Bioinformatics. 2013 May 15; 29 (10): 1325–1332. PubMed Abstract | Publisher Full Text | Free Full Text 46. Ison J, Rapacki K, Ménager H, et al. : Tools and data services registry: a community effort to document bioinformatics resources. Nucleic Acids Res. 2015; 44 : D38–D47. PubMed Abstract | Publisher Full Text | Free Full Text 47. Goble C, Soiland-Reyes S, Bacall F, et al. : Implementing FAIR Digital Objects in the EOSC-Life Workflow Collaboratory. Zenodo. 2021. https://zenodo.org/record/4605654 48. Lindgreen S, Adair KL, Gardner PP: An evaluation of the accuracy and speed of metagenome analysis tools. Sci. Rep. 2016 Jan 18; 6 (1): 1–14. Publisher Full Text 49. Wu Z, Wang Y, Zeng J, et al. : Constructing metagenome-assembled genomes for almost all components in a real bacterial consortium for binning benchmarking. BMC Genomics. 2022 Nov 10; 23 (1): 1–19. Publisher Full Text 50. Poussin C, Khachatryan L, Sierro N, et al. : Crowdsourced benchmarking of taxonomic metagenome profilers: lessons learned from the sbv IMPROVER Microbiomics challenge. BMC Genomics. 2022 Aug 30; 23 (1): 1–19. Publisher Full Text 51. Almeida A, Mitchell AL, Tarkowska A, et al. : Benchmarking taxonomic assignments based on 16S rRNA gene profiling of the microbiota from commonly sampled environments. Gigascience. 2018 May 11; 7 (5): giy054. Publisher Full Text 52. O’Sullivan DM, Doyle RM, Temisak S, et al. : An inter-laboratory study to investigate the impact of the bioinformatics component on microbiome analysis using mock communities. Sci. Rep. 2021 May 19; 11 (1): 10590. PubMed Abstract | Publisher Full Text | Free Full Text 53. Sczyrba A, Hofmann P, Belmann P, et al. : Critical Assessment of Metagenome Interpretation-a benchmark of metagenomics software. Nat. Methods. 2017 Nov; 14 (11): 1063–1071. PubMed Abstract | Publisher Full Text | Free Full Text 54. Fritz A, Hofmann P, Majda S, et al. : CAMISIM: simulating metagenomes and microbial communities. Microbiome. 2019 Feb 8; 7 (1): 17. PubMed Abstract | Publisher Full Text | Free Full Text 55. Meyer F, Hofmann P, Belmann P, et al. : AMBER: Assessment of Metagenome BinnERs. Gigascience. 2018 Jun 1; 7 (6). PubMed Abstract | Publisher Full Text | Free Full Text 56. Meyer F, Fritz A, Deng ZL, et al. : Critical Assessment of Metagenome Interpretation: the second round of challenges. Nat. Methods. 2022 Apr; 19 (4): 429–440. PubMed Abstract | Publisher Full Text | Free Full Text 57. Perkel JM: Workflow systems turn raw data into scientific knowledge. Nature. 2019 Sep; 573 (7772): 149–150. PubMed Abstract | Publisher Full Text 58. Implementing FAIR Digital Objects in the EOSC-Life Workflow Collaboratory: [cited 2023 Oct 30]. Reference Source 59. Zafeiropoulos H, Beracochea M, Ninidakis S, et al. : metaGOflow: a workflow for the analysis of marine Genomic Observatories shotgun metagenomics data. Gigascience. 2022 Dec 28; 12 . PubMed Abstract | Publisher Full Text | Free Full Text 60. Soiland-Reyes S, Sefton P, Crosas M, et al. : Packaging research artefacts with RO-Crate. Data Sci. 2022; 5 (2): 97–138. Publisher Full Text 61. Beard N, Bacall F, Nenadic A, et al. : TeSS: A Platform for Discovering Life-Science Training Opportunities. Bioinformatics. 2020; 36 (10): 3290–3291. PubMed Abstract | Publisher Full Text | Free Full Text 62. Field D, Amaral-Zettler L, Cochrane G, et al. : The Genomic Standards Consortium. PLoS Biol. 2011 Jun; 9 (6): e1001088. PubMed Abstract | Publisher Full Text | Free Full Text 63. Yilmaz P, Kottmann R, Field D, et al. : Minimum information about a marker gene sequence (MIMARKS) and minimum information about any (x) sequence (MIxS) specifications. Nat. Biotechnol. 2011 May 6; 29 (5): 415–420. PubMed Abstract | Publisher Full Text | Free Full Text 64. McDonald D, Clemente JC, Kuczynski J, et al. : The Biological Observation Matrix (BIOM) format or: how I learned to stop worrying and love the ome-ome. Gigascience. 2012 Jul 12; 1 (1): 2047–217X – 1–7. PubMed Abstract | Publisher Full Text | Free Full Text 65. Meisner A, Kostic T, Vernooij M, et al. : The global microbiome research landscape: mapping of research, infrastructures, policies and institutions in 2021. MicrobiomeSupport Consortium. 2022. 66. Van Den Bossche T, Kunath BJ, Schallert K, et al. : Critical Assessment of MetaProteome Investigation (CAMPI): a multi-laboratory comparison of established workflows. Nat. Commun. 2021 Dec 15; 12 (1): 1–15. Publisher Full Text 67. Eloe-Fadrosh EA, Ahmed F, Anubhav A, et al. : The National Microbiome Data Collaborative Data Portal: an integrated multi-omics microbiome data resource. Nucleic Acids Res. 2021 Oct 30; 50 (D1): D828–D836. 68. Parks DH, Imelfort M, Skennerton CT, et al. : CheckM: assessing the quality of microbial genomes recovered from isolates, single cells, and metagenomes. Genome Res. 2015 Jul; 25 (7): 1043–1055. PubMed Abstract | Publisher Full Text | Free Full Text 69. Parks DH, Chuvochina M, Rinke C, et al. : GTDB: an ongoing census of bacterial and archaeal diversity through a phylogenetically consistent, rank normalized and complete genome-based taxonomy. Nucleic Acids Res. 2021 Sep 14; 50 (D1): D785–D794. PubMed Abstract | Publisher Full Text | Free Full Text 70. Chaumeil PA, Mussig AJ, Hugenholtz P, et al. : GTDB-Tk: a toolkit to classify genomes with the Genome Taxonomy Database. Bioinformatics. 2019 Nov 15; 36 (6): 1925–1927. PubMed Abstract | Publisher Full Text 71. Keegan KP, Glass EM, Meyer F: MG-RAST, a Metagenomics Service for Analysis of Microbial Community Structure and Function. Methods Mol. Biol. 2016; 1399 : 207–233. PubMed Abstract | Publisher Full Text 72. Chen IMA, Chu K, Palaniappan K, et al. : The IMG/M data management and analysis system v.7: content updates and new features. Nucleic Acids Res. 2023 Jan 6; 51 (D1): D723–D732. PubMed Abstract | Publisher Full Text | Free Full Text 73. Camargo AP, Nayfach S, Chen IMA, et al. : IMG/VR v4: an expanded database of uncultivated virus genomes within a framework of extensive functional, taxonomic, and ecological metadata. Nucleic Acids Res. 2023 Jan 6; 51 (D1): D733–D743. PubMed Abstract | Publisher Full Text | Free Full Text 74. Amid C, Alako BTF, Balavenkataraman Kadhirvelu V, et al. : The European Nucleotide Archive in 2019. Nucleic Acids Res. 2020 Jan 8; 48 (D1): D70–D76. PubMed Abstract | Publisher Full Text 75. CAMI II: identifying best practices and issues for metagenomics software. Nat. Methods. 2022 Apr; 19 (4): 412–413. PubMed Abstract | Publisher Full Text 76. Sommers P, Chatterjee A, Varsani A, et al. : Integrating Viral Metagenomics into an Ecological Framework. Annu Rev Virol. 2021 Sep 29; 8 (1): 133–158. PubMed Abstract | Publisher Full Text 77. Roux S, Camargo AP, Coutinho FH, et al. : iPHoP: An integrated machine learning framework to maximize host prediction for metagenome-derived viruses of archaea and bacteria. PLoS Biol. 2023 Apr; 21 (4): e3002083. PubMed Abstract | Publisher Full Text | Free Full Text 78. Camargo AP, Roux S, Schulz F, et al. : Identification of mobile genetic elements with geNomad. Nat. Biotechnol. 2023 Sep 21: 1–10. Publisher Full Text 79. Genome-resolving metagenomics reveals wild western capercaillies (Tetrao urogallus) as avian hosts for antibiotic-resistance bacteria and their interactions with the gut-virome community. Microbiol. Res. 2023 Jun 1; 271 : 127372. PubMed Abstract | Publisher Full Text 80. Fredslund F, Borchert MS, Poulsen JCN, et al. : Structure of a hyperthermostable carbonic anhydrase identified from an active hydrothermal vent chimney. Enzym. Microb. Technol. 2018 Jul; 114 : 48–54. PubMed Abstract | Publisher Full Text 81. Schiebenhoefer H, Van Den Bossche T, Fuchs S, et al. : Challenges and promise at the interface of metaproteomics and genomics: an overview of recent progress in metaproteogenomic data analysis. Expert Rev. Proteomics. 2019 May; 16 (5): 375–390. PubMed Abstract | Publisher Full Text 82. Felden J, Möller L, Schindler U, et al. : PANGAEA - Data Publisher for Earth & Environmental Science. Scientific Data. 2023 Jun 2; 10 (1): 1–9. Publisher Full Text 83. Argelaguet R, Velten B, Arnol D, et al. : Multi-Omics Factor Analysis—a framework for unsupervised integration of multi-omics data sets. Mol. Syst. Biol. 2018 Jun 1; 14 (6): e8124. PubMed Abstract | Publisher Full Text | Free Full Text Comments on this article Comments (0) Version 2 VERSION 2 PUBLISHED 08 Jan 2024 ADD YOUR COMMENT Comment Author details Author details 1 European Bioinformatics Institute, European Molecular Biology Laboratory, Hinxton, UK 2 Institute of Biomembranes, Bioenergetics and Molecular Biotechnologies, Bari, Italy 3 ELIXIR Hub, Hixton, UK 4 Station Biologique de Roscoff, CNRS/Sorbonne Universite, Roscoff, France 5 Centro de Ciências do Mar, Universidade do Algarve, Faro, Portugal 6 Edmund Mach Foundation Research and Innovation Centre, San Michele all'Adige, Trentino-South Tyrol, Italy 7 Systems and Synthetic Biology, Wageningen University & Research, Wageningen, Gelderland, The Netherlands 8 Department of Biosciences, Biotechnologies and Biopharmaceutics, University of Bari, Bari, Italy 9 Faculty of Medicine, University of Ljubljana, Ljubljana, Slovenia 10 Berlin Institute of Health Charité, Universitätsmedizin Berlin, Berlin, Germany 11 Luxembourg Centre for Systems Biomedicine, University of Luxembourg, Esch-sur-Alzette, Luxembourg 12 Institut Français de Bioinformatique, CNRS, Evry, France 13 Genomics Metabolics, Genoscope, Institut François-Jacob / CEA / CNRS / Université Evry / Université Paris-Saclay, Evry, France 14 Institute of Marine Biology, Biotechnology and Aquaculture, Hellenic Centre for Marine Research, Heraklion, Greece 15 Department of Soil, Plant and Food Sciences (Di.S.S.P.A.), University of Bari, Bari, Italy 16 VIB, UGent Center for Medical Biotechnology, Ghent, Belgium 17 Department of Biomolecular Medicine, Faculty of Medicine and Health Sciences, Ghent, Belgium 18 UiT The Arctic University of Norway, Tromsø, Norway 19 Research Federation for the study of Global Ocean Systems Ecology and Evolution, Paris, France 20 Bioinformatics Group, Department of Computer Science, Albert-Ludwigs-University Freiburg, Freiburg, Germany Robert D. Finn Roles: Writing – Original Draft Preparation, Writing – Review & Editing Bachir Balech Roles: Writing – Original Draft Preparation, Writing – Review & Editing Josephine Burgin Roles: Writing – Original Draft Preparation, Writing – Review & Editing Physilia Chua Roles: Writing – Original Draft Preparation, Writing – Review & Editing Erwan Corre Roles: Writing – Original Draft Preparation, Writing – Review & Editing Cymon J. Cox Roles: Writing – Original Draft Preparation, Writing – Review & Editing Claudio Donati Roles: Writing – Original Draft Preparation, Writing – Review & Editing Vitor Martins dos Santos Roles: Writing – Original Draft Preparation, Writing – Review & Editing Bruno Fosso Roles: Writing – Original Draft Preparation, Writing – Review & Editing John Hancock Roles: Writing – Original Draft Preparation, Writing – Review & Editing Katharina F. Heil Roles: Writing – Original Draft Preparation, Writing – Review & Editing Naveed Ishaque Roles: Writing – Original Draft Preparation, Writing – Review & Editing Varsha Kale Roles: Writing – Original Draft Preparation, Writing – Review & Editing Benoit J. Kunath Roles: Writing – Original Draft Preparation, Writing – Review & Editing Claudine Médigue Roles: Writing – Original Draft Preparation, Writing – Review & Editing Evangelos Pafilis Roles: Writing – Original Draft Preparation, Writing – Review & Editing Graziano Pesole Roles: Writing – Original Draft Preparation, Writing – Review & Editing Lorna Richardson Roles: Writing – Original Draft Preparation, Writing – Review & Editing Monica Santamaria Roles: Writing – Original Draft Preparation, Writing – Review & Editing Tim Van Den Bossche Roles: Writing – Original Draft Preparation, Writing – Review & Editing Juan Antonio Vizcaíno Roles: Writing – Original Draft Preparation, Writing – Review & Editing Haris Zafeiropoulos Roles: Writing – Original Draft Preparation, Writing – Review & Editing Nils P. Willassen Roles: Writing – Original Draft Preparation, Writing – Review & Editing Eric Pelletier Roles: Writing – Original Draft Preparation, Writing – Review & Editing Bérénice Batut Roles: Writing – Original Draft Preparation, Writing – Review & Editing Competing interests No competing interests were disclosed. Grant information CJC received Portuguese national funds from the Foundation for Science and Technology (FCT) through projects UIDB/04326/2020, UIDP/04326/2020, and LA/P/0101/2020. T.V.D.B. acknowledges funding from the Research Foundation Flanders (FWO) [1286824N]. GP acknowledges funding from MUR (Italy), CnrBiomics (grant number PIR01_00017) and ELIXIRxNextGenIT (grant number IR0000010). JUV received the C19/BM/13684739 grant, funded by National Research Fund Luxembourg (FNR). VK was supported by a Biotechnology and Biological Sciences Research Council [BB/V01868X/1]. LR and RDF were supported by EMBL core funds. The funders had no role in study design, data collection and analysis, decision to publish, or preparation of the manuscript. Article Versions (2) version 2 Revised Published: 08 Sep 2025, 13:50 https://doi.org/10.12688/f1000research.144515.2 version 1 Published: 08 Jan 2024, 13:50 https://doi.org/10.12688/f1000research.144515.1 Copyright © 2024 Finn RD et al . This is an open access article distributed under the terms of the Creative Commons Attribution License , which permits unrestricted use, distribution, and reproduction in any medium, provided the original work is properly cited. Download Export To Sciwheel Bibtex EndNote ProCite Ref. Manager (RIS) Sente metrics Views Downloads F1000Research - - PubMed Central info_outline Data from PMC are received and updated monthly. - - Citations open_in_new 0 open_in_new 0 open_in_new SEE MORE DETAILS CITE how to cite this article Finn RD, Balech B, Burgin J et al. Establishing the ELIXIR Microbiome Community [version 1; peer review: 1 approved, 1 approved with reservations] . F1000Research 2024, 13 (ELIXIR):50 ( https://doi.org/10.12688/f1000research.144515.1 ) NOTE: If applicable, it is important to ensure the information in square brackets after the title is included in all citations of this article. COPY CITATION DETAILS track receive updates on this article Track an article to receive email alerts on any updates to this article. TRACK THIS ARTICLE Share Open Peer Review Current Reviewer Status: ? Key to Reviewer Statuses VIEW HIDE Approved The paper is scientifically sound in its current form and only minor, if any, improvements are suggested Approved with reservations A number of small changes, sometimes more significant revisions are required to address specific details and improve the papers academic merit. Not approved Fundamental flaws in the paper seriously undermine the findings and conclusions Version 1 VERSION 1 PUBLISHED 08 Jan 2024 Views 0 Cite How to cite this report: Heinken A. Reviewer Report For: Establishing the ELIXIR Microbiome Community [version 1; peer review: 1 approved, 1 approved with reservations] . F1000Research 2024, 13 (ELIXIR):50 ( https://doi.org/10.5256/f1000research.158321.r251099 ) The direct URL for this report is: https://f1000research.com/articles/13-50/v1#referee-response-251099 NOTE: it is important to ensure the information in square brackets after the title is included in this citation. Close Copy Citation Details Reviewer Report 11 May 2024 Almut Heinken , University of Lorraine, Lorraine, France Approved VIEWS 0 https://doi.org/10.5256/f1000research.158321.r251099 In this article, the ELIXIR Microbiome community is described. The community was established in 2015 as the Marine Metagenomics community and was recently expanded to the scope of metagenomics research for multiple areas such as human health, agriculture, and ecology. ... Continue reading READ ALL In this article, the ELIXIR Microbiome community is described. The community was established in 2015 as the Marine Metagenomics community and was recently expanded to the scope of metagenomics research for multiple areas such as human health, agriculture, and ecology. The authors then describe their perspective for the community. One focus area will be the promotion of best practices for metagenomics analyses by benchmarking tools. Interoperability between different tools will also be facilitated. Another focus will be encouraging data sharing and reuse and providing data storage platforms. Finally, links to existing ELIXIR initiatives in related areas as well as with international initiatives will be established. Overall, this is a very clear, concise, and informative review. The scope and aims of ELIXIR Microbiome are well-described and detailed. The short-term and long-term objectives are also clearly described. Specific comments: I appreciate the links to other ELIXIR initiatives in Table 2. I would be particularly interested in more detail on the integration with the Systems Biology community. How would the ability to reuse multi-omics data for systems biology approaches be improved? It is a bit unclear to me which is the proposed data platforms, tools, and training will be freely available to non-ELIXIR members. Regarding providing metadata of samples, how will GDPR regulations be handled for human samples? Is the topic of the opinion article discussed accurately in the context of the current literature? Yes Are all factual statements correct and adequately supported by citations? Yes Are arguments sufficiently supported by evidence from the published literature? Yes Are the conclusions drawn balanced and justified on the basis of the presented arguments? Yes Competing Interests: No competing interests were disclosed. Reviewer Expertise: Systems biology I confirm that I have read this submission and believe that I have an appropriate level of expertise to confirm that it is of an acceptable scientific standard. Close READ LESS CITE CITE HOW TO CITE THIS REPORT Heinken A. Reviewer Report For: Establishing the ELIXIR Microbiome Community [version 1; peer review: 1 approved, 1 approved with reservations] . F1000Research 2024, 13 (ELIXIR):50 ( https://doi.org/10.5256/f1000research.158321.r251099 ) The direct URL for this report is: https://f1000research.com/articles/13-50/v1#referee-response-251099 NOTE: it is important to ensure the information in square brackets after the title is included in all citations of this article. COPY CITATION DETAILS Report a concern Author Response 10 Sep 2025 Bérénice Batut , Bioinformatics Group, Department of Computer Science, Albert-Ludwigs-University Freiburg, Freiburg, Germany 10 Sep 2025 Author Response In this article, the ELIXIR Microbiome community is described. The community was established in 2015 as the Marine Metagenomics community and was recently expanded to the scope of metagenomics research ... Continue reading In this article, the ELIXIR Microbiome community is described. The community was established in 2015 as the Marine Metagenomics community and was recently expanded to the scope of metagenomics research for multiple areas such as human health, agriculture, and ecology. The authors then describe their perspective for the community. One focus area will be the promotion of best practices for metagenomics analyses by benchmarking tools. Interoperability between different tools will also be facilitated. Another focus will be encouraging data sharing and reuse and providing data storage platforms. Finally, links to existing ELIXIR initiatives in related areas as well as with international initiatives will be established. Overall, this is a very clear, concise, and informative review. The scope and aims of ELIXIR Microbiome are well-described and detailed. The short-term and long-term objectives are also clearly described. We would like to thank the reviewer for their positive comments concerning the ELIXIR microbiome community papers. Below we address their specific comments. Specific comments: I appreciate the links to other ELIXIR initiatives in Table 2. I would be particularly interested in more detail on the integration with the Systems Biology community. How would the ability to reuse multi-omics data for systems biology approaches be improved? We have added a few sentences to expand how the integration might be improved. It is a bit unclear to me which is the proposed data platforms, tools, and training will be freely available to non-ELIXIR members. All of the proposed activities are freely available to non-ELIXIR members. ELIXIR does not put explicit boundaries on who can use, but engagement with some ELIXIR events may be preferentially given to scientists coming from ELIXIR member states and funds from ELIXIR funding schemes would be restricted to member states. Regarding providing metadata of samples, how will GDPR regulations be handled for human samples? The landscape concerning human microbiome samples is complicated. Currently there is little consensus across Europe whether human microbiomes should be under controlled access. Similarly, the GDPR landscape is also complicated and not entirely independent. There are already established routes for suppression of data (should an individual wish to be forgotten), which can be propagated to other databases (e.g. MGnify will remove analyses associated with a suppressed sequence dataset). This is clearly an area that will need to be developed as part of the ELIXR community, as no specific solutions have been agreed, we would prefer not to comment about how they will be handled in this manuscript. In this article, the ELIXIR Microbiome community is described. The community was established in 2015 as the Marine Metagenomics community and was recently expanded to the scope of metagenomics research for multiple areas such as human health, agriculture, and ecology. The authors then describe their perspective for the community. One focus area will be the promotion of best practices for metagenomics analyses by benchmarking tools. Interoperability between different tools will also be facilitated. Another focus will be encouraging data sharing and reuse and providing data storage platforms. Finally, links to existing ELIXIR initiatives in related areas as well as with international initiatives will be established. Overall, this is a very clear, concise, and informative review. The scope and aims of ELIXIR Microbiome are well-described and detailed. The short-term and long-term objectives are also clearly described. We would like to thank the reviewer for their positive comments concerning the ELIXIR microbiome community papers. Below we address their specific comments. Specific comments: I appreciate the links to other ELIXIR initiatives in Table 2. I would be particularly interested in more detail on the integration with the Systems Biology community. How would the ability to reuse multi-omics data for systems biology approaches be improved? We have added a few sentences to expand how the integration might be improved. It is a bit unclear to me which is the proposed data platforms, tools, and training will be freely available to non-ELIXIR members. All of the proposed activities are freely available to non-ELIXIR members. ELIXIR does not put explicit boundaries on who can use, but engagement with some ELIXIR events may be preferentially given to scientists coming from ELIXIR member states and funds from ELIXIR funding schemes would be restricted to member states. Regarding providing metadata of samples, how will GDPR regulations be handled for human samples? The landscape concerning human microbiome samples is complicated. Currently there is little consensus across Europe whether human microbiomes should be under controlled access. Similarly, the GDPR landscape is also complicated and not entirely independent. There are already established routes for suppression of data (should an individual wish to be forgotten), which can be propagated to other databases (e.g. MGnify will remove analyses associated with a suppressed sequence dataset). This is clearly an area that will need to be developed as part of the ELIXR community, as no specific solutions have been agreed, we would prefer not to comment about how they will be handled in this manuscript. Competing Interests: No competing interests were disclosed. Close Report a concern Respond or Comment COMMENTS ON THIS REPORT Author Response 10 Sep 2025 Bérénice Batut , Bioinformatics Group, Department of Computer Science, Albert-Ludwigs-University Freiburg, Freiburg, Germany 10 Sep 2025 Author Response In this article, the ELIXIR Microbiome community is described. The community was established in 2015 as the Marine Metagenomics community and was recently expanded to the scope of metagenomics research ... Continue reading In this article, the ELIXIR Microbiome community is described. The community was established in 2015 as the Marine Metagenomics community and was recently expanded to the scope of metagenomics research for multiple areas such as human health, agriculture, and ecology. The authors then describe their perspective for the community. One focus area will be the promotion of best practices for metagenomics analyses by benchmarking tools. Interoperability between different tools will also be facilitated. Another focus will be encouraging data sharing and reuse and providing data storage platforms. Finally, links to existing ELIXIR initiatives in related areas as well as with international initiatives will be established. Overall, this is a very clear, concise, and informative review. The scope and aims of ELIXIR Microbiome are well-described and detailed. The short-term and long-term objectives are also clearly described. We would like to thank the reviewer for their positive comments concerning the ELIXIR microbiome community papers. Below we address their specific comments. Specific comments: I appreciate the links to other ELIXIR initiatives in Table 2. I would be particularly interested in more detail on the integration with the Systems Biology community. How would the ability to reuse multi-omics data for systems biology approaches be improved? We have added a few sentences to expand how the integration might be improved. It is a bit unclear to me which is the proposed data platforms, tools, and training will be freely available to non-ELIXIR members. All of the proposed activities are freely available to non-ELIXIR members. ELIXIR does not put explicit boundaries on who can use, but engagement with some ELIXIR events may be preferentially given to scientists coming from ELIXIR member states and funds from ELIXIR funding schemes would be restricted to member states. Regarding providing metadata of samples, how will GDPR regulations be handled for human samples? The landscape concerning human microbiome samples is complicated. Currently there is little consensus across Europe whether human microbiomes should be under controlled access. Similarly, the GDPR landscape is also complicated and not entirely independent. There are already established routes for suppression of data (should an individual wish to be forgotten), which can be propagated to other databases (e.g. MGnify will remove analyses associated with a suppressed sequence dataset). This is clearly an area that will need to be developed as part of the ELIXR community, as no specific solutions have been agreed, we would prefer not to comment about how they will be handled in this manuscript. In this article, the ELIXIR Microbiome community is described. The community was established in 2015 as the Marine Metagenomics community and was recently expanded to the scope of metagenomics research for multiple areas such as human health, agriculture, and ecology. The authors then describe their perspective for the community. One focus area will be the promotion of best practices for metagenomics analyses by benchmarking tools. Interoperability between different tools will also be facilitated. Another focus will be encouraging data sharing and reuse and providing data storage platforms. Finally, links to existing ELIXIR initiatives in related areas as well as with international initiatives will be established. Overall, this is a very clear, concise, and informative review. The scope and aims of ELIXIR Microbiome are well-described and detailed. The short-term and long-term objectives are also clearly described. We would like to thank the reviewer for their positive comments concerning the ELIXIR microbiome community papers. Below we address their specific comments. Specific comments: I appreciate the links to other ELIXIR initiatives in Table 2. I would be particularly interested in more detail on the integration with the Systems Biology community. How would the ability to reuse multi-omics data for systems biology approaches be improved? We have added a few sentences to expand how the integration might be improved. It is a bit unclear to me which is the proposed data platforms, tools, and training will be freely available to non-ELIXIR members. All of the proposed activities are freely available to non-ELIXIR members. ELIXIR does not put explicit boundaries on who can use, but engagement with some ELIXIR events may be preferentially given to scientists coming from ELIXIR member states and funds from ELIXIR funding schemes would be restricted to member states. Regarding providing metadata of samples, how will GDPR regulations be handled for human samples? The landscape concerning human microbiome samples is complicated. Currently there is little consensus across Europe whether human microbiomes should be under controlled access. Similarly, the GDPR landscape is also complicated and not entirely independent. There are already established routes for suppression of data (should an individual wish to be forgotten), which can be propagated to other databases (e.g. MGnify will remove analyses associated with a suppressed sequence dataset). This is clearly an area that will need to be developed as part of the ELIXR community, as no specific solutions have been agreed, we would prefer not to comment about how they will be handled in this manuscript. Competing Interests: No competing interests were disclosed. Close Report a concern COMMENT ON THIS REPORT Views 0 Cite How to cite this report: Pauvert C. Reviewer Report For: Establishing the ELIXIR Microbiome Community [version 1; peer review: 1 approved, 1 approved with reservations] . F1000Research 2024, 13 (ELIXIR):50 ( https://doi.org/10.5256/f1000research.158321.r244855 ) The direct URL for this report is: https://f1000research.com/articles/13-50/v1#referee-response-244855 NOTE: it is important to ensure the information in square brackets after the title is included in this citation. Close Copy Citation Details Reviewer Report 14 Mar 2024 Charlie Pauvert , University Hospital of RWTH, Aachen, Germany Approved with Reservations VIEWS 0 https://doi.org/10.5256/f1000research.158321.r244855 The authors make the case for mutating the ELIXIR Marine Metagenomics Community initiative into an ELIXIR Microbiome Community. It is indeed timely to have an such a pan-European initiative especially as it draws on previous experience, albeit a more narrow ... Continue reading READ ALL The authors make the case for mutating the ELIXIR Marine Metagenomics Community initiative into an ELIXIR Microbiome Community. It is indeed timely to have an such a pan-European initiative especially as it draws on previous experience, albeit a more narrow biome. The emphasis on multi-omics integration is also a strong suit as it is a complex topic where the microbiome research community would need support and infrastructure. The authors map out existing (but not all) similar initiatives, even if it is unclear how they would work together. In my opinion, a couple of claims should be strengthened and the argumentative power of this white paper would benefit from minor reorganization, a text slim down and extra proofreading. To this end, I highlight the strengths and weaknesses of the manuscript below. ## Strengths - Table 3 is a great idea to map out the others initiatives but needs a bit rework to be straight to the point. - In the Data section, the part between "Throughout the lifetime of the ELIXIR Marine Metagenomics Community" and the end of the paragraph is really relevant and draws from the experience of the ELIXIR Community. These sentences should be more highlighted, maybe upstream. I do wonder how the authors considered how the others initiatives (mentioned in Table 3) could also contribute to the addition of metadata for a global curation effort, possibly via the mentioned Contextual Data Clearinghouse. - The data re-analysis and integration between PRIDE and MGnify (named MetaPUF) is a good use case of the efforts that the new ELIXIR Community can promote. - I did appreciate the clear objective at the end of the Compute section, to help with scaling up analyses as well as the plan for the interaction with the ELIXIR Training Platform which promise to give a boost in training the current and next generation of microbiome researchers. - I liked how the authors highlighted (using the example of viral sequences databases) that time and resources should not be wasted in duplicating works, and that initiatives such as the ELIXIR Microbiome Community can promote this. - The authors thought ahead to promote and develop methods to limit confounding factors when re-using datasets, especially in multi-omics settings, and added this to one of the ELIXIR Microbiome Community challenges. ## Weaknesses ### Introduction - (minor) The microbiome definition used in the paper was revisited in Review reference 1 to include the interactions as well as others molecules than genomes as part of the microbiome. I suggest to use this definition given the emphasis on additional omics made in the paper. - (major) The first paragraph ends on challenges of our research community including how to make data FAIR. Given the experience gained by the ELIXIR Marine Metagenomics Community, I would have appreciated a sentence on why making our data FAIR is important (e.g., transparency to make the research more reproducible, accountability given the (public) source of funding). See major comment in next section. ### The scope of the ELIXIR Microbiome Community - (major) These following arguments for FAIR data should be in the Introduction section. "Moreover, when wishing to contextualise the results with similar experiments, the way a dataset has been produced and processed must be transparent to establish whether it is comparable (e.g. amplified sequence variants can only be compared when the same amplified regions are compared). Furthermore, when different methods are applied, best practices in data stewardship are required to ensure that the connectivity of the derived sequence data products, together with functional and taxonomic assertions are kept in context of the original sample/sequencing effort and associated contextual metadata." - (minor) The very first sentence of this section blames researchers for misuse of a term "[...] regularly (incorrectly) used [...]". I suggest to rephrase the sentence to simply state the differences, or back up the misuse of the term with references or a survey. - (minor) "Fundamentally, the ELIXIR Microbiome Community is about providing the necessary infrastructures required to perform analysis of nucleotide sequence data derived from a microbiome" - (minor) "Finally, and possible unique to this ELIXIR Community, is the variety of researchers". I am not sure that these features are unique to the ELIXIR Community but rather to the field of microbiome research. I suggest to tone it down by simply pointing out that the ELIXIR Community gathers a variety of researchers. ### Table 1 - (major) Metabolomics is listed in Table 1 but is omitted from the narrative in the paragraph. - (major) The Table 1 is not really used but could actually reduce some redundancies in the manuscript by providing a one-stop-shop to explain and detail these techniques. ### Figure 1 - (major) The figure is early on in the manuscript and it is unclear which of the ELIXIR platforms and communities are already established, or foreseen. Especially since "current and future ELIXIR activities" is mentioned in the manuscript before referencing this figure. The Table 2 does not add more information in that regard. - (major) There is a need for different types of arrows as the same arrow represent interactions between communities or processes. - (minor) The "Data" node is quite generic, I was wondering whether it meant public repositories, institute repositories or both. Please be precise. - (minor) I guess the type of metadata illustrated in the figure is restricted to biological metadata, that is indeed collected during sampling, however, technical metadata such as the sequencing method used, the type of instrument or the library layout, are collected during the processing of the sample not only the sampling itself. - (minor) There should be an arrow from "Archive" back to "Data" when the data produced is deposited and then contributes back to public repositories? ### Interactions with other ELIXIR Communities * (major) The sentence "Similarly, many of the biodiversity approaches use marker gene amplification for studying environmental DNA (eDNA)" is redundant with the previous one and introduces a different nomenclature that was not used before (marker gene amplication vs metabarcoding) and the differences, if any, are not explained. I would suggest to remove the sentence. * (major) The term "Isolation of genomes" is misleading and I guess the authors used it as a shorthand for "The isolation of bacteria, its DNA extraction, genome sequencing and their annotation". Please rephrase to avoid misinterpretation. Plus I would argue that these steps are also done when deconstructing microbiomes via cultivation strategies and are therefore not so out-of-scope. * (minor) "yet each one of these areas is far greater in scientific scope" feels exaggerated and vague. It should be rephrased. A suggestion is: "yet each one of these areas is too complex to be tacked individually" * (minor) There is no link nor transition between the paragraph that starts with "In summary, " before the Table 3 and the paragraph after that starts with "Similarly, microbiome research". Please rephrase or edit to connect the two sections. Plus, whilst this is good to have a concrete example in the "Similarly" paragraph, the paragraph before was very broad and doing a summary. I would suggest to try reordering the two paragraphs and bring the example earlier for a smoother transition. * (minor) Typo "Similarly, microbiome research has many translation al aspects" * (minor) The PET acronym is detailed but actually used only once, I am unsure if this is necessary. ### Table 2 - (major) The text in the "Interaction" column needs rework as it does not use a consistent wording and could be more to the point, especially as it is a complement to the main text: - The Food and Nutrition entry is a question. - The Galaxy entry has an unnecessary return carriage, and a unspecific "ongoing evaluation study" that needs to be clarified. - The Plant Science entry starts with a generic sentence that could be in the introduction or removed for clarity. Plus I would add a clarification that "plants maintain or not their microbial communities across generations." - (minor) The status (e.g., currently active, inactive, planned, etc.) of each ELIXIR Communities would have been appreciated as it is missing also from the Figure 1. - (minor) The mention of "the field" for the Federated human data community is vague in a manuscript about gathering communities, which research field is implied? If this is the human microbiome research field as a whole, please indicate. ### Table 3 - (major) Similarly to the Table 2 major comment, there is a lack of consistent wording in the "Aim" column that makes the Table 3 not as impactful as it should and could be. Maybe the authors could extract common/distinct features from each of these initiatives as an alternative way to the "Aim" column. A couple of suggestions for these features would be: is it a national initiative?, does it relates to data storage, data analysis, training? is it linked to ELIXIR? - (major) I was surprised not to see the NMDC listed in the table 3, especially when it is discussed in the main text. What about the NCCR Microbiomes initiative in Switzerland? I can understand that some initiatives are not included for space reason, but maybe state it in the legend of the table. - (minor) Is this table sorted? It seems not, but it could be by acronyms or names. - (minor) The aim for the NFDI4Microbiota is way too big a paragraph. The authors should reduce it for conciseness. - (minor) The Metaproteomics Initiative entry has a hyperlink and a reference when none of the others have. Please homogenise. - (minor) Some entries have country listed and some not. Please homogenise. - (minor) I am not questioning the existence of the European Reference Genome Atlas here, but how best to phrase its relevance to ELIXIR in the manuscript. There seems to have no prokaryotes genomes in their atlas, however, there seems to be a trove of fungi and protists genomes which are usually said to be understudied in microbiome. So I think there is a missed opportunity for the authors here to make the most out of this entry. - (minor) "With its headquarters in Bari (Apulia region)," seems irrelevant in the context of the table, please remove. ### Data - (minor) The end of the following sentence is redundant as the INSDC was introduced earlier already "INSDC, which in collaboration with the National Institute of Genetics DNA DataBank of Japan (DDBJ) and the United States National Center for Biotechnology’s (NCBI) GenBank and Sequence Read Archive (SRA), facilitate the deposition and global exchange of sequence data.". Please adjust accordingly. - (minor) "A current challenge facing the field is connecting different multi ‘omics data that have been derived from the same sample." Is this going to be tackled by the ELIXIR Microbiome Community? If so, I would state that this is part of its objectives. - (minor) The end of sentence "[...]overarching context to the experiment, which can be important for meta-analyses." seems like an euphemism, I would suggest to replace with "[...]overarching context to the experiment, re-analyses or meta-analyses." to include re-analyses as well. - (minor) "We will continue to promote such approaches, enriching metadata wherever possible." Is this going to be done via the CDCH? - (minor) "The ELIXIR Microbiome Community will also work to move the Marine Metagenomics domain in the RDMKit towards a more general Microbiome domain." What is the RDMKit? It is not explained, nor cited nor mentioned again. ### Tools - (minor) "will increase their use of BioContainers" should be "will increase the use of BioContainers" - (minor) "In order to make tools findable by the end users, the Community" - (minor) "workflow descriptions (e.g. Snakemake, CWL, Nextflow)" None of them have their references cited, is it an omission or space limitation? - (minor) "A current joint effort between the Microbiome and Galaxy Communities" ### Benchmarking - (major) "Benchmarking" it is the only item at this hierarchy level, meaning that this is useless for structuring the text. Please edit. - (minor) Review reference 2 published a recent review with guidelines to learn from that could have its place in this paragraph. - (minor) in the sentence: "As the Microbiome Community establishes, we will develop a broader understanding of the requirements of the Community , feed this to the Tools Platform, as well as seek opportunities to interact with the Tools Platform to capture the diversity of tools and their utility via such benchmarking activities.", is this the ELIXIR Microbiome Community, or the wide community of microbiome researchers? Is it to mean that the ELIXIR Microbiome Community is going to act as an interface between microbiome researchers and ELIXIR Tools/Infrastructure? ### Compute - (minor) The first sentence would fit better in the introduction. ### Interoperability - (minor) The reference 57 should be at the end of the sentence starting with "This effort was paralleled" not in the middle. - (minor) Reference 58 should be removed at it is a duplicate of reference 47. - (minor) Use the full text "Global Alliance for Genomics and Health" instead of GA4GH. - (minor) "is in the process of applying to be a n ELIXIR Recommended Interoperability Resource." - (minor) In the sentence: "This will require the development of new data Interoperability layers for data resources that are not normally focused in Microbiome data" I think the authors meant "used" instead of "focused", and "microbiome" instead of "Microbiome". ### Training - (minor) In "Platforms such as MGnify support large-scale services for most, if not all, steps of a microbiome study", I would suggest to remove the "if not all". - (minor) "with areas of expertise covering different environments, ‘omics approaches and data analysis pathways." I think the authors meant "learning paths", and I would refrain from using "pathways" as it also has a biological meaning. ### Context with other international initiatives - (major) Whilst I appreciated the emphasis that no initiative exists on its own, I feel the first paragraph on the Genom ic (please correct the typo) Standards Consortium feels lengthy for a manuscript whose topic is not the GSC. I would advise to summarize. In this respect, the second paragraph is particularly relevant to a tangible collaboration between ELIXIR Microbiome Community and GSC. - (major) The NMDC is discussed in this section but not part of the Table 3. - (minor) Is the mentioned M5 project still active as the website's last update is 2012? - (minor) "Combining the activities on standards concerning workflows [...]" does this means adding and providing workflows to the microbiome research community? - (minor) Given the emphasis on the fact that MicrobiomeSupport was a program, the authors could update the readers and indicate that it is now MicrobiomeSupport Association. ### Interaction with other key data resources beyond ELIXIR - (major) There is an order issue with the main text that a proofread could solve, as MG-RAST is explained and cited in the first paragraph but already mentioned upstream of the main text in the Interoperability section. ### Specific challenges and objectives of the ELIXIR Microbiome Community - (major) The strong claim "it is widely accepted that current short-read assembly-based methods do not generally work as well for soil microbiomes" would probably need at least one reference. - (major) "(iii) there is no centralised database collecting the millions of viral sequences". It seems to be the case indeed, and there are databases (~24) out there as recently compiled in Review reference 3. How ELIXIR Microbiome Community plans to integrate/aggregate these resources in a non-duplicating manner? - (minor) The word "through" is superfluous in the the sentence that starts with "This current limitation, [...]" and can be removed. - (minor) The CAMI was already explained and cited above, so the already defined acronym can be used. - (minor) Precise the area in :"Additionally, another key area of development of taxonomy [...]". - (minor) The sentence "Viruses, particularly those that infect bacteria, are found ubiquitously in all environments and play critical roles in community dynamics." belongs in an introduction, not so downstream of the manuscript. - (minor) The sentence starting the sixth paragraph could be precised as "The increase in metagenomic assemblies has resulted in a parallel increase in the number of predicted protein sequences, with sets of non-redundant proteins now in the billions." - (minor) Fix typo in "that are undetectable by current sequence based methods." - (minor) Would it make sense to also be able to access representative, of clusters for instance? "[...] develop new infrastructural frameworks for accessing slices of the data or adequate representatives based on the requirements." - (minor) What is the " expanded Microbiome Community"? It was never mentioned before. - (minor) A few comments on the sentence: "Within the Community, we will develop and promote standards around the analysis provenance (analytical metadata),". In my opinion and how it was already stated in the manuscript, it would make more sense to promote existing standards first and then develop if need be. There is no mention of other type of metadata, so the "analytical metadata" precision seems superfluous. I would suggest: "Within the Community, we will promote and develop standards regarding the analysis provenance," - (minor) Precise the term forms in "[...] ensuring that comput ing resources are accessible for performing the different forms of data analysis [...]", do the authors mean types of/steps in the data analysis? - (minor) Same argument as before regarding reinventing the wheel, I would swap the part of the sentence: "This may require the extensions to existing databases or development of new ones , but it requires an agreement from the research community to adopt them." - (minor) This part "Metaproteomics aims to elucidate the functional and taxonomic interplay of proteins in microbiomes," should have been in the Table 1, or to reuse the Table 1 here. - (minor) The tenth paragraph of this section starts with the mention of multiple major challenges, but detail "only" one of them. I would suggest to mention some of the others challenges. - (minor) The "MIA" method is just a hyperlink, without any reference. Either cite the website accordingly or add the reference. ### Table 4 - (major) I was surprised to see that the objective "Foster international collaborations between other resources providers and databases to ensure global harmonisation of e-infrastructures for microbiome research" was long-term, as I would have imagined that a gap analysis would be short-term to ensure we do not reinvent the wheel, especially given the others initiatives discussed in the manuscript. - (minor) If the Objective column starts with action verbs (which is a good idea), then it should be "Survey the needs" instead of "Survey of needs". - (minor) Specify the type of workflow with "Address knowledge gaps in generating and adopting data analysis workflows" - (minor) The verb is missing in " Teach a dvanced containerisation and cloud deployment" - (minor) The verb is missing in " Promote data analysis through the use of services". Is this ELIXIR services in general or specific ELIXIR Microbiome Community services? - (minor) Use "Share" instead of "Sharing" in the "Co-ordinate" entry. - (minor) The "Industry connection" entry does not fit the action verb pattern. A suggestion would be "Use ELIXIR and Node forums to understand pharmaceutical and biotechnological demands and current limitations impacting this sector." as the first sentence felt generic. - (minor) The verb is missing in " Design targeted training for different microbiome communities" - (minor) Use "findability" instead of "discoverability" for consistency and to fit with the FAIR principles. - (minor) Reorder the sentence to start with the action verb: " E stablish new standards for microbiome research, particularly with respect to data analysis reporting and contextual metadata reporting in conjunction with GSC " - (minor) Correct "Established" to "Establish" in the "Promoting new approaches" entry. Editorial comments: - (major) "Community" is used in upper-case and this is unclear in many instances whether the ELIXIR Marine Metagenomics Community is referred to, the ELIXIR Microbiome Community, or the broader microbiome research community. I suggest to use consistently the full term for the sake of transparency. Abbreviations like EMMC and EMC could be even more misleading in my opinion. - (major) The structure of the white paper is not evident as the hierarchy is indicated only by change in font size, and some sections are quite lengthy for a white paper that is supposed to be concise. I understand that this is a constraint from the Editor, but see Review reference 4 for a white paper with a more clearer structure. An alternative could be to use numbered sections. - References - (minor) The very first reference is oddly formatted in the text creating an artificial and confusing end of sentence. - (minor) Title of reference 9 is truncated and should be "Methods included: standardizing computational reuse and portability with the Common Workflow Language" - (minor) The superscript numbers of the numeric style of bibliography are in many instances after the final dot (see reference 2, 12-16, 17-19, 36, 37), after a comma (see reference 10, 20, 23), a bracket (see reference 60) or semi colon (see reference 65) when they should be before any of these symbols. - (minor) In the paragraph "Interactions with other ELIXIR Communities", we jump from reference 24 to 29 when the numeric style of bibliography (that was chosen by the authors) is expected to mirror the mentions in the manuscript. Please either have the Table 2 earlier in the paper, or change the order of the references. Is the topic of the opinion article discussed accurately in the context of the current literature? Yes Are all factual statements correct and adequately supported by citations? Partly Are arguments sufficiently supported by evidence from the published literature? Partly Are the conclusions drawn balanced and justified on the basis of the presented arguments? Partly References 1. Berg G, Rybakova D, Fischer D, Cernava T, et al.: Microbiome definition re-visited: old concepts and new challenges. Microbiome . 2020; 8 (1). Publisher Full Text 2. Ritsch M, Cassman NA, Saghaei S, Marz M: Navigating the Landscape: A Comprehensive Review of Current Virus Databases. Viruses . 2023; 15 (9). PubMed Abstract | Publisher Full Text 3. Brooks TG, Lahens NF, Mrčela A, Grant GR: Challenges and best practices in omics benchmarking. Nat Rev Genet . 2024. PubMed Abstract | Publisher Full Text 4. Vizcaíno JA, Walzer M, Jiménez RC, Bittremieux W, et al.: A community proposal to integrate proteomics activities in ELIXIR. F1000Res . 2017; 6 . PubMed Abstract | Publisher Full Text Competing Interests: No competing interests were disclosed. Reviewer Expertise: microbiome, bioinformatics, reproducible research, data and metadata standards I confirm that I have read this submission and believe that I have an appropriate level of expertise to confirm that it is of an acceptable scientific standard, however I have significant reservations, as outlined above. Close READ LESS CITE CITE HOW TO CITE THIS REPORT Pauvert C. Reviewer Report For: Establishing the ELIXIR Microbiome Community [version 1; peer review: 1 approved, 1 approved with reservations] . F1000Research 2024, 13 (ELIXIR):50 ( https://doi.org/10.5256/f1000research.158321.r244855 ) The direct URL for this report is: https://f1000research.com/articles/13-50/v1#referee-response-244855 NOTE: it is important to ensure the information in square brackets after the title is included in all citations of this article. COPY CITATION DETAILS Report a concern Author Response 10 Sep 2025 Bérénice Batut , Bioinformatics Group, Department of Computer Science, Albert-Ludwigs-University Freiburg, Freiburg, Germany 10 Sep 2025 Author Response The authors make the case for mutating the ELIXIR Marine Metagenomics Community initiative into an ELIXIR Microbiome Community. It is indeed timely to have an such a pan-European initiative especially ... Continue reading The authors make the case for mutating the ELIXIR Marine Metagenomics Community initiative into an ELIXIR Microbiome Community. It is indeed timely to have an such a pan-European initiative especially as it draws on previous experience, albeit a more narrow biome. The emphasis on multi-omics integration is also a strong suit as it is a complex topic where the microbiome research community would need support and infrastructure. The authors map out existing (but not all) similar initiatives, even if it is unclear how they would work together. In my opinion, a couple of claims should be strengthened and the argumentative power of this white paper would benefit from minor reorganization, a text slim down and extra proofreading. To this end, I highlight the strengths and weaknesses of the manuscript below. ## Strengths - Table 3 is a great idea to map out the others initiatives but needs a bit rework to be straight to the point. - In the Data section, the part between "Throughout the lifetime of the ELIXIR Marine Metagenomics Community" and the end of the paragraph is really relevant and draws from the experience of the ELIXIR Community. These sentences should be more highlighted, maybe upstream. I do wonder how the authors considered how the others initiatives (mentioned in Table 3) could also contribute to the addition of metadata for a global curation effort, possibly via the mentioned Contextual Data Clearinghouse. - The data re-analysis and integration between PRIDE and MGnify (named MetaPUF) is a good use case of the efforts that the new ELIXIR Community can promote. - I did appreciate the clear objective at the end of the Compute section, to help with scaling up analyses as well as the plan for the interaction with the ELIXIR Training Platform which promise to give a boost in training the current and next generation of microbiome researchers. - I liked how the authors highlighted (using the example of viral sequences databases) that time and resources should not be wasted in duplicating works, and that initiatives such as the ELIXIR Microbiome Community can promote this. - The authors thought ahead to promote and develop methods to limit confounding factors when re-using datasets, especially in multi-omics settings, and added this to one of the ELIXIR Microbiome Community challenges. ## Weaknesses ### Introduction - (minor) The microbiome definition used in the paper was revisited in Review reference 1 to include the interactions as well as others molecules than genomes as part of the microbiome. I suggest to use this definition given the emphasis on additional omics made in the paper. We have added that the definition includes other molecules in addition to genomes. - (major) The first paragraph ends on challenges of our research community including how to make data FAIR. Given the experience gained by the ELIXIR Marine Metagenomics Community, I would have appreciated a sentence on why making our data FAIR is important (e.g., transparency to make the research more reproducible, accountability given the (public) source of funding). See major comment in next section. Thank you for this suggestion. We have added a few sentences to this effect in the manuscript. ### The scope of the ELIXIR Microbiome Community - (major) These following arguments for FAIR data should be in the Introduction section. "Moreover, when wishing to contextualise the results with similar experiments, the way a dataset has been produced and processed must be transparent to establish whether it is comparable (e.g. amplified sequence variants can only be compared when the same amplified regions are compared). Furthermore, when different methods are applied, best practices in data stewardship are required to ensure that the connectivity of the derived sequence data products, together with functional and taxonomic assertions are kept in context of the original sample/sequencing effort and associated contextual metadata." We have merged this part of the paragraph into the Introduction. - (minor) The very first sentence of this section blames researchers for misuse of a term "[...] regularly (incorrectly) used [...]". I suggest to rephrase the sentence to simply state the differences, or back up the misuse of the term with references or a survey. We have qualified this with an example of INSDC mislabelling. - (minor) "Fundamentally, the ELIXIR Microbiome Community is about providing the necessary infrastructures required to perform analysis of nucleotide sequence data derived from a microbiome" “Nucleotide” has been added to this sentence. - (minor) "Finally, and possible unique to this ELIXIR Community, is the variety of researchers". I am not sure that these features are unique to the ELIXIR Community but rather to the field of microbiome research. I suggest to tone it down by simply pointing out that the ELIXIR Community gathers a variety of researchers. Modified accordingly. ### Table 1 - (major) Metabolomics is listed in Table 1 but is omitted from the narrative in the paragraph. We have added a sentence in the paragraph. - (major) The Table 1 is not really used but could actually reduce some redundancies in the manuscript by providing a one-stop-shop to explain and detail these techniques. We have increased the cross linking to the table. ### Figure 1 - (major) The figure is early on in the manuscript and it is unclear which of the ELIXIR platforms and communities are already established, or foreseen. Especially since "current and future ELIXIR activities" is mentioned in the manuscript before referencing this figure. The Table 2 does not add more information in that regard. - (major) There is a need for different types of arrows as the same arrow represent interactions between communities or processes. Modified accordingly. - (minor) The "Data" node is quite generic, I was wondering whether it meant public repositories, institute repositories or both. Please be precise. "Data" refers here to the ELIXIR Data Platform - (minor) I guess the type of metadata illustrated in the figure is restricted to biological metadata, that is indeed collected during sampling, however, technical metadata such as the sequencing method used, the type of instrument or the library layout, are collected during the processing of the sample not only the sampling itself. This is everything from phenotypic, sample conditions, experimental methods - (minor) There should be an arrow from "Archive" back to "Data" when the data produced is deposited and then contributes back to public repositories? ### Interactions with other ELIXIR Communities * (major) The sentence "Similarly, many of the biodiversity approaches use marker gene amplification for studying environmental DNA (eDNA)" is redundant with the previous one and introduces a different nomenclature that was not used before (marker gene amplication vs metabarcoding) and the differences, if any, are not explained. I would suggest to remove the sentence. We have modified the sentence to refer to barcoding rather than removing the sentence. We feel it is important to mention eDNA. * (major) The term "Isolation of genomes" is misleading and I guess the authors used it as a shorthand for "The isolation of bacteria, its DNA extraction, genome sequencing and their annotation". Please rephrase to avoid misinterpretation. Plus I would argue that these steps are also done when deconstructing microbiomes via cultivation strategies and are therefore not so out-of-scope. This sentence has been removed as the reviewer is correct that cultivation of bacteria (and other organisms) from microbiomes is becoming an increasing trend. * (minor) "yet each one of these areas is far greater in scientific scope" feels exaggerated and vague. It should be rephrased. A suggestion is: "yet each one of these areas is too complex to be tacked individually" We have amended the sentence accordingly. * (minor) There is no link nor transition between the paragraph that starts with "In summary, " before the Table 3 and the paragraph after that starts with "Similarly, microbiome research". Please rephrase or edit to connect the two sections. Plus, whilst this is good to have a concrete example in the "Similarly" paragraph, the paragraph before was very broad and doing a summary. I would suggest to try reordering the two paragraphs and bring the example earlier for a smoother transition. The paragraphs have been reordered as suggested and slightly amended to improve readability. * (minor) Typo "Similarly, microbiome research has many translation al aspects" Fixed * (minor) The PET acronym is detailed but actually used only once, I am unsure if this is necessary. There is another instance of this abbreviation, so we have retained this abbreviation. ### Table 2 - (major) The text in the "Interaction" column needs rework as it does not use a consistent wording and could be more to the point, especially as it is a complement to the main text: - The Food and Nutrition entry is a question. - The Galaxy entry has an unnecessary return carriage, and a unspecific "ongoing evaluation study" that needs to be clarified. - The Plant Science entry starts with a generic sentence that could be in the introduction or removed for clarity. Plus I would add a clarification that "plants maintain or not their microbial communities across generations." - (minor) The status (e.g., currently active, inactive, planned, etc.) of each ELIXIR Communities would have been appreciated as it is missing also from the Figure 1. - (minor) The mention of "the field" for the Federated human data community is vague in a manuscript about gathering communities, which research field is implied? If this is the human microbiome research field as a whole, please indicate. In response to all of the above comments, we have substantially reworked table 2. The only comment that we have not addressed so specifically, is whether an activity is in progress or not, because the activities ebb and flow, and it is not always easy to say when there is a start or an end. However, being somewhat more focused in the content, we feel this revised table provides a stronger view of the direction that the community wants to take with the other communities. ### Table 3 - (major) Similarly to the Table 2 major comment, there is a lack of consistent wording in the "Aim" column that makes the Table 3 not as impactful as it should and could be. Maybe the authors could extract common/distinct features from each of these initiatives as an alternative way to the "Aim" column. A couple of suggestions for these features would be: is it a national initiative?, does it relates to data storage, data analysis, training? is it linked to ELIXIR? - (major) I was surprised not to see the NMDC listed in the table 3, especially when it is discussed in the main text. What about the NCCR Microbiomes initiative in Switzerland? I can understand that some initiatives are not included for space reason, but maybe state it in the legend of the table. - (minor) Is this table sorted? It seems not, but it could be by acronyms or names. - (minor) The aim for the NFDI4Microbiota is way too big a paragraph. The authors should reduce it for conciseness. - (minor) The Metaproteomics Initiative entry has a hyperlink and a reference when none of the others have. Please homogenise. - (minor) Some entries have country listed and some not. Please homogenise. - (minor) I am not questioning the existence of the European Reference Genome Atlas here, but how best to phrase its relevance to ELIXIR in the manuscript. There seems to have no prokaryotes genomes in their atlas, however, there seems to be a trove of fungi and protists genomes which are usually said to be understudied in microbiome. So I think there is a missed opportunity for the authors here to make the most out of this entry. - (minor) "With its headquarters in Bari (Apulia region)," seems irrelevant in the context of the table, please remove. As with Table 2, we have substantially reworked Table 3, including ordering by the reach of the effort, harmonising the language and reducing text. We have also included the relevance of the activity to the community. We have not included NMDC, as this is a US effort, and while important to the community, this table is focused on Europe efforts. ### Data - (minor) The end of the following sentence is redundant as the INSDC was introduced earlier already "INSDC, which in collaboration with the National Institute of Genetics DNA DataBank of Japan (DDBJ) and the United States National Center for Biotechnology’s (NCBI) GenBank and Sequence Read Archive (SRA), facilitate the deposition and global exchange of sequence data.". Please adjust accordingly. We have removed the redundancy here. - (minor) "A current challenge facing the field is connecting different multi ‘omics data that have been derived from the same sample." Is this going to be tackled by the ELIXIR Microbiome Community? If so, I would state that this is part of its objectives. We feel this is likely to be a joint effort across different communities. Thus, we have retained the sentence as is. - (minor) The end of sentence "[...]overarching context to the experiment, which can be important for meta-analyses." seems like an euphemism, I would suggest to replace with "[...]overarching context to the experiment, re-analyses or meta-analyses." to include re-analyses as well. We have added “re-analyses or” as suggested. - (minor) "We will continue to promote such approaches, enriching metadata wherever possible." Is this going to be done via the CDCH? While CDCH offers one approach, we believe that there are many different ways that metadata may be enriched, first by making scientists more aware of the need of submitting meta, promoting different, more specific checklists (e.g. STORMS or MicroB3) and mining metadata from published literature. This, specifically highlighting CDCH would not be appropriate. - (minor) "The ELIXIR Microbiome Community will also work to move the Marine Metagenomics domain in the RDMKit towards a more general Microbiome domain." What is the RDMKit? It is not explained, nor cited nor mentioned again. This has been expanded to the ELIXIR Research Data Management Kit and a link has been added to the text. ### Tools - (minor) "will increase their use of BioContainers" should be "will increase the use of BioContainers" Corrected. - (minor) "In order to make tools findable by the end users, the Community" Done. - (minor) "workflow descriptions (e.g. Snakemake, CWL, Nextflow)" None of them have their references cited, is it an omission or space limitation? Added references - (minor) "A current joint effort between the Microbiome and Galaxy Communities" Fixed ### Benchmarking - (major) "Benchmarking" it is the only item at this hierarchy level, meaning that this is useless for structuring the text. Please edit. Thank you for pointing this out. We have changed the sub-sub-heading into an introductory sentence to this paragraph. - (minor) Review reference 2 published a recent review with guidelines to learn from that could have its place in this paragraph. - (minor) in the sentence: "As the Microbiome Community establishes, we will develop a broader understanding of the requirements of the Community , feed this to the Tools Platform, as well as seek opportunities to interact with the Tools Platform to capture the diversity of tools and their utility via such benchmarking activities.", is this the ELIXIR Microbiome Community, or the wide community of microbiome researchers? Is it to mean that the ELIXIR Microbiome Community is going to act as an interface between microbiome researchers and ELIXIR Tools/Infrastructure? We clarified this to mean the wider community of microbiome researchers. ### Compute - (minor) The first sentence would fit better in the introduction. ### Interoperability - (minor) The reference 57 should be at the end of the sentence starting with "This effort was paralleled" not in the middle. This citation has been moved to the end of the sentence. - (minor) Reference 58 should be removed at it is a duplicate of reference 47. References have been fixed - (minor) Use the full text "Global Alliance for Genomics and Health" instead of GA4GH. This abbreviation has been expanded. - (minor) "is in the process of applying to be a n ELIXIR Recommended Interoperability Resource." ELIXIR has been added to this sentence. - (minor) In the sentence: "This will require the development of new data Interoperability layers for data resources that are not normally focused in Microbiome data" I think the authors meant "used" instead of "focused", and "microbiome" instead of "Microbiome". We have updated the sentence accordingly. ### Training - (minor) In "Platforms such as MGnify support large-scale services for most, if not all, steps of a microbiome study", I would suggest to remove the "if not all". Agreed. - (minor) "with areas of expertise covering different environments, ‘omics approaches and data analysis pathways." I think the authors meant "learning paths", and I would refrain from using "pathways" as it also has a biological meaning. Actually, we do mean data analysis, but have changed to data analysis strategies. This is in reference to the fact that there may be multiple different ways of analysing a data type, for example metagenomics can be analysed using tools such as Kraken or MetaPhlAn4, to full assembly, gene calling and functional analysis. ### Context with other international initiatives - (major) Whilst I appreciated the emphasis that no initiative exists on its own, I feel the first paragraph on the Genom ic (please correct the typo) Standards Consortium feels lengthy for a manuscript whose topic is not the GSC. I would advise to summarize. In this respect, the second paragraph is particularly relevant to a tangible collaboration between ELIXIR Microbiome Community and GSC. The typo has been corrected. - (major) The NMDC is discussed in this section but not part of the Table 3. Table 3 lists pan-European efforts, whereas the NMDC is a US specific initiative. Thus, we have not added this to the table. - (minor) Is the mentioned M5 project still active as the website's last update is 2012? While the website has not been updated, there are ongoing efforts to expand and revitalise this effort. RDF is a member of the GSC board and is promoting aspects of the M5 initiative. While the name may change, we feel this is useful to keep. However, as part of condensing the overall length of the manuscript, we have removed the reference to this initiative. - (minor) "Combining the activities on standards concerning workflows [...]" does this means adding and providing workflows to the microbiome research community? - (minor) Given the emphasis on the fact that MicrobiomeSupport was a program, the authors could update the readers and indicate that it is now MicrobiomeSupport Association. ### Interaction with other key data resources beyond ELIXIR - (major) There is an order issue with the main text that a proofread could solve, as MG-RAST is explained and cited in the first paragraph but already mentioned upstream of the main text in the Interoperability section. This has now been fixed. ### Specific challenges and objectives of the ELIXIR Microbiome Community - (major) The strong claim "it is widely accepted that current short-read assembly-based methods do not generally work as well for soil microbiomes" would probably need at least one reference. Added - (major) "(iii) there is no centralised database collecting the millions of viral sequences". It seems to be the case indeed, and there are databases (~24) out there as recently compiled in Review reference 3. How ELIXIR Microbiome Community plans to integrate/aggregate these resources in a non-duplicating manner? - (minor) The word "through" is superfluous in the the sentence that starts with "This current limitation, [...]" and can be removed. We have removed the sentence “through”. - (minor) The CAMI was already explained and cited above, so the already defined acronym can be used. This has been fixed. - (minor) Precise the area in :"Additionally, another key area of development of taxonomy [...]". Added - (minor) The sentence "Viruses, particularly those that infect bacteria, are found ubiquitously in all environments and play critical roles in community dynamics." belongs in an introduction, not so downstream of the manuscript. - (minor) The sentence starting the sixth paragraph could be precised as "The increase in metagenomic assemblies has resulted in a parallel increase in the number of predicted protein sequences, with sets of non-redundant proteins now in the billions." Added “predicted” to the sentence. - (minor) Fix typo in "that are undetectable by current sequence based methods." Fixed - sequenced -> sequence - (minor) Would it make sense to also be able to access representative, of clusters for instance? "[...] develop new infrastructural frameworks for accessing slices of the data or adequate representatives based on the requirements." Added - (minor) What is the " expanded Microbiome Community"? It was never mentioned before. Removed “expanded” from the sentence. - (minor) A few comments on the sentence: "Within the Community, we will develop and promote standards around the analysis provenance (analytical metadata),". In my opinion and how it was already stated in the manuscript, it would make more sense to promote existing standards first and then develop if need be. There is no mention of other type of metadata, so the "analytical metadata" precision seems superfluous. I would suggest: "Within the Community, we will promote and develop standards regarding the analysis provenance," We agree with the reviewer's comment and have re-ordered accordingly. - (minor) Precise the term forms in "[...] ensuring that comput ing resources are accessible for performing the different forms of data analysis [...]", do the authors mean types of/steps in the data analysis? “compute resources” is an accepted resource. We have removed the “different forms of” as we feel it is a little superfluous. We were meaning the different strategies, but it does not add to the sentence. - (minor) Same argument as before regarding reinventing the wheel, I would swap the part of the sentence: "This may require the extensions to existing databases or development of new ones , but it requires an agreement from the research community to adopt them." We agree, and have amended accordingly. - (minor) This part "Metaproteomics aims to elucidate the functional and taxonomic interplay of proteins in microbiomes," should have been in the Table 1, or to reuse the Table 1 here. We have highlighted this in table 1. We have kept the text here to provide the context for the rest of the sentence. - (minor) The tenth paragraph of this section starts with the mention of multiple major challenges, but detail "only" one of them. I would suggest to mention some of the others challenges. - (minor) The "MIA" method is just a hyperlink, without any reference. Either cite the website accordingly or add the reference. There is not a reference, and the hyperlink goes to the GitHub site, as requested by the authors. We have included “ https://github.com/microbiome/mia ” for completeness sake. ### Table 4 - (major) I was surprised to see that the objective "Foster international collaborations between other resources providers and databases to ensure global harmonisation of e-infrastructures for microbiome research" was long-term, as I would have imagined that a gap analysis would be short-term to ensure we do not reinvent the wheel, especially given the others initiatives discussed in the manuscript. - (minor) If the Objective column starts with action verbs (which is a good idea), then it should be "Survey the needs" instead of "Survey of needs". We have amended accordingly. - (minor) Specify the type of workflow with "Address knowledge gaps in generating and adopting data analysis workflows" Added - (minor) The verb is missing in " Teach a dvanced containerisation and cloud deployment" Added. - (minor) The verb is missing in " Promote data analysis through the use of services". Is this ELIXIR services in general or specific ELIXIR Microbiome Community services? Added Promote - generalised to ELIXIR services. - (minor) Use "Share" instead of "Sharing" in the "Co-ordinate" entry. Done - (minor) The "Industry connection" entry does not fit the action verb pattern. A suggestion would be "Use ELIXIR and Node forums to understand pharmaceutical and biotechnological demands and current limitations impacting this sector." as the first sentence felt generic. Done - (minor) The verb is missing in " Design targeted training for different microbiome communities" Done - (minor) Use "findability" instead of "discoverability" for consistency and to fit with the FAIR principles. Done . - (minor) Reorder the sentence to start with the action verb: " E stablish new standards for microbiome research, particularly with respect to data analysis reporting and contextual metadata reporting in conjunction with GSC " Done - (minor) Correct "Established" to "Establish" in the "Promoting new approaches" entry. Done Editorial comments: - (major) "Community" is used in upper-case and this is unclear in many instances whether the ELIXIR Marine Metagenomics Community is referred to, the ELIXIR Microbiome Community, or the broader microbiome research community. I suggest to use consistently the full term for the sake of transparency. Abbreviations like EMMC and EMC could be even more misleading in my opinion. We have prefixed all instances of “Community” with their explicit community name. - (major) The structure of the white paper is not evident as the hierarchy is indicated only by change in font size, and some sections are quite lengthy for a white paper that is supposed to be concise. I understand that this is a constraint from the Editor, but see Review reference 4 for a white paper with a more clearer structure. An alternative could be to use numbered sections. We have introduced section numbers to aid the structuring of the manuscript. - References - (minor) The very first reference is oddly formatted in the text creating an artificial and confusing end of sentence. - (minor) Title of reference 9 is truncated and should be "Methods included: standardizing computational reuse and portability with the Common Workflow Language" - (minor) The superscript numbers of the numeric style of bibliography are in many instances after the final dot (see reference 2, 12-16, 17-19, 36, 37), after a comma (see reference 10, 20, 23), a bracket (see reference 60) or semi colon (see reference 65) when they should be before any of these symbols. - (minor) In the paragraph "Interactions with other ELIXIR Communities", we jump from reference 24 to 29 when the numeric style of bibliography (that was chosen by the authors) is expected to mirror the mentions in the manuscript. Please either have the Table 2 earlier in the paper, or change the order of the references. References have been fixed The authors make the case for mutating the ELIXIR Marine Metagenomics Community initiative into an ELIXIR Microbiome Community. It is indeed timely to have an such a pan-European initiative especially as it draws on previous experience, albeit a more narrow biome. The emphasis on multi-omics integration is also a strong suit as it is a complex topic where the microbiome research community would need support and infrastructure. The authors map out existing (but not all) similar initiatives, even if it is unclear how they would work together. In my opinion, a couple of claims should be strengthened and the argumentative power of this white paper would benefit from minor reorganization, a text slim down and extra proofreading. To this end, I highlight the strengths and weaknesses of the manuscript below. ## Strengths - Table 3 is a great idea to map out the others initiatives but needs a bit rework to be straight to the point. - In the Data section, the part between "Throughout the lifetime of the ELIXIR Marine Metagenomics Community" and the end of the paragraph is really relevant and draws from the experience of the ELIXIR Community. These sentences should be more highlighted, maybe upstream. I do wonder how the authors considered how the others initiatives (mentioned in Table 3) could also contribute to the addition of metadata for a global curation effort, possibly via the mentioned Contextual Data Clearinghouse. - The data re-analysis and integration between PRIDE and MGnify (named MetaPUF) is a good use case of the efforts that the new ELIXIR Community can promote. - I did appreciate the clear objective at the end of the Compute section, to help with scaling up analyses as well as the plan for the interaction with the ELIXIR Training Platform which promise to give a boost in training the current and next generation of microbiome researchers. - I liked how the authors highlighted (using the example of viral sequences databases) that time and resources should not be wasted in duplicating works, and that initiatives such as the ELIXIR Microbiome Community can promote this. - The authors thought ahead to promote and develop methods to limit confounding factors when re-using datasets, especially in multi-omics settings, and added this to one of the ELIXIR Microbiome Community challenges. ## Weaknesses ### Introduction - (minor) The microbiome definition used in the paper was revisited in Review reference 1 to include the interactions as well as others molecules than genomes as part of the microbiome. I suggest to use this definition given the emphasis on additional omics made in the paper. We have added that the definition includes other molecules in addition to genomes. - (major) The first paragraph ends on challenges of our research community including how to make data FAIR. Given the experience gained by the ELIXIR Marine Metagenomics Community, I would have appreciated a sentence on why making our data FAIR is important (e.g., transparency to make the research more reproducible, accountability given the (public) source of funding). See major comment in next section. Thank you for this suggestion. We have added a few sentences to this effect in the manuscript. ### The scope of the ELIXIR Microbiome Community - (major) These following arguments for FAIR data should be in the Introduction section. "Moreover, when wishing to contextualise the results with similar experiments, the way a dataset has been produced and processed must be transparent to establish whether it is comparable (e.g. amplified sequence variants can only be compared when the same amplified regions are compared). Furthermore, when different methods are applied, best practices in data stewardship are required to ensure that the connectivity of the derived sequence data products, together with functional and taxonomic assertions are kept in context of the original sample/sequencing effort and associated contextual metadata." We have merged this part of the paragraph into the Introduction. - (minor) The very first sentence of this section blames researchers for misuse of a term "[...] regularly (incorrectly) used [...]". I suggest to rephrase the sentence to simply state the differences, or back up the misuse of the term with references or a survey. We have qualified this with an example of INSDC mislabelling. - (minor) "Fundamentally, the ELIXIR Microbiome Community is about providing the necessary infrastructures required to perform analysis of nucleotide sequence data derived from a microbiome" “Nucleotide” has been added to this sentence. - (minor) "Finally, and possible unique to this ELIXIR Community, is the variety of researchers". I am not sure that these features are unique to the ELIXIR Community but rather to the field of microbiome research. I suggest to tone it down by simply pointing out that the ELIXIR Community gathers a variety of researchers. Modified accordingly. ### Table 1 - (major) Metabolomics is listed in Table 1 but is omitted from the narrative in the paragraph. We have added a sentence in the paragraph. - (major) The Table 1 is not really used but could actually reduce some redundancies in the manuscript by providing a one-stop-shop to explain and detail these techniques. We have increased the cross linking to the table. ### Figure 1 - (major) The figure is early on in the manuscript and it is unclear which of the ELIXIR platforms and communities are already established, or foreseen. Especially since "current and future ELIXIR activities" is mentioned in the manuscript before referencing this figure. The Table 2 does not add more information in that regard. - (major) There is a need for different types of arrows as the same arrow represent interactions between communities or processes. Modified accordingly. - (minor) The "Data" node is quite generic, I was wondering whether it meant public repositories, institute repositories or both. Please be precise. "Data" refers here to the ELIXIR Data Platform - (minor) I guess the type of metadata illustrated in the figure is restricted to biological metadata, that is indeed collected during sampling, however, technical metadata such as the sequencing method used, the type of instrument or the library layout, are collected during the processing of the sample not only the sampling itself. This is everything from phenotypic, sample conditions, experimental methods - (minor) There should be an arrow from "Archive" back to "Data" when the data produced is deposited and then contributes back to public repositories? ### Interactions with other ELIXIR Communities * (major) The sentence "Similarly, many of the biodiversity approaches use marker gene amplification for studying environmental DNA (eDNA)" is redundant with the previous one and introduces a different nomenclature that was not used before (marker gene amplication vs metabarcoding) and the differences, if any, are not explained. I would suggest to remove the sentence. We have modified the sentence to refer to barcoding rather than removing the sentence. We feel it is important to mention eDNA. * (major) The term "Isolation of genomes" is misleading and I guess the authors used it as a shorthand for "The isolation of bacteria, its DNA extraction, genome sequencing and their annotation". Please rephrase to avoid misinterpretation. Plus I would argue that these steps are also done when deconstructing microbiomes via cultivation strategies and are therefore not so out-of-scope. This sentence has been removed as the reviewer is correct that cultivation of bacteria (and other organisms) from microbiomes is becoming an increasing trend. * (minor) "yet each one of these areas is far greater in scientific scope" feels exaggerated and vague. It should be rephrased. A suggestion is: "yet each one of these areas is too complex to be tacked individually" We have amended the sentence accordingly. * (minor) There is no link nor transition between the paragraph that starts with "In summary, " before the Table 3 and the paragraph after that starts with "Similarly, microbiome research". Please rephrase or edit to connect the two sections. Plus, whilst this is good to have a concrete example in the "Similarly" paragraph, the paragraph before was very broad and doing a summary. I would suggest to try reordering the two paragraphs and bring the example earlier for a smoother transition. The paragraphs have been reordered as suggested and slightly amended to improve readability. * (minor) Typo "Similarly, microbiome research has many translation al aspects" Fixed * (minor) The PET acronym is detailed but actually used only once, I am unsure if this is necessary. There is another instance of this abbreviation, so we have retained this abbreviation. ### Table 2 - (major) The text in the "Interaction" column needs rework as it does not use a consistent wording and could be more to the point, especially as it is a complement to the main text: - The Food and Nutrition entry is a question. - The Galaxy entry has an unnecessary return carriage, and a unspecific "ongoing evaluation study" that needs to be clarified. - The Plant Science entry starts with a generic sentence that could be in the introduction or removed for clarity. Plus I would add a clarification that "plants maintain or not their microbial communities across generations." - (minor) The status (e.g., currently active, inactive, planned, etc.) of each ELIXIR Communities would have been appreciated as it is missing also from the Figure 1. - (minor) The mention of "the field" for the Federated human data community is vague in a manuscript about gathering communities, which research field is implied? If this is the human microbiome research field as a whole, please indicate. In response to all of the above comments, we have substantially reworked table 2. The only comment that we have not addressed so specifically, is whether an activity is in progress or not, because the activities ebb and flow, and it is not always easy to say when there is a start or an end. However, being somewhat more focused in the content, we feel this revised table provides a stronger view of the direction that the community wants to take with the other communities. ### Table 3 - (major) Similarly to the Table 2 major comment, there is a lack of consistent wording in the "Aim" column that makes the Table 3 not as impactful as it should and could be. Maybe the authors could extract common/distinct features from each of these initiatives as an alternative way to the "Aim" column. A couple of suggestions for these features would be: is it a national initiative?, does it relates to data storage, data analysis, training? is it linked to ELIXIR? - (major) I was surprised not to see the NMDC listed in the table 3, especially when it is discussed in the main text. What about the NCCR Microbiomes initiative in Switzerland? I can understand that some initiatives are not included for space reason, but maybe state it in the legend of the table. - (minor) Is this table sorted? It seems not, but it could be by acronyms or names. - (minor) The aim for the NFDI4Microbiota is way too big a paragraph. The authors should reduce it for conciseness. - (minor) The Metaproteomics Initiative entry has a hyperlink and a reference when none of the others have. Please homogenise. - (minor) Some entries have country listed and some not. Please homogenise. - (minor) I am not questioning the existence of the European Reference Genome Atlas here, but how best to phrase its relevance to ELIXIR in the manuscript. There seems to have no prokaryotes genomes in their atlas, however, there seems to be a trove of fungi and protists genomes which are usually said to be understudied in microbiome. So I think there is a missed opportunity for the authors here to make the most out of this entry. - (minor) "With its headquarters in Bari (Apulia region)," seems irrelevant in the context of the table, please remove. As with Table 2, we have substantially reworked Table 3, including ordering by the reach of the effort, harmonising the language and reducing text. We have also included the relevance of the activity to the community. We have not included NMDC, as this is a US effort, and while important to the community, this table is focused on Europe efforts. ### Data - (minor) The end of the following sentence is redundant as the INSDC was introduced earlier already "INSDC, which in collaboration with the National Institute of Genetics DNA DataBank of Japan (DDBJ) and the United States National Center for Biotechnology’s (NCBI) GenBank and Sequence Read Archive (SRA), facilitate the deposition and global exchange of sequence data.". Please adjust accordingly. We have removed the redundancy here. - (minor) "A current challenge facing the field is connecting different multi ‘omics data that have been derived from the same sample." Is this going to be tackled by the ELIXIR Microbiome Community? If so, I would state that this is part of its objectives. We feel this is likely to be a joint effort across different communities. Thus, we have retained the sentence as is. - (minor) The end of sentence "[...]overarching context to the experiment, which can be important for meta-analyses." seems like an euphemism, I would suggest to replace with "[...]overarching context to the experiment, re-analyses or meta-analyses." to include re-analyses as well. We have added “re-analyses or” as suggested. - (minor) "We will continue to promote such approaches, enriching metadata wherever possible." Is this going to be done via the CDCH? While CDCH offers one approach, we believe that there are many different ways that metadata may be enriched, first by making scientists more aware of the need of submitting meta, promoting different, more specific checklists (e.g. STORMS or MicroB3) and mining metadata from published literature. This, specifically highlighting CDCH would not be appropriate. - (minor) "The ELIXIR Microbiome Community will also work to move the Marine Metagenomics domain in the RDMKit towards a more general Microbiome domain." What is the RDMKit? It is not explained, nor cited nor mentioned again. This has been expanded to the ELIXIR Research Data Management Kit and a link has been added to the text. ### Tools - (minor) "will increase their use of BioContainers" should be "will increase the use of BioContainers" Corrected. - (minor) "In order to make tools findable by the end users, the Community" Done. - (minor) "workflow descriptions (e.g. Snakemake, CWL, Nextflow)" None of them have their references cited, is it an omission or space limitation? Added references - (minor) "A current joint effort between the Microbiome and Galaxy Communities" Fixed ### Benchmarking - (major) "Benchmarking" it is the only item at this hierarchy level, meaning that this is useless for structuring the text. Please edit. Thank you for pointing this out. We have changed the sub-sub-heading into an introductory sentence to this paragraph. - (minor) Review reference 2 published a recent review with guidelines to learn from that could have its place in this paragraph. - (minor) in the sentence: "As the Microbiome Community establishes, we will develop a broader understanding of the requirements of the Community , feed this to the Tools Platform, as well as seek opportunities to interact with the Tools Platform to capture the diversity of tools and their utility via such benchmarking activities.", is this the ELIXIR Microbiome Community, or the wide community of microbiome researchers? Is it to mean that the ELIXIR Microbiome Community is going to act as an interface between microbiome researchers and ELIXIR Tools/Infrastructure? We clarified this to mean the wider community of microbiome researchers. ### Compute - (minor) The first sentence would fit better in the introduction. ### Interoperability - (minor) The reference 57 should be at the end of the sentence starting with "This effort was paralleled" not in the middle. This citation has been moved to the end of the sentence. - (minor) Reference 58 should be removed at it is a duplicate of reference 47. References have been fixed - (minor) Use the full text "Global Alliance for Genomics and Health" instead of GA4GH. This abbreviation has been expanded. - (minor) "is in the process of applying to be a n ELIXIR Recommended Interoperability Resource." ELIXIR has been added to this sentence. - (minor) In the sentence: "This will require the development of new data Interoperability layers for data resources that are not normally focused in Microbiome data" I think the authors meant "used" instead of "focused", and "microbiome" instead of "Microbiome". We have updated the sentence accordingly. ### Training - (minor) In "Platforms such as MGnify support large-scale services for most, if not all, steps of a microbiome study", I would suggest to remove the "if not all". Agreed. - (minor) "with areas of expertise covering different environments, ‘omics approaches and data analysis pathways." I think the authors meant "learning paths", and I would refrain from using "pathways" as it also has a biological meaning. Actually, we do mean data analysis, but have changed to data analysis strategies. This is in reference to the fact that there may be multiple different ways of analysing a data type, for example metagenomics can be analysed using tools such as Kraken or MetaPhlAn4, to full assembly, gene calling and functional analysis. ### Context with other international initiatives - (major) Whilst I appreciated the emphasis that no initiative exists on its own, I feel the first paragraph on the Genom ic (please correct the typo) Standards Consortium feels lengthy for a manuscript whose topic is not the GSC. I would advise to summarize. In this respect, the second paragraph is particularly relevant to a tangible collaboration between ELIXIR Microbiome Community and GSC. The typo has been corrected. - (major) The NMDC is discussed in this section but not part of the Table 3. Table 3 lists pan-European efforts, whereas the NMDC is a US specific initiative. Thus, we have not added this to the table. - (minor) Is the mentioned M5 project still active as the website's last update is 2012? While the website has not been updated, there are ongoing efforts to expand and revitalise this effort. RDF is a member of the GSC board and is promoting aspects of the M5 initiative. While the name may change, we feel this is useful to keep. However, as part of condensing the overall length of the manuscript, we have removed the reference to this initiative. - (minor) "Combining the activities on standards concerning workflows [...]" does this means adding and providing workflows to the microbiome research community? - (minor) Given the emphasis on the fact that MicrobiomeSupport was a program, the authors could update the readers and indicate that it is now MicrobiomeSupport Association. ### Interaction with other key data resources beyond ELIXIR - (major) There is an order issue with the main text that a proofread could solve, as MG-RAST is explained and cited in the first paragraph but already mentioned upstream of the main text in the Interoperability section. This has now been fixed. ### Specific challenges and objectives of the ELIXIR Microbiome Community - (major) The strong claim "it is widely accepted that current short-read assembly-based methods do not generally work as well for soil microbiomes" would probably need at least one reference. Added - (major) "(iii) there is no centralised database collecting the millions of viral sequences". It seems to be the case indeed, and there are databases (~24) out there as recently compiled in Review reference 3. How ELIXIR Microbiome Community plans to integrate/aggregate these resources in a non-duplicating manner? - (minor) The word "through" is superfluous in the the sentence that starts with "This current limitation, [...]" and can be removed. We have removed the sentence “through”. - (minor) The CAMI was already explained and cited above, so the already defined acronym can be used. This has been fixed. - (minor) Precise the area in :"Additionally, another key area of development of taxonomy [...]". Added - (minor) The sentence "Viruses, particularly those that infect bacteria, are found ubiquitously in all environments and play critical roles in community dynamics." belongs in an introduction, not so downstream of the manuscript. - (minor) The sentence starting the sixth paragraph could be precised as "The increase in metagenomic assemblies has resulted in a parallel increase in the number of predicted protein sequences, with sets of non-redundant proteins now in the billions." Added “predicted” to the sentence. - (minor) Fix typo in "that are undetectable by current sequence based methods." Fixed - sequenced -> sequence - (minor) Would it make sense to also be able to access representative, of clusters for instance? "[...] develop new infrastructural frameworks for accessing slices of the data or adequate representatives based on the requirements." Added - (minor) What is the " expanded Microbiome Community"? It was never mentioned before. Removed “expanded” from the sentence. - (minor) A few comments on the sentence: "Within the Community, we will develop and promote standards around the analysis provenance (analytical metadata),". In my opinion and how it was already stated in the manuscript, it would make more sense to promote existing standards first and then develop if need be. There is no mention of other type of metadata, so the "analytical metadata" precision seems superfluous. I would suggest: "Within the Community, we will promote and develop standards regarding the analysis provenance," We agree with the reviewer's comment and have re-ordered accordingly. - (minor) Precise the term forms in "[...] ensuring that comput ing resources are accessible for performing the different forms of data analysis [...]", do the authors mean types of/steps in the data analysis? “compute resources” is an accepted resource. We have removed the “different forms of” as we feel it is a little superfluous. We were meaning the different strategies, but it does not add to the sentence. - (minor) Same argument as before regarding reinventing the wheel, I would swap the part of the sentence: "This may require the extensions to existing databases or development of new ones , but it requires an agreement from the research community to adopt them." We agree, and have amended accordingly. - (minor) This part "Metaproteomics aims to elucidate the functional and taxonomic interplay of proteins in microbiomes," should have been in the Table 1, or to reuse the Table 1 here. We have highlighted this in table 1. We have kept the text here to provide the context for the rest of the sentence. - (minor) The tenth paragraph of this section starts with the mention of multiple major challenges, but detail "only" one of them. I would suggest to mention some of the others challenges. - (minor) The "MIA" method is just a hyperlink, without any reference. Either cite the website accordingly or add the reference. There is not a reference, and the hyperlink goes to the GitHub site, as requested by the authors. We have included “ https://github.com/microbiome/mia ” for completeness sake. ### Table 4 - (major) I was surprised to see that the objective "Foster international collaborations between other resources providers and databases to ensure global harmonisation of e-infrastructures for microbiome research" was long-term, as I would have imagined that a gap analysis would be short-term to ensure we do not reinvent the wheel, especially given the others initiatives discussed in the manuscript. - (minor) If the Objective column starts with action verbs (which is a good idea), then it should be "Survey the needs" instead of "Survey of needs". We have amended accordingly. - (minor) Specify the type of workflow with "Address knowledge gaps in generating and adopting data analysis workflows" Added - (minor) The verb is missing in " Teach a dvanced containerisation and cloud deployment" Added. - (minor) The verb is missing in " Promote data analysis through the use of services". Is this ELIXIR services in general or specific ELIXIR Microbiome Community services? Added Promote - generalised to ELIXIR services. - (minor) Use "Share" instead of "Sharing" in the "Co-ordinate" entry. Done - (minor) The "Industry connection" entry does not fit the action verb pattern. A suggestion would be "Use ELIXIR and Node forums to understand pharmaceutical and biotechnological demands and current limitations impacting this sector." as the first sentence felt generic. Done - (minor) The verb is missing in " Design targeted training for different microbiome communities" Done - (minor) Use "findability" instead of "discoverability" for consistency and to fit with the FAIR principles. Done . - (minor) Reorder the sentence to start with the action verb: " E stablish new standards for microbiome research, particularly with respect to data analysis reporting and contextual metadata reporting in conjunction with GSC " Done - (minor) Correct "Established" to "Establish" in the "Promoting new approaches" entry. Done Editorial comments: - (major) "Community" is used in upper-case and this is unclear in many instances whether the ELIXIR Marine Metagenomics Community is referred to, the ELIXIR Microbiome Community, or the broader microbiome research community. I suggest to use consistently the full term for the sake of transparency. Abbreviations like EMMC and EMC could be even more misleading in my opinion. We have prefixed all instances of “Community” with their explicit community name. - (major) The structure of the white paper is not evident as the hierarchy is indicated only by change in font size, and some sections are quite lengthy for a white paper that is supposed to be concise. I understand that this is a constraint from the Editor, but see Review reference 4 for a white paper with a more clearer structure. An alternative could be to use numbered sections. We have introduced section numbers to aid the structuring of the manuscript. - References - (minor) The very first reference is oddly formatted in the text creating an artificial and confusing end of sentence. - (minor) Title of reference 9 is truncated and should be "Methods included: standardizing computational reuse and portability with the Common Workflow Language" - (minor) The superscript numbers of the numeric style of bibliography are in many instances after the final dot (see reference 2, 12-16, 17-19, 36, 37), after a comma (see reference 10, 20, 23), a bracket (see reference 60) or semi colon (see reference 65) when they should be before any of these symbols. - (minor) In the paragraph "Interactions with other ELIXIR Communities", we jump from reference 24 to 29 when the numeric style of bibliography (that was chosen by the authors) is expected to mirror the mentions in the manuscript. Please either have the Table 2 earlier in the paper, or change the order of the references. References have been fixed Competing Interests: No competing interests were disclosed. Close Report a concern Respond or Comment COMMENTS ON THIS REPORT Author Response 10 Sep 2025 Bérénice Batut , Bioinformatics Group, Department of Computer Science, Albert-Ludwigs-University Freiburg, Freiburg, Germany 10 Sep 2025 Author Response The authors make the case for mutating the ELIXIR Marine Metagenomics Community initiative into an ELIXIR Microbiome Community. It is indeed timely to have an such a pan-European initiative especially ... Continue reading The authors make the case for mutating the ELIXIR Marine Metagenomics Community initiative into an ELIXIR Microbiome Community. It is indeed timely to have an such a pan-European initiative especially as it draws on previous experience, albeit a more narrow biome. The emphasis on multi-omics integration is also a strong suit as it is a complex topic where the microbiome research community would need support and infrastructure. The authors map out existing (but not all) similar initiatives, even if it is unclear how they would work together. In my opinion, a couple of claims should be strengthened and the argumentative power of this white paper would benefit from minor reorganization, a text slim down and extra proofreading. To this end, I highlight the strengths and weaknesses of the manuscript below. ## Strengths - Table 3 is a great idea to map out the others initiatives but needs a bit rework to be straight to the point. - In the Data section, the part between "Throughout the lifetime of the ELIXIR Marine Metagenomics Community" and the end of the paragraph is really relevant and draws from the experience of the ELIXIR Community. These sentences should be more highlighted, maybe upstream. I do wonder how the authors considered how the others initiatives (mentioned in Table 3) could also contribute to the addition of metadata for a global curation effort, possibly via the mentioned Contextual Data Clearinghouse. - The data re-analysis and integration between PRIDE and MGnify (named MetaPUF) is a good use case of the efforts that the new ELIXIR Community can promote. - I did appreciate the clear objective at the end of the Compute section, to help with scaling up analyses as well as the plan for the interaction with the ELIXIR Training Platform which promise to give a boost in training the current and next generation of microbiome researchers. - I liked how the authors highlighted (using the example of viral sequences databases) that time and resources should not be wasted in duplicating works, and that initiatives such as the ELIXIR Microbiome Community can promote this. - The authors thought ahead to promote and develop methods to limit confounding factors when re-using datasets, especially in multi-omics settings, and added this to one of the ELIXIR Microbiome Community challenges. ## Weaknesses ### Introduction - (minor) The microbiome definition used in the paper was revisited in Review reference 1 to include the interactions as well as others molecules than genomes as part of the microbiome. I suggest to use this definition given the emphasis on additional omics made in the paper. We have added that the definition includes other molecules in addition to genomes. - (major) The first paragraph ends on challenges of our research community including how to make data FAIR. Given the experience gained by the ELIXIR Marine Metagenomics Community, I would have appreciated a sentence on why making our data FAIR is important (e.g., transparency to make the research more reproducible, accountability given the (public) source of funding). See major comment in next section. Thank you for this suggestion. We have added a few sentences to this effect in the manuscript. ### The scope of the ELIXIR Microbiome Community - (major) These following arguments for FAIR data should be in the Introduction section. "Moreover, when wishing to contextualise the results with similar experiments, the way a dataset has been produced and processed must be transparent to establish whether it is comparable (e.g. amplified sequence variants can only be compared when the same amplified regions are compared). Furthermore, when different methods are applied, best practices in data stewardship are required to ensure that the connectivity of the derived sequence data products, together with functional and taxonomic assertions are kept in context of the original sample/sequencing effort and associated contextual metadata." We have merged this part of the paragraph into the Introduction. - (minor) The very first sentence of this section blames researchers for misuse of a term "[...] regularly (incorrectly) used [...]". I suggest to rephrase the sentence to simply state the differences, or back up the misuse of the term with references or a survey. We have qualified this with an example of INSDC mislabelling. - (minor) "Fundamentally, the ELIXIR Microbiome Community is about providing the necessary infrastructures required to perform analysis of nucleotide sequence data derived from a microbiome" “Nucleotide” has been added to this sentence. - (minor) "Finally, and possible unique to this ELIXIR Community, is the variety of researchers". I am not sure that these features are unique to the ELIXIR Community but rather to the field of microbiome research. I suggest to tone it down by simply pointing out that the ELIXIR Community gathers a variety of researchers. Modified accordingly. ### Table 1 - (major) Metabolomics is listed in Table 1 but is omitted from the narrative in the paragraph. We have added a sentence in the paragraph. - (major) The Table 1 is not really used but could actually reduce some redundancies in the manuscript by providing a one-stop-shop to explain and detail these techniques. We have increased the cross linking to the table. ### Figure 1 - (major) The figure is early on in the manuscript and it is unclear which of the ELIXIR platforms and communities are already established, or foreseen. Especially since "current and future ELIXIR activities" is mentioned in the manuscript before referencing this figure. The Table 2 does not add more information in that regard. - (major) There is a need for different types of arrows as the same arrow represent interactions between communities or processes. Modified accordingly. - (minor) The "Data" node is quite generic, I was wondering whether it meant public repositories, institute repositories or both. Please be precise. "Data" refers here to the ELIXIR Data Platform - (minor) I guess the type of metadata illustrated in the figure is restricted to biological metadata, that is indeed collected during sampling, however, technical metadata such as the sequencing method used, the type of instrument or the library layout, are collected during the processing of the sample not only the sampling itself. This is everything from phenotypic, sample conditions, experimental methods - (minor) There should be an arrow from "Archive" back to "Data" when the data produced is deposited and then contributes back to public repositories? ### Interactions with other ELIXIR Communities * (major) The sentence "Similarly, many of the biodiversity approaches use marker gene amplification for studying environmental DNA (eDNA)" is redundant with the previous one and introduces a different nomenclature that was not used before (marker gene amplication vs metabarcoding) and the differences, if any, are not explained. I would suggest to remove the sentence. We have modified the sentence to refer to barcoding rather than removing the sentence. We feel it is important to mention eDNA. * (major) The term "Isolation of genomes" is misleading and I guess the authors used it as a shorthand for "The isolation of bacteria, its DNA extraction, genome sequencing and their annotation". Please rephrase to avoid misinterpretation. Plus I would argue that these steps are also done when deconstructing microbiomes via cultivation strategies and are therefore not so out-of-scope. This sentence has been removed as the reviewer is correct that cultivation of bacteria (and other organisms) from microbiomes is becoming an increasing trend. * (minor) "yet each one of these areas is far greater in scientific scope" feels exaggerated and vague. It should be rephrased. A suggestion is: "yet each one of these areas is too complex to be tacked individually" We have amended the sentence accordingly. * (minor) There is no link nor transition between the paragraph that starts with "In summary, " before the Table 3 and the paragraph after that starts with "Similarly, microbiome research". Please rephrase or edit to connect the two sections. Plus, whilst this is good to have a concrete example in the "Similarly" paragraph, the paragraph before was very broad and doing a summary. I would suggest to try reordering the two paragraphs and bring the example earlier for a smoother transition. The paragraphs have been reordered as suggested and slightly amended to improve readability. * (minor) Typo "Similarly, microbiome research has many translation al aspects" Fixed * (minor) The PET acronym is detailed but actually used only once, I am unsure if this is necessary. There is another instance of this abbreviation, so we have retained this abbreviation. ### Table 2 - (major) The text in the "Interaction" column needs rework as it does not use a consistent wording and could be more to the point, especially as it is a complement to the main text: - The Food and Nutrition entry is a question. - The Galaxy entry has an unnecessary return carriage, and a unspecific "ongoing evaluation study" that needs to be clarified. - The Plant Science entry starts with a generic sentence that could be in the introduction or removed for clarity. Plus I would add a clarification that "plants maintain or not their microbial communities across generations." - (minor) The status (e.g., currently active, inactive, planned, etc.) of each ELIXIR Communities would have been appreciated as it is missing also from the Figure 1. - (minor) The mention of "the field" for the Federated human data community is vague in a manuscript about gathering communities, which research field is implied? If this is the human microbiome research field as a whole, please indicate. In response to all of the above comments, we have substantially reworked table 2. The only comment that we have not addressed so specifically, is whether an activity is in progress or not, because the activities ebb and flow, and it is not always easy to say when there is a start or an end. However, being somewhat more focused in the content, we feel this revised table provides a stronger view of the direction that the community wants to take with the other communities. ### Table 3 - (major) Similarly to the Table 2 major comment, there is a lack of consistent wording in the "Aim" column that makes the Table 3 not as impactful as it should and could be. Maybe the authors could extract common/distinct features from each of these initiatives as an alternative way to the "Aim" column. A couple of suggestions for these features would be: is it a national initiative?, does it relates to data storage, data analysis, training? is it linked to ELIXIR? - (major) I was surprised not to see the NMDC listed in the table 3, especially when it is discussed in the main text. What about the NCCR Microbiomes initiative in Switzerland? I can understand that some initiatives are not included for space reason, but maybe state it in the legend of the table. - (minor) Is this table sorted? It seems not, but it could be by acronyms or names. - (minor) The aim for the NFDI4Microbiota is way too big a paragraph. The authors should reduce it for conciseness. - (minor) The Metaproteomics Initiative entry has a hyperlink and a reference when none of the others have. Please homogenise. - (minor) Some entries have country listed and some not. Please homogenise. - (minor) I am not questioning the existence of the European Reference Genome Atlas here, but how best to phrase its relevance to ELIXIR in the manuscript. There seems to have no prokaryotes genomes in their atlas, however, there seems to be a trove of fungi and protists genomes which are usually said to be understudied in microbiome. So I think there is a missed opportunity for the authors here to make the most out of this entry. - (minor) "With its headquarters in Bari (Apulia region)," seems irrelevant in the context of the table, please remove. As with Table 2, we have substantially reworked Table 3, including ordering by the reach of the effort, harmonising the language and reducing text. We have also included the relevance of the activity to the community. We have not included NMDC, as this is a US effort, and while important to the community, this table is focused on Europe efforts. ### Data - (minor) The end of the following sentence is redundant as the INSDC was introduced earlier already "INSDC, which in collaboration with the National Institute of Genetics DNA DataBank of Japan (DDBJ) and the United States National Center for Biotechnology’s (NCBI) GenBank and Sequence Read Archive (SRA), facilitate the deposition and global exchange of sequence data.". Please adjust accordingly. We have removed the redundancy here. - (minor) "A current challenge facing the field is connecting different multi ‘omics data that have been derived from the same sample." Is this going to be tackled by the ELIXIR Microbiome Community? If so, I would state that this is part of its objectives. We feel this is likely to be a joint effort across different communities. Thus, we have retained the sentence as is. - (minor) The end of sentence "[...]overarching context to the experiment, which can be important for meta-analyses." seems like an euphemism, I would suggest to replace with "[...]overarching context to the experiment, re-analyses or meta-analyses." to include re-analyses as well. We have added “re-analyses or” as suggested. - (minor) "We will continue to promote such approaches, enriching metadata wherever possible." Is this going to be done via the CDCH? While CDCH offers one approach, we believe that there are many different ways that metadata may be enriched, first by making scientists more aware of the need of submitting meta, promoting different, more specific checklists (e.g. STORMS or MicroB3) and mining metadata from published literature. This, specifically highlighting CDCH would not be appropriate. - (minor) "The ELIXIR Microbiome Community will also work to move the Marine Metagenomics domain in the RDMKit towards a more general Microbiome domain." What is the RDMKit? It is not explained, nor cited nor mentioned again. This has been expanded to the ELIXIR Research Data Management Kit and a link has been added to the text. ### Tools - (minor) "will increase their use of BioContainers" should be "will increase the use of BioContainers" Corrected. - (minor) "In order to make tools findable by the end users, the Community" Done. - (minor) "workflow descriptions (e.g. Snakemake, CWL, Nextflow)" None of them have their references cited, is it an omission or space limitation? Added references - (minor) "A current joint effort between the Microbiome and Galaxy Communities" Fixed ### Benchmarking - (major) "Benchmarking" it is the only item at this hierarchy level, meaning that this is useless for structuring the text. Please edit. Thank you for pointing this out. We have changed the sub-sub-heading into an introductory sentence to this paragraph. - (minor) Review reference 2 published a recent review with guidelines to learn from that could have its place in this paragraph. - (minor) in the sentence: "As the Microbiome Community establishes, we will develop a broader understanding of the requirements of the Community , feed this to the Tools Platform, as well as seek opportunities to interact with the Tools Platform to capture the diversity of tools and their utility via such benchmarking activities.", is this the ELIXIR Microbiome Community, or the wide community of microbiome researchers? Is it to mean that the ELIXIR Microbiome Community is going to act as an interface between microbiome researchers and ELIXIR Tools/Infrastructure? We clarified this to mean the wider community of microbiome researchers. ### Compute - (minor) The first sentence would fit better in the introduction. ### Interoperability - (minor) The reference 57 should be at the end of the sentence starting with "This effort was paralleled" not in the middle. This citation has been moved to the end of the sentence. - (minor) Reference 58 should be removed at it is a duplicate of reference 47. References have been fixed - (minor) Use the full text "Global Alliance for Genomics and Health" instead of GA4GH. This abbreviation has been expanded. - (minor) "is in the process of applying to be a n ELIXIR Recommended Interoperability Resource." ELIXIR has been added to this sentence. - (minor) In the sentence: "This will require the development of new data Interoperability layers for data resources that are not normally focused in Microbiome data" I think the authors meant "used" instead of "focused", and "microbiome" instead of "Microbiome". We have updated the sentence accordingly. ### Training - (minor) In "Platforms such as MGnify support large-scale services for most, if not all, steps of a microbiome study", I would suggest to remove the "if not all". Agreed. - (minor) "with areas of expertise covering different environments, ‘omics approaches and data analysis pathways." I think the authors meant "learning paths", and I would refrain from using "pathways" as it also has a biological meaning. Actually, we do mean data analysis, but have changed to data analysis strategies. This is in reference to the fact that there may be multiple different ways of analysing a data type, for example metagenomics can be analysed using tools such as Kraken or MetaPhlAn4, to full assembly, gene calling and functional analysis. ### Context with other international initiatives - (major) Whilst I appreciated the emphasis that no initiative exists on its own, I feel the first paragraph on the Genom ic (please correct the typo) Standards Consortium feels lengthy for a manuscript whose topic is not the GSC. I would advise to summarize. In this respect, the second paragraph is particularly relevant to a tangible collaboration between ELIXIR Microbiome Community and GSC. The typo has been corrected. - (major) The NMDC is discussed in this section but not part of the Table 3. Table 3 lists pan-European efforts, whereas the NMDC is a US specific initiative. Thus, we have not added this to the table. - (minor) Is the mentioned M5 project still active as the website's last update is 2012? While the website has not been updated, there are ongoing efforts to expand and revitalise this effort. RDF is a member of the GSC board and is promoting aspects of the M5 initiative. While the name may change, we feel this is useful to keep. However, as part of condensing the overall length of the manuscript, we have removed the reference to this initiative. - (minor) "Combining the activities on standards concerning workflows [...]" does this means adding and providing workflows to the microbiome research community? - (minor) Given the emphasis on the fact that MicrobiomeSupport was a program, the authors could update the readers and indicate that it is now MicrobiomeSupport Association. ### Interaction with other key data resources beyond ELIXIR - (major) There is an order issue with the main text that a proofread could solve, as MG-RAST is explained and cited in the first paragraph but already mentioned upstream of the main text in the Interoperability section. This has now been fixed. ### Specific challenges and objectives of the ELIXIR Microbiome Community - (major) The strong claim "it is widely accepted that current short-read assembly-based methods do not generally work as well for soil microbiomes" would probably need at least one reference. Added - (major) "(iii) there is no centralised database collecting the millions of viral sequences". It seems to be the case indeed, and there are databases (~24) out there as recently compiled in Review reference 3. How ELIXIR Microbiome Community plans to integrate/aggregate these resources in a non-duplicating manner? - (minor) The word "through" is superfluous in the the sentence that starts with "This current limitation, [...]" and can be removed. We have removed the sentence “through”. - (minor) The CAMI was already explained and cited above, so the already defined acronym can be used. This has been fixed. - (minor) Precise the area in :"Additionally, another key area of development of taxonomy [...]". Added - (minor) The sentence "Viruses, particularly those that infect bacteria, are found ubiquitously in all environments and play critical roles in community dynamics." belongs in an introduction, not so downstream of the manuscript. - (minor) The sentence starting the sixth paragraph could be precised as "The increase in metagenomic assemblies has resulted in a parallel increase in the number of predicted protein sequences, with sets of non-redundant proteins now in the billions." Added “predicted” to the sentence. - (minor) Fix typo in "that are undetectable by current sequence based methods." Fixed - sequenced -> sequence - (minor) Would it make sense to also be able to access representative, of clusters for instance? "[...] develop new infrastructural frameworks for accessing slices of the data or adequate representatives based on the requirements." Added - (minor) What is the " expanded Microbiome Community"? It was never mentioned before. Removed “expanded” from the sentence. - (minor) A few comments on the sentence: "Within the Community, we will develop and promote standards around the analysis provenance (analytical metadata),". In my opinion and how it was already stated in the manuscript, it would make more sense to promote existing standards first and then develop if need be. There is no mention of other type of metadata, so the "analytical metadata" precision seems superfluous. I would suggest: "Within the Community, we will promote and develop standards regarding the analysis provenance," We agree with the reviewer's comment and have re-ordered accordingly. - (minor) Precise the term forms in "[...] ensuring that comput ing resources are accessible for performing the different forms of data analysis [...]", do the authors mean types of/steps in the data analysis? “compute resources” is an accepted resource. We have removed the “different forms of” as we feel it is a little superfluous. We were meaning the different strategies, but it does not add to the sentence. - (minor) Same argument as before regarding reinventing the wheel, I would swap the part of the sentence: "This may require the extensions to existing databases or development of new ones , but it requires an agreement from the research community to adopt them." We agree, and have amended accordingly. - (minor) This part "Metaproteomics aims to elucidate the functional and taxonomic interplay of proteins in microbiomes," should have been in the Table 1, or to reuse the Table 1 here. We have highlighted this in table 1. We have kept the text here to provide the context for the rest of the sentence. - (minor) The tenth paragraph of this section starts with the mention of multiple major challenges, but detail "only" one of them. I would suggest to mention some of the others challenges. - (minor) The "MIA" method is just a hyperlink, without any reference. Either cite the website accordingly or add the reference. There is not a reference, and the hyperlink goes to the GitHub site, as requested by the authors. We have included “ https://github.com/microbiome/mia ” for completeness sake. ### Table 4 - (major) I was surprised to see that the objective "Foster international collaborations between other resources providers and databases to ensure global harmonisation of e-infrastructures for microbiome research" was long-term, as I would have imagined that a gap analysis would be short-term to ensure we do not reinvent the wheel, especially given the others initiatives discussed in the manuscript. - (minor) If the Objective column starts with action verbs (which is a good idea), then it should be "Survey the needs" instead of "Survey of needs". We have amended accordingly. - (minor) Specify the type of workflow with "Address knowledge gaps in generating and adopting data analysis workflows" Added - (minor) The verb is missing in " Teach a dvanced containerisation and cloud deployment" Added. - (minor) The verb is missing in " Promote data analysis through the use of services". Is this ELIXIR services in general or specific ELIXIR Microbiome Community services? Added Promote - generalised to ELIXIR services. - (minor) Use "Share" instead of "Sharing" in the "Co-ordinate" entry. Done - (minor) The "Industry connection" entry does not fit the action verb pattern. A suggestion would be "Use ELIXIR and Node forums to understand pharmaceutical and biotechnological demands and current limitations impacting this sector." as the first sentence felt generic. Done - (minor) The verb is missing in " Design targeted training for different microbiome communities" Done - (minor) Use "findability" instead of "discoverability" for consistency and to fit with the FAIR principles. Done . - (minor) Reorder the sentence to start with the action verb: " E stablish new standards for microbiome research, particularly with respect to data analysis reporting and contextual metadata reporting in conjunction with GSC " Done - (minor) Correct "Established" to "Establish" in the "Promoting new approaches" entry. Done Editorial comments: - (major) "Community" is used in upper-case and this is unclear in many instances whether the ELIXIR Marine Metagenomics Community is referred to, the ELIXIR Microbiome Community, or the broader microbiome research community. I suggest to use consistently the full term for the sake of transparency. Abbreviations like EMMC and EMC could be even more misleading in my opinion. We have prefixed all instances of “Community” with their explicit community name. - (major) The structure of the white paper is not evident as the hierarchy is indicated only by change in font size, and some sections are quite lengthy for a white paper that is supposed to be concise. I understand that this is a constraint from the Editor, but see Review reference 4 for a white paper with a more clearer structure. An alternative could be to use numbered sections. We have introduced section numbers to aid the structuring of the manuscript. - References - (minor) The very first reference is oddly formatted in the text creating an artificial and confusing end of sentence. - (minor) Title of reference 9 is truncated and should be "Methods included: standardizing computational reuse and portability with the Common Workflow Language" - (minor) The superscript numbers of the numeric style of bibliography are in many instances after the final dot (see reference 2, 12-16, 17-19, 36, 37), after a comma (see reference 10, 20, 23), a bracket (see reference 60) or semi colon (see reference 65) when they should be before any of these symbols. - (minor) In the paragraph "Interactions with other ELIXIR Communities", we jump from reference 24 to 29 when the numeric style of bibliography (that was chosen by the authors) is expected to mirror the mentions in the manuscript. Please either have the Table 2 earlier in the paper, or change the order of the references. References have been fixed The authors make the case for mutating the ELIXIR Marine Metagenomics Community initiative into an ELIXIR Microbiome Community. It is indeed timely to have an such a pan-European initiative especially as it draws on previous experience, albeit a more narrow biome. The emphasis on multi-omics integration is also a strong suit as it is a complex topic where the microbiome research community would need support and infrastructure. The authors map out existing (but not all) similar initiatives, even if it is unclear how they would work together. In my opinion, a couple of claims should be strengthened and the argumentative power of this white paper would benefit from minor reorganization, a text slim down and extra proofreading. To this end, I highlight the strengths and weaknesses of the manuscript below. ## Strengths - Table 3 is a great idea to map out the others initiatives but needs a bit rework to be straight to the point. - In the Data section, the part between "Throughout the lifetime of the ELIXIR Marine Metagenomics Community" and the end of the paragraph is really relevant and draws from the experience of the ELIXIR Community. These sentences should be more highlighted, maybe upstream. I do wonder how the authors considered how the others initiatives (mentioned in Table 3) could also contribute to the addition of metadata for a global curation effort, possibly via the mentioned Contextual Data Clearinghouse. - The data re-analysis and integration between PRIDE and MGnify (named MetaPUF) is a good use case of the efforts that the new ELIXIR Community can promote. - I did appreciate the clear objective at the end of the Compute section, to help with scaling up analyses as well as the plan for the interaction with the ELIXIR Training Platform which promise to give a boost in training the current and next generation of microbiome researchers. - I liked how the authors highlighted (using the example of viral sequences databases) that time and resources should not be wasted in duplicating works, and that initiatives such as the ELIXIR Microbiome Community can promote this. - The authors thought ahead to promote and develop methods to limit confounding factors when re-using datasets, especially in multi-omics settings, and added this to one of the ELIXIR Microbiome Community challenges. ## Weaknesses ### Introduction - (minor) The microbiome definition used in the paper was revisited in Review reference 1 to include the interactions as well as others molecules than genomes as part of the microbiome. I suggest to use this definition given the emphasis on additional omics made in the paper. We have added that the definition includes other molecules in addition to genomes. - (major) The first paragraph ends on challenges of our research community including how to make data FAIR. Given the experience gained by the ELIXIR Marine Metagenomics Community, I would have appreciated a sentence on why making our data FAIR is important (e.g., transparency to make the research more reproducible, accountability given the (public) source of funding). See major comment in next section. Thank you for this suggestion. We have added a few sentences to this effect in the manuscript. ### The scope of the ELIXIR Microbiome Community - (major) These following arguments for FAIR data should be in the Introduction section. "Moreover, when wishing to contextualise the results with similar experiments, the way a dataset has been produced and processed must be transparent to establish whether it is comparable (e.g. amplified sequence variants can only be compared when the same amplified regions are compared). Furthermore, when different methods are applied, best practices in data stewardship are required to ensure that the connectivity of the derived sequence data products, together with functional and taxonomic assertions are kept in context of the original sample/sequencing effort and associated contextual metadata." We have merged this part of the paragraph into the Introduction. - (minor) The very first sentence of this section blames researchers for misuse of a term "[...] regularly (incorrectly) used [...]". I suggest to rephrase the sentence to simply state the differences, or back up the misuse of the term with references or a survey. We have qualified this with an example of INSDC mislabelling. - (minor) "Fundamentally, the ELIXIR Microbiome Community is about providing the necessary infrastructures required to perform analysis of nucleotide sequence data derived from a microbiome" “Nucleotide” has been added to this sentence. - (minor) "Finally, and possible unique to this ELIXIR Community, is the variety of researchers". I am not sure that these features are unique to the ELIXIR Community but rather to the field of microbiome research. I suggest to tone it down by simply pointing out that the ELIXIR Community gathers a variety of researchers. Modified accordingly. ### Table 1 - (major) Metabolomics is listed in Table 1 but is omitted from the narrative in the paragraph. We have added a sentence in the paragraph. - (major) The Table 1 is not really used but could actually reduce some redundancies in the manuscript by providing a one-stop-shop to explain and detail these techniques. We have increased the cross linking to the table. ### Figure 1 - (major) The figure is early on in the manuscript and it is unclear which of the ELIXIR platforms and communities are already established, or foreseen. Especially since "current and future ELIXIR activities" is mentioned in the manuscript before referencing this figure. The Table 2 does not add more information in that regard. - (major) There is a need for different types of arrows as the same arrow represent interactions between communities or processes. Modified accordingly. - (minor) The "Data" node is quite generic, I was wondering whether it meant public repositories, institute repositories or both. Please be precise. "Data" refers here to the ELIXIR Data Platform - (minor) I guess the type of metadata illustrated in the figure is restricted to biological metadata, that is indeed collected during sampling, however, technical metadata such as the sequencing method used, the type of instrument or the library layout, are collected during the processing of the sample not only the sampling itself. This is everything from phenotypic, sample conditions, experimental methods - (minor) There should be an arrow from "Archive" back to "Data" when the data produced is deposited and then contributes back to public repositories? ### Interactions with other ELIXIR Communities * (major) The sentence "Similarly, many of the biodiversity approaches use marker gene amplification for studying environmental DNA (eDNA)" is redundant with the previous one and introduces a different nomenclature that was not used before (marker gene amplication vs metabarcoding) and the differences, if any, are not explained. I would suggest to remove the sentence. We have modified the sentence to refer to barcoding rather than removing the sentence. We feel it is important to mention eDNA. * (major) The term "Isolation of genomes" is misleading and I guess the authors used it as a shorthand for "The isolation of bacteria, its DNA extraction, genome sequencing and their annotation". Please rephrase to avoid misinterpretation. Plus I would argue that these steps are also done when deconstructing microbiomes via cultivation strategies and are therefore not so out-of-scope. This sentence has been removed as the reviewer is correct that cultivation of bacteria (and other organisms) from microbiomes is becoming an increasing trend. * (minor) "yet each one of these areas is far greater in scientific scope" feels exaggerated and vague. It should be rephrased. A suggestion is: "yet each one of these areas is too complex to be tacked individually" We have amended the sentence accordingly. * (minor) There is no link nor transition between the paragraph that starts with "In summary, " before the Table 3 and the paragraph after that starts with "Similarly, microbiome research". Please rephrase or edit to connect the two sections. Plus, whilst this is good to have a concrete example in the "Similarly" paragraph, the paragraph before was very broad and doing a summary. I would suggest to try reordering the two paragraphs and bring the example earlier for a smoother transition. The paragraphs have been reordered as suggested and slightly amended to improve readability. * (minor) Typo "Similarly, microbiome research has many translation al aspects" Fixed * (minor) The PET acronym is detailed but actually used only once, I am unsure if this is necessary. There is another instance of this abbreviation, so we have retained this abbreviation. ### Table 2 - (major) The text in the "Interaction" column needs rework as it does not use a consistent wording and could be more to the point, especially as it is a complement to the main text: - The Food and Nutrition entry is a question. - The Galaxy entry has an unnecessary return carriage, and a unspecific "ongoing evaluation study" that needs to be clarified. - The Plant Science entry starts with a generic sentence that could be in the introduction or removed for clarity. Plus I would add a clarification that "plants maintain or not their microbial communities across generations." - (minor) The status (e.g., currently active, inactive, planned, etc.) of each ELIXIR Communities would have been appreciated as it is missing also from the Figure 1. - (minor) The mention of "the field" for the Federated human data community is vague in a manuscript about gathering communities, which research field is implied? If this is the human microbiome research field as a whole, please indicate. In response to all of the above comments, we have substantially reworked table 2. The only comment that we have not addressed so specifically, is whether an activity is in progress or not, because the activities ebb and flow, and it is not always easy to say when there is a start or an end. However, being somewhat more focused in the content, we feel this revised table provides a stronger view of the direction that the community wants to take with the other communities. ### Table 3 - (major) Similarly to the Table 2 major comment, there is a lack of consistent wording in the "Aim" column that makes the Table 3 not as impactful as it should and could be. Maybe the authors could extract common/distinct features from each of these initiatives as an alternative way to the "Aim" column. A couple of suggestions for these features would be: is it a national initiative?, does it relates to data storage, data analysis, training? is it linked to ELIXIR? - (major) I was surprised not to see the NMDC listed in the table 3, especially when it is discussed in the main text. What about the NCCR Microbiomes initiative in Switzerland? I can understand that some initiatives are not included for space reason, but maybe state it in the legend of the table. - (minor) Is this table sorted? It seems not, but it could be by acronyms or names. - (minor) The aim for the NFDI4Microbiota is way too big a paragraph. The authors should reduce it for conciseness. - (minor) The Metaproteomics Initiative entry has a hyperlink and a reference when none of the others have. Please homogenise. - (minor) Some entries have country listed and some not. Please homogenise. - (minor) I am not questioning the existence of the European Reference Genome Atlas here, but how best to phrase its relevance to ELIXIR in the manuscript. There seems to have no prokaryotes genomes in their atlas, however, there seems to be a trove of fungi and protists genomes which are usually said to be understudied in microbiome. So I think there is a missed opportunity for the authors here to make the most out of this entry. - (minor) "With its headquarters in Bari (Apulia region)," seems irrelevant in the context of the table, please remove. As with Table 2, we have substantially reworked Table 3, including ordering by the reach of the effort, harmonising the language and reducing text. We have also included the relevance of the activity to the community. We have not included NMDC, as this is a US effort, and while important to the community, this table is focused on Europe efforts. ### Data - (minor) The end of the following sentence is redundant as the INSDC was introduced earlier already "INSDC, which in collaboration with the National Institute of Genetics DNA DataBank of Japan (DDBJ) and the United States National Center for Biotechnology’s (NCBI) GenBank and Sequence Read Archive (SRA), facilitate the deposition and global exchange of sequence data.". Please adjust accordingly. We have removed the redundancy here. - (minor) "A current challenge facing the field is connecting different multi ‘omics data that have been derived from the same sample." Is this going to be tackled by the ELIXIR Microbiome Community? If so, I would state that this is part of its objectives. We feel this is likely to be a joint effort across different communities. Thus, we have retained the sentence as is. - (minor) The end of sentence "[...]overarching context to the experiment, which can be important for meta-analyses." seems like an euphemism, I would suggest to replace with "[...]overarching context to the experiment, re-analyses or meta-analyses." to include re-analyses as well. We have added “re-analyses or” as suggested. - (minor) "We will continue to promote such approaches, enriching metadata wherever possible." Is this going to be done via the CDCH? While CDCH offers one approach, we believe that there are many different ways that metadata may be enriched, first by making scientists more aware of the need of submitting meta, promoting different, more specific checklists (e.g. STORMS or MicroB3) and mining metadata from published literature. This, specifically highlighting CDCH would not be appropriate. - (minor) "The ELIXIR Microbiome Community will also work to move the Marine Metagenomics domain in the RDMKit towards a more general Microbiome domain." What is the RDMKit? It is not explained, nor cited nor mentioned again. This has been expanded to the ELIXIR Research Data Management Kit and a link has been added to the text. ### Tools - (minor) "will increase their use of BioContainers" should be "will increase the use of BioContainers" Corrected. - (minor) "In order to make tools findable by the end users, the Community" Done. - (minor) "workflow descriptions (e.g. Snakemake, CWL, Nextflow)" None of them have their references cited, is it an omission or space limitation? Added references - (minor) "A current joint effort between the Microbiome and Galaxy Communities" Fixed ### Benchmarking - (major) "Benchmarking" it is the only item at this hierarchy level, meaning that this is useless for structuring the text. Please edit. Thank you for pointing this out. We have changed the sub-sub-heading into an introductory sentence to this paragraph. - (minor) Review reference 2 published a recent review with guidelines to learn from that could have its place in this paragraph. - (minor) in the sentence: "As the Microbiome Community establishes, we will develop a broader understanding of the requirements of the Community , feed this to the Tools Platform, as well as seek opportunities to interact with the Tools Platform to capture the diversity of tools and their utility via such benchmarking activities.", is this the ELIXIR Microbiome Community, or the wide community of microbiome researchers? Is it to mean that the ELIXIR Microbiome Community is going to act as an interface between microbiome researchers and ELIXIR Tools/Infrastructure? We clarified this to mean the wider community of microbiome researchers. ### Compute - (minor) The first sentence would fit better in the introduction. ### Interoperability - (minor) The reference 57 should be at the end of the sentence starting with "This effort was paralleled" not in the middle. This citation has been moved to the end of the sentence. - (minor) Reference 58 should be removed at it is a duplicate of reference 47. References have been fixed - (minor) Use the full text "Global Alliance for Genomics and Health" instead of GA4GH. This abbreviation has been expanded. - (minor) "is in the process of applying to be a n ELIXIR Recommended Interoperability Resource." ELIXIR has been added to this sentence. - (minor) In the sentence: "This will require the development of new data Interoperability layers for data resources that are not normally focused in Microbiome data" I think the authors meant "used" instead of "focused", and "microbiome" instead of "Microbiome". We have updated the sentence accordingly. ### Training - (minor) In "Platforms such as MGnify support large-scale services for most, if not all, steps of a microbiome study", I would suggest to remove the "if not all". Agreed. - (minor) "with areas of expertise covering different environments, ‘omics approaches and data analysis pathways." I think the authors meant "learning paths", and I would refrain from using "pathways" as it also has a biological meaning. Actually, we do mean data analysis, but have changed to data analysis strategies. This is in reference to the fact that there may be multiple different ways of analysing a data type, for example metagenomics can be analysed using tools such as Kraken or MetaPhlAn4, to full assembly, gene calling and functional analysis. ### Context with other international initiatives - (major) Whilst I appreciated the emphasis that no initiative exists on its own, I feel the first paragraph on the Genom ic (please correct the typo) Standards Consortium feels lengthy for a manuscript whose topic is not the GSC. I would advise to summarize. In this respect, the second paragraph is particularly relevant to a tangible collaboration between ELIXIR Microbiome Community and GSC. The typo has been corrected. - (major) The NMDC is discussed in this section but not part of the Table 3. Table 3 lists pan-European efforts, whereas the NMDC is a US specific initiative. Thus, we have not added this to the table. - (minor) Is the mentioned M5 project still active as the website's last update is 2012? While the website has not been updated, there are ongoing efforts to expand and revitalise this effort. RDF is a member of the GSC board and is promoting aspects of the M5 initiative. While the name may change, we feel this is useful to keep. However, as part of condensing the overall length of the manuscript, we have removed the reference to this initiative. - (minor) "Combining the activities on standards concerning workflows [...]" does this means adding and providing workflows to the microbiome research community? - (minor) Given the emphasis on the fact that MicrobiomeSupport was a program, the authors could update the readers and indicate that it is now MicrobiomeSupport Association. ### Interaction with other key data resources beyond ELIXIR - (major) There is an order issue with the main text that a proofread could solve, as MG-RAST is explained and cited in the first paragraph but already mentioned upstream of the main text in the Interoperability section. This has now been fixed. ### Specific challenges and objectives of the ELIXIR Microbiome Community - (major) The strong claim "it is widely accepted that current short-read assembly-based methods do not generally work as well for soil microbiomes" would probably need at least one reference. Added - (major) "(iii) there is no centralised database collecting the millions of viral sequences". It seems to be the case indeed, and there are databases (~24) out there as recently compiled in Review reference 3. How ELIXIR Microbiome Community plans to integrate/aggregate these resources in a non-duplicating manner? - (minor) The word "through" is superfluous in the the sentence that starts with "This current limitation, [...]" and can be removed. We have removed the sentence “through”. - (minor) The CAMI was already explained and cited above, so the already defined acronym can be used. This has been fixed. - (minor) Precise the area in :"Additionally, another key area of development of taxonomy [...]". Added - (minor) The sentence "Viruses, particularly those that infect bacteria, are found ubiquitously in all environments and play critical roles in community dynamics." belongs in an introduction, not so downstream of the manuscript. - (minor) The sentence starting the sixth paragraph could be precised as "The increase in metagenomic assemblies has resulted in a parallel increase in the number of predicted protein sequences, with sets of non-redundant proteins now in the billions." Added “predicted” to the sentence. - (minor) Fix typo in "that are undetectable by current sequence based methods." Fixed - sequenced -> sequence - (minor) Would it make sense to also be able to access representative, of clusters for instance? "[...] develop new infrastructural frameworks for accessing slices of the data or adequate representatives based on the requirements." Added - (minor) What is the " expanded Microbiome Community"? It was never mentioned before. Removed “expanded” from the sentence. - (minor) A few comments on the sentence: "Within the Community, we will develop and promote standards around the analysis provenance (analytical metadata),". In my opinion and how it was already stated in the manuscript, it would make more sense to promote existing standards first and then develop if need be. There is no mention of other type of metadata, so the "analytical metadata" precision seems superfluous. I would suggest: "Within the Community, we will promote and develop standards regarding the analysis provenance," We agree with the reviewer's comment and have re-ordered accordingly. - (minor) Precise the term forms in "[...] ensuring that comput ing resources are accessible for performing the different forms of data analysis [...]", do the authors mean types of/steps in the data analysis? “compute resources” is an accepted resource. We have removed the “different forms of” as we feel it is a little superfluous. We were meaning the different strategies, but it does not add to the sentence. - (minor) Same argument as before regarding reinventing the wheel, I would swap the part of the sentence: "This may require the extensions to existing databases or development of new ones , but it requires an agreement from the research community to adopt them." We agree, and have amended accordingly. - (minor) This part "Metaproteomics aims to elucidate the functional and taxonomic interplay of proteins in microbiomes," should have been in the Table 1, or to reuse the Table 1 here. We have highlighted this in table 1. We have kept the text here to provide the context for the rest of the sentence. - (minor) The tenth paragraph of this section starts with the mention of multiple major challenges, but detail "only" one of them. I would suggest to mention some of the others challenges. - (minor) The "MIA" method is just a hyperlink, without any reference. Either cite the website accordingly or add the reference. There is not a reference, and the hyperlink goes to the GitHub site, as requested by the authors. We have included “ https://github.com/microbiome/mia ” for completeness sake. ### Table 4 - (major) I was surprised to see that the objective "Foster international collaborations between other resources providers and databases to ensure global harmonisation of e-infrastructures for microbiome research" was long-term, as I would have imagined that a gap analysis would be short-term to ensure we do not reinvent the wheel, especially given the others initiatives discussed in the manuscript. - (minor) If the Objective column starts with action verbs (which is a good idea), then it should be "Survey the needs" instead of "Survey of needs". We have amended accordingly. - (minor) Specify the type of workflow with "Address knowledge gaps in generating and adopting data analysis workflows" Added - (minor) The verb is missing in " Teach a dvanced containerisation and cloud deployment" Added. - (minor) The verb is missing in " Promote data analysis through the use of services". Is this ELIXIR services in general or specific ELIXIR Microbiome Community services? Added Promote - generalised to ELIXIR services. - (minor) Use "Share" instead of "Sharing" in the "Co-ordinate" entry. Done - (minor) The "Industry connection" entry does not fit the action verb pattern. A suggestion would be "Use ELIXIR and Node forums to understand pharmaceutical and biotechnological demands and current limitations impacting this sector." as the first sentence felt generic. Done - (minor) The verb is missing in " Design targeted training for different microbiome communities" Done - (minor) Use "findability" instead of "discoverability" for consistency and to fit with the FAIR principles. Done . - (minor) Reorder the sentence to start with the action verb: " E stablish new standards for microbiome research, particularly with respect to data analysis reporting and contextual metadata reporting in conjunction with GSC " Done - (minor) Correct "Established" to "Establish" in the "Promoting new approaches" entry. Done Editorial comments: - (major) "Community" is used in upper-case and this is unclear in many instances whether the ELIXIR Marine Metagenomics Community is referred to, the ELIXIR Microbiome Community, or the broader microbiome research community. I suggest to use consistently the full term for the sake of transparency. Abbreviations like EMMC and EMC could be even more misleading in my opinion. We have prefixed all instances of “Community” with their explicit community name. - (major) The structure of the white paper is not evident as the hierarchy is indicated only by change in font size, and some sections are quite lengthy for a white paper that is supposed to be concise. I understand that this is a constraint from the Editor, but see Review reference 4 for a white paper with a more clearer structure. An alternative could be to use numbered sections. We have introduced section numbers to aid the structuring of the manuscript. - References - (minor) The very first reference is oddly formatted in the text creating an artificial and confusing end of sentence. - (minor) Title of reference 9 is truncated and should be "Methods included: standardizing computational reuse and portability with the Common Workflow Language" - (minor) The superscript numbers of the numeric style of bibliography are in many instances after the final dot (see reference 2, 12-16, 17-19, 36, 37), after a comma (see reference 10, 20, 23), a bracket (see reference 60) or semi colon (see reference 65) when they should be before any of these symbols. - (minor) In the paragraph "Interactions with other ELIXIR Communities", we jump from reference 24 to 29 when the numeric style of bibliography (that was chosen by the authors) is expected to mirror the mentions in the manuscript. Please either have the Table 2 earlier in the paper, or change the order of the references. References have been fixed Competing Interests: No competing interests were disclosed. Close Report a concern COMMENT ON THIS REPORT Comments on this article Comments (0) Version 2 VERSION 2 PUBLISHED 08 Jan 2024 ADD YOUR COMMENT Comment keyboard_arrow_left keyboard_arrow_right Open Peer Review Reviewer Status info_outline Alongside their report, reviewers assign a status to the article: Approved The paper is scientifically sound in its current form and only minor, if any, improvements are suggested Approved with reservations A number of small changes, sometimes more significant revisions are required to address specific details and improve the papers academic merit. Not approved Fundamental flaws in the paper seriously undermine the findings and conclusions Reviewer Reports Invited Reviewers 1 2 3 Version 2 (revision) 08 Sep 25 read read Version 1 08 Jan 24 read read Charlie Pauvert , University Hospital of RWTH, Aachen, Germany Almut Heinken , University of Lorraine, Lorraine, France Lauren Lui , E O Lawrence Berkeley National Laboratory, Berkeley, USA Comments on this article All Comments (0) Add a comment Sign up for content alerts Sign Up You are now signed up to receive this alert Browse by related subjects keyboard_arrow_left Back to all reports Reviewer Report 0 Views copyright © 2025 Lui L. This is an open access peer review report distributed under the terms of the Creative Commons Attribution License , which permits unrestricted use, distribution, and reproduction in any medium, provided the original work is properly cited. 06 Oct 2025 | for Version 2 Lauren Lui , E O Lawrence Berkeley National Laboratory, Berkeley, California, USA 0 Views copyright © 2025 Lui L. This is an open access peer review report distributed under the terms of the Creative Commons Attribution License , which permits unrestricted use, distribution, and reproduction in any medium, provided the original work is properly cited. format_quote Cite this report speaker_notes Responses (0) Approved With Reservations info_outline Alongside their report, reviewers assign a status to the article: Approved The paper is scientifically sound in its current form and only minor, if any, improvements are suggested Approved with reservations A number of small changes, sometimes more significant revisions are required to address specific details and improve the papers academic merit. Not approved Fundamental flaws in the paper seriously undermine the findings and conclusions Summary Finn et al . have written a comprehensive white paper on the establishment of the ELIXIR Microbiome Community. They outline the purpose of establishing the ELIXIR Microbiome Community, the history of its formation, interaction with other ELIXIR communities, plans for data management and computational resources, and challenges in implementation. I found the manuscript informative about the intent of the ELIXIR Microbiome Community initiatives, and it was useful to have the historical and international context included. There is a lot of information packed into this manuscript! I only have minor issues with the manuscript that I think can easily be addressed. Major Comments On page 6, the authors state that the ELIXIR Microbiome Community is about providing the necessary infrastructure for “nucleotide sequence data” from a microbiome. In the context of the entire manuscript this statement is confusing because proteomics and metabolomics are also discussed in relation to the entire ELIXIR metacommunity. Should this statement be revised or can there be some clarification added here? Why I put “no” for “Are all factual statements correct and adequately supported by citations?” and why I put “no” for “Are arguments sufficiently supported by evidence from the published literature” There are some blanket statements in the manuscript that are consistent with the literature but there aren’t any citations. There are a few statements that are incorrect. Introduction. There are many statements here that could be supported by literature references. For example, the statement that viruses that infect bacteria play critical roles in community dynamics could easily have a few references. This statement could also be expanded to include archaeal and eukaryotic viruses. For example, viruses that infect eukaryotic algae are also very important in the dynamics of algal blooms. The statement about the impact of dysbiosis of microbiomes could also have a few references. One Health. One Health is used in the manuscript but not defined nor does it have a citation. Please describe briefly and include a citation. Page 8. Microbial dark matter. There’s at least two places that mentions “microbial dark matter” but don’t define it (pages 8, and 17). Capitalization is not consistent between these two instances. Please include a reference and/or define this term. Page 16. Most of the content on this page could use more references: Section on long reads. This secion could include citations on long read metagenomics studies, notably the publication of Microflora Danica from Mads Albertsen’s lab (https://www.nature.com/articles/s41564-025-02062-z). I acknowledge that this came out fairly recently and the authors may have drafted this section before this publication came out, but there are many other long read metagenome papers that can also be cited. My second issue with this section is that long reads don’t really mitigate the computational burden – you still can’t assemble the metagenomes on a laptop and nanopore basecalling needs GPUs. Long reads assist with genome assembly because the reads are longer than genomic repeats, eliminating the ambiguities of short read assembly. This is discussed by Ryan Wick on his github page (https://github.com/rrwick/Unicycler?tab=readme-ov-file) and in the introduction and conclusion of this paper https://pmc.ncbi.nlm.nih.gov/articles/PMC8172020/. Paragraph on MAGs. MAGs should be defined as “metagenome-assembled genomes” not “environmental genomes.” MAGs can also be from clinical metagenomes. In the same sentence as above, the statement “has allowed the identification of thousands of specific functions” is unclear as to what it refers to. If this refers to protein gene function, then all of these would be by computational prediction and not really identification of new functions. A citation would be helpful here. Paragraph about classification of MAGs. Please add “taxonomic” in front of “classification” in the first sentence for clarity There’s been huge changes in prokaryotic taxonomy in recent years, and there are no citations to the committee reports here or opinion pieces of the use of “candidatus” when naming organisms. It is worth a few sentences of discussion here or added references. Paragraph on viruses Please add a few references for the statements about the challenges of analyzing viral microbiomes. The statement that there is no centralized database is incorrect. IMG/VR, which is mentioned in the article, has millions of viral genomes. Some statements may need to be qualified or rephrased. Page 11, the statement “Even relatively simple workflows that perform metagenomics assembly are computationally heavy.” This statement is heavily qualified based on someone’s experience with and access to computational resources, as well as the amount of data to be processed. It would also be better to use something like “resource-intensive” instead of “heavy.” If this statement could also be changed to something like “most metagenomes cannot be assembled on a laptop” to provide more context for the reader. Page 17, last paragraph. The sentence about the challenges in metaproteomics mentions proteins that are not identified in corresponding metagenomes and calls them “so called unidentified proteins of unknown function.” This seems to be a bit of a non sequitur because proteins of unknown function are those that cannot have a function assigned computationally, not because they are not found in the genomic data. It is also possible that "so called unidentified proteins of unknown function" is a separate list item - if so then numbers should be added into this sentence for clarity. Page 14, the statement beginning with “For example, it is widely accepted that current short-read assembly-based methods…” The accuracy of this statement depends on access to computational resources. Terabase scale metagenomes have been done ( https://www.nature.com/articles/s41598-020-67416-5 ). Perhaps change “computationally intractable” to “highly computationally resource intensive.” Minor comments Use of “ELIXIR Microbiome Community” In most places, “ELIXIR Microbiome Community” is used, but in some places it is only “Microbiome Community.” Please change all of these instances to “ELIXIR Microbiome Community” for consistency. ELIXIR Platforms Please add a table of what each of the platforms do or describe briefly in the text (top of page 5) ASV ASV is the abbreviation for “amplicon sequence variant” not “amplified sequence variant” Typos Page 11, “non-sequenced” should be “non-sequence”’ Page 11, bottom of paragraph 4, should microbiome be lower case? Citations INSDC needs a citation (bottom of page 5) Figure 1 I assume that the starting point for this diagram is the “Microbiomes in their environment” box. This could either be a unique color (so something other than white, blue, or orange) or larger than the other boxes to visually indicate that this is the starting point of the data/material workflow. Abbreviations Given the nature of this article, there are many acronyms and the definitions of a few were missed. GBIF – page 8 Country abbreviations in Table 3 (FR, IT, DE) Grammar There are some grammatical issues sprinkled throughout the manuscript but are easily fixed: Colons and semicolons are misused throughout this manuscript given the large number of lists. Colons only come after complete clauses. Semi-colons can only separate two independent clauses. For example, on page 4, the last sentence of the first paragraph should be something like “Key challenges facing the microbiome research community are how to (i) appropriately store the data, (ii) informatically process, integrate, compare and interpret microbiome-derived data, and (iii) how to make the data findable, accessible, interoperable, and reusable, i.e., FAIR.” The colon was removed and commas used to separate the list items. This is also an issue in the third paragraph of section 2.2.2.5 Training, first paragraph of section 3, page 16 fifth paragraph, fourth paragraph on page 17, last row of table 3. Use commas after i.e. and e.g. I think that last sentence of the third paragraph on page 6 may need to have an “and” in front of “to biotechnology”. Page 17, last paragraph. The sentence about the challenges in metaproteomics could use numbers for the items. As it is currently structured, the sentence is confusing to read. Page 18. Second paragraph would benefit from the use of the oxford comma. Example is add a comma after “image data.” The last sentence of this paragraph is awkwardly worded. Is the topic of the opinion article discussed accurately in the context of the current literature? Yes Are all factual statements correct and adequately supported by citations? No Are arguments sufficiently supported by evidence from the published literature? No Are the conclusions drawn balanced and justified on the basis of the presented arguments? Yes Competing Interests No competing interests were disclosed. Reviewer Expertise Microbial ecology, Bioinformatics, Metagenomics I confirm that I have read this submission and believe that I have an appropriate level of expertise to confirm that it is of an acceptable scientific standard, however I have significant reservations, as outlined above. reply Respond to this report Responses (0) Lui L. Peer Review Report For: Establishing the ELIXIR Microbiome Community [version 1; peer review: 1 approved, 1 approved with reservations] . F1000Research 2024, 13 (ELIXIR):50 ( https://doi.org/10.5256/f1000research.185171.r412934) NOTE: it is important to ensure the information in square brackets after the title is included in this citation. The direct URL for this report is: https://f1000research.com/articles/13-50/v2#referee-response-412934 keyboard_arrow_left Back to all reports Reviewer Report 0 Views copyright © 2025 Pauvert C. This is an open access peer review report distributed under the terms of the Creative Commons Attribution License , which permits unrestricted use, distribution, and reproduction in any medium, provided the original work is properly cited. 16 Sep 2025 | for Version 2 Charlie Pauvert , University Hospital of RWTH, Aachen, Germany 0 Views copyright © 2025 Pauvert C. This is an open access peer review report distributed under the terms of the Creative Commons Attribution License , which permits unrestricted use, distribution, and reproduction in any medium, provided the original work is properly cited. format_quote Cite this report speaker_notes Responses (0) Approved info_outline Alongside their report, reviewers assign a status to the article: Approved The paper is scientifically sound in its current form and only minor, if any, improvements are suggested Approved with reservations A number of small changes, sometimes more significant revisions are required to address specific details and improve the papers academic merit. Not approved Fundamental flaws in the paper seriously undermine the findings and conclusions Thanks to the authors for addressing all my comments and suggestions. They successfully restructured and amended the tables 2 and 3 and the article for an increased impact to highlight how the scientific community can benefit from the ELIXIR Microbiome Community. On a small minor note, it seems possible to acknowledge the authors of the MIA package more than the URL by citing their work see https://microbiome.github.io/mia/authors.html#citation (and the added citation #1 via the form). References 1. Borman, T Ernst, F Shetty, S Lahti, et al.: mia: Microbiome analysis. R package version 1.17.9. https://microbiome.github.io/mia/ . 2025. Competing Interests No competing interests were disclosed. Reviewer Expertise microbiome, bioinformatics, reproducible research, data and metadata standards I confirm that I have read this submission and believe that I have an appropriate level of expertise to confirm that it is of an acceptable scientific standard. reply Respond to this report Responses (0) Pauvert C. Peer Review Report For: Establishing the ELIXIR Microbiome Community [version 1; peer review: 1 approved, 1 approved with reservations] . F1000Research 2024, 13 (ELIXIR):50 ( https://doi.org/10.5256/f1000research.185171.r412360) NOTE: it is important to ensure the information in square brackets after the title is included in this citation. The direct URL for this report is: https://f1000research.com/articles/13-50/v2#referee-response-412360 keyboard_arrow_left Back to all reports Reviewer Report 0 Views copyright © 2024 Heinken A. This is an open access peer review report distributed under the terms of the Creative Commons Attribution License , which permits unrestricted use, distribution, and reproduction in any medium, provided the original work is properly cited. 11 May 2024 | for Version 1 Almut Heinken , University of Lorraine, Lorraine, France 0 Views copyright © 2024 Heinken A. This is an open access peer review report distributed under the terms of the Creative Commons Attribution License , which permits unrestricted use, distribution, and reproduction in any medium, provided the original work is properly cited. format_quote Cite this report speaker_notes Responses (1) Approved info_outline Alongside their report, reviewers assign a status to the article: Approved The paper is scientifically sound in its current form and only minor, if any, improvements are suggested Approved with reservations A number of small changes, sometimes more significant revisions are required to address specific details and improve the papers academic merit. Not approved Fundamental flaws in the paper seriously undermine the findings and conclusions In this article, the ELIXIR Microbiome community is described. The community was established in 2015 as the Marine Metagenomics community and was recently expanded to the scope of metagenomics research for multiple areas such as human health, agriculture, and ecology. The authors then describe their perspective for the community. One focus area will be the promotion of best practices for metagenomics analyses by benchmarking tools. Interoperability between different tools will also be facilitated. Another focus will be encouraging data sharing and reuse and providing data storage platforms. Finally, links to existing ELIXIR initiatives in related areas as well as with international initiatives will be established. Overall, this is a very clear, concise, and informative review. The scope and aims of ELIXIR Microbiome are well-described and detailed. The short-term and long-term objectives are also clearly described. Specific comments: I appreciate the links to other ELIXIR initiatives in Table 2. I would be particularly interested in more detail on the integration with the Systems Biology community. How would the ability to reuse multi-omics data for systems biology approaches be improved? It is a bit unclear to me which is the proposed data platforms, tools, and training will be freely available to non-ELIXIR members. Regarding providing metadata of samples, how will GDPR regulations be handled for human samples? Is the topic of the opinion article discussed accurately in the context of the current literature? Yes Are all factual statements correct and adequately supported by citations? Yes Are arguments sufficiently supported by evidence from the published literature? Yes Are the conclusions drawn balanced and justified on the basis of the presented arguments? Yes Competing Interests No competing interests were disclosed. Reviewer Expertise Systems biology I confirm that I have read this submission and believe that I have an appropriate level of expertise to confirm that it is of an acceptable scientific standard. reply Respond to this report Responses (1) Author Response 10 Sep 2025 Bérénice Batut, Bioinformatics Group, Department of Computer Science, Albert-Ludwigs-University Freiburg, Freiburg, Germany In this article, the ELIXIR Microbiome community is described. The community was established in 2015 as the Marine Metagenomics community and was recently expanded to the scope of metagenomics research for multiple areas such as human health, agriculture, and ecology. The authors then describe their perspective for the community. One focus area will be the promotion of best practices for metagenomics analyses by benchmarking tools. Interoperability between different tools will also be facilitated. Another focus will be encouraging data sharing and reuse and providing data storage platforms. Finally, links to existing ELIXIR initiatives in related areas as well as with international initiatives will be established. Overall, this is a very clear, concise, and informative review. The scope and aims of ELIXIR Microbiome are well-described and detailed. The short-term and long-term objectives are also clearly described. We would like to thank the reviewer for their positive comments concerning the ELIXIR microbiome community papers. Below we address their specific comments. Specific comments: I appreciate the links to other ELIXIR initiatives in Table 2. I would be particularly interested in more detail on the integration with the Systems Biology community. How would the ability to reuse multi-omics data for systems biology approaches be improved? We have added a few sentences to expand how the integration might be improved. It is a bit unclear to me which is the proposed data platforms, tools, and training will be freely available to non-ELIXIR members. All of the proposed activities are freely available to non-ELIXIR members. ELIXIR does not put explicit boundaries on who can use, but engagement with some ELIXIR events may be preferentially given to scientists coming from ELIXIR member states and funds from ELIXIR funding schemes would be restricted to member states. Regarding providing metadata of samples, how will GDPR regulations be handled for human samples? The landscape concerning human microbiome samples is complicated. Currently there is little consensus across Europe whether human microbiomes should be under controlled access. Similarly, the GDPR landscape is also complicated and not entirely independent. There are already established routes for suppression of data (should an individual wish to be forgotten), which can be propagated to other databases (e.g. MGnify will remove analyses associated with a suppressed sequence dataset). This is clearly an area that will need to be developed as part of the ELIXR community, as no specific solutions have been agreed, we would prefer not to comment about how they will be handled in this manuscript. View more View less Competing Interests No competing interests were disclosed. reply Respond Report a concern Heinken A. Peer Review Report For: Establishing the ELIXIR Microbiome Community [version 1; peer review: 1 approved, 1 approved with reservations] . F1000Research 2024, 13 (ELIXIR):50 ( https://doi.org/10.5256/f1000research.158321.r251099) NOTE: it is important to ensure the information in square brackets after the title is included in this citation. The direct URL for this report is: https://f1000research.com/articles/13-50/v1#referee-response-251099 keyboard_arrow_left Back to all reports Reviewer Report 0 Views copyright © 2024 Pauvert C. This is an open access peer review report distributed under the terms of the Creative Commons Attribution License , which permits unrestricted use, distribution, and reproduction in any medium, provided the original work is properly cited. 14 Mar 2024 | for Version 1 Charlie Pauvert , University Hospital of RWTH, Aachen, Germany 0 Views copyright © 2024 Pauvert C. This is an open access peer review report distributed under the terms of the Creative Commons Attribution License , which permits unrestricted use, distribution, and reproduction in any medium, provided the original work is properly cited. format_quote Cite this report speaker_notes Responses (1) Approved With Reservations info_outline Alongside their report, reviewers assign a status to the article: Approved The paper is scientifically sound in its current form and only minor, if any, improvements are suggested Approved with reservations A number of small changes, sometimes more significant revisions are required to address specific details and improve the papers academic merit. Not approved Fundamental flaws in the paper seriously undermine the findings and conclusions The authors make the case for mutating the ELIXIR Marine Metagenomics Community initiative into an ELIXIR Microbiome Community. It is indeed timely to have an such a pan-European initiative especially as it draws on previous experience, albeit a more narrow biome. The emphasis on multi-omics integration is also a strong suit as it is a complex topic where the microbiome research community would need support and infrastructure. The authors map out existing (but not all) similar initiatives, even if it is unclear how they would work together. In my opinion, a couple of claims should be strengthened and the argumentative power of this white paper would benefit from minor reorganization, a text slim down and extra proofreading. To this end, I highlight the strengths and weaknesses of the manuscript below. ## Strengths - Table 3 is a great idea to map out the others initiatives but needs a bit rework to be straight to the point. - In the Data section, the part between "Throughout the lifetime of the ELIXIR Marine Metagenomics Community" and the end of the paragraph is really relevant and draws from the experience of the ELIXIR Community. These sentences should be more highlighted, maybe upstream. I do wonder how the authors considered how the others initiatives (mentioned in Table 3) could also contribute to the addition of metadata for a global curation effort, possibly via the mentioned Contextual Data Clearinghouse. - The data re-analysis and integration between PRIDE and MGnify (named MetaPUF) is a good use case of the efforts that the new ELIXIR Community can promote. - I did appreciate the clear objective at the end of the Compute section, to help with scaling up analyses as well as the plan for the interaction with the ELIXIR Training Platform which promise to give a boost in training the current and next generation of microbiome researchers. - I liked how the authors highlighted (using the example of viral sequences databases) that time and resources should not be wasted in duplicating works, and that initiatives such as the ELIXIR Microbiome Community can promote this. - The authors thought ahead to promote and develop methods to limit confounding factors when re-using datasets, especially in multi-omics settings, and added this to one of the ELIXIR Microbiome Community challenges. ## Weaknesses ### Introduction - (minor) The microbiome definition used in the paper was revisited in Review reference 1 to include the interactions as well as others molecules than genomes as part of the microbiome. I suggest to use this definition given the emphasis on additional omics made in the paper. - (major) The first paragraph ends on challenges of our research community including how to make data FAIR. Given the experience gained by the ELIXIR Marine Metagenomics Community, I would have appreciated a sentence on why making our data FAIR is important (e.g., transparency to make the research more reproducible, accountability given the (public) source of funding). See major comment in next section. ### The scope of the ELIXIR Microbiome Community - (major) These following arguments for FAIR data should be in the Introduction section. "Moreover, when wishing to contextualise the results with similar experiments, the way a dataset has been produced and processed must be transparent to establish whether it is comparable (e.g. amplified sequence variants can only be compared when the same amplified regions are compared). Furthermore, when different methods are applied, best practices in data stewardship are required to ensure that the connectivity of the derived sequence data products, together with functional and taxonomic assertions are kept in context of the original sample/sequencing effort and associated contextual metadata." - (minor) The very first sentence of this section blames researchers for misuse of a term "[...] regularly (incorrectly) used [...]". I suggest to rephrase the sentence to simply state the differences, or back up the misuse of the term with references or a survey. - (minor) "Fundamentally, the ELIXIR Microbiome Community is about providing the necessary infrastructures required to perform analysis of nucleotide sequence data derived from a microbiome" - (minor) "Finally, and possible unique to this ELIXIR Community, is the variety of researchers". I am not sure that these features are unique to the ELIXIR Community but rather to the field of microbiome research. I suggest to tone it down by simply pointing out that the ELIXIR Community gathers a variety of researchers. ### Table 1 - (major) Metabolomics is listed in Table 1 but is omitted from the narrative in the paragraph. - (major) The Table 1 is not really used but could actually reduce some redundancies in the manuscript by providing a one-stop-shop to explain and detail these techniques. ### Figure 1 - (major) The figure is early on in the manuscript and it is unclear which of the ELIXIR platforms and communities are already established, or foreseen. Especially since "current and future ELIXIR activities" is mentioned in the manuscript before referencing this figure. The Table 2 does not add more information in that regard. - (major) There is a need for different types of arrows as the same arrow represent interactions between communities or processes. - (minor) The "Data" node is quite generic, I was wondering whether it meant public repositories, institute repositories or both. Please be precise. - (minor) I guess the type of metadata illustrated in the figure is restricted to biological metadata, that is indeed collected during sampling, however, technical metadata such as the sequencing method used, the type of instrument or the library layout, are collected during the processing of the sample not only the sampling itself. - (minor) There should be an arrow from "Archive" back to "Data" when the data produced is deposited and then contributes back to public repositories? ### Interactions with other ELIXIR Communities * (major) The sentence "Similarly, many of the biodiversity approaches use marker gene amplification for studying environmental DNA (eDNA)" is redundant with the previous one and introduces a different nomenclature that was not used before (marker gene amplication vs metabarcoding) and the differences, if any, are not explained. I would suggest to remove the sentence. * (major) The term "Isolation of genomes" is misleading and I guess the authors used it as a shorthand for "The isolation of bacteria, its DNA extraction, genome sequencing and their annotation". Please rephrase to avoid misinterpretation. Plus I would argue that these steps are also done when deconstructing microbiomes via cultivation strategies and are therefore not so out-of-scope. * (minor) "yet each one of these areas is far greater in scientific scope" feels exaggerated and vague. It should be rephrased. A suggestion is: "yet each one of these areas is too complex to be tacked individually" * (minor) There is no link nor transition between the paragraph that starts with "In summary, " before the Table 3 and the paragraph after that starts with "Similarly, microbiome research". Please rephrase or edit to connect the two sections. Plus, whilst this is good to have a concrete example in the "Similarly" paragraph, the paragraph before was very broad and doing a summary. I would suggest to try reordering the two paragraphs and bring the example earlier for a smoother transition. * (minor) Typo "Similarly, microbiome research has many translation al aspects" * (minor) The PET acronym is detailed but actually used only once, I am unsure if this is necessary. ### Table 2 - (major) The text in the "Interaction" column needs rework as it does not use a consistent wording and could be more to the point, especially as it is a complement to the main text: - The Food and Nutrition entry is a question. - The Galaxy entry has an unnecessary return carriage, and a unspecific "ongoing evaluation study" that needs to be clarified. - The Plant Science entry starts with a generic sentence that could be in the introduction or removed for clarity. Plus I would add a clarification that "plants maintain or not their microbial communities across generations." - (minor) The status (e.g., currently active, inactive, planned, etc.) of each ELIXIR Communities would have been appreciated as it is missing also from the Figure 1. - (minor) The mention of "the field" for the Federated human data community is vague in a manuscript about gathering communities, which research field is implied? If this is the human microbiome research field as a whole, please indicate. ### Table 3 - (major) Similarly to the Table 2 major comment, there is a lack of consistent wording in the "Aim" column that makes the Table 3 not as impactful as it should and could be. Maybe the authors could extract common/distinct features from each of these initiatives as an alternative way to the "Aim" column. A couple of suggestions for these features would be: is it a national initiative?, does it relates to data storage, data analysis, training? is it linked to ELIXIR? - (major) I was surprised not to see the NMDC listed in the table 3, especially when it is discussed in the main text. What about the NCCR Microbiomes initiative in Switzerland? I can understand that some initiatives are not included for space reason, but maybe state it in the legend of the table. - (minor) Is this table sorted? It seems not, but it could be by acronyms or names. - (minor) The aim for the NFDI4Microbiota is way too big a paragraph. The authors should reduce it for conciseness. - (minor) The Metaproteomics Initiative entry has a hyperlink and a reference when none of the others have. Please homogenise. - (minor) Some entries have country listed and some not. Please homogenise. - (minor) I am not questioning the existence of the European Reference Genome Atlas here, but how best to phrase its relevance to ELIXIR in the manuscript. There seems to have no prokaryotes genomes in their atlas, however, there seems to be a trove of fungi and protists genomes which are usually said to be understudied in microbiome. So I think there is a missed opportunity for the authors here to make the most out of this entry. - (minor) "With its headquarters in Bari (Apulia region)," seems irrelevant in the context of the table, please remove. ### Data - (minor) The end of the following sentence is redundant as the INSDC was introduced earlier already "INSDC, which in collaboration with the National Institute of Genetics DNA DataBank of Japan (DDBJ) and the United States National Center for Biotechnology’s (NCBI) GenBank and Sequence Read Archive (SRA), facilitate the deposition and global exchange of sequence data.". Please adjust accordingly. - (minor) "A current challenge facing the field is connecting different multi ‘omics data that have been derived from the same sample." Is this going to be tackled by the ELIXIR Microbiome Community? If so, I would state that this is part of its objectives. - (minor) The end of sentence "[...]overarching context to the experiment, which can be important for meta-analyses." seems like an euphemism, I would suggest to replace with "[...]overarching context to the experiment, re-analyses or meta-analyses." to include re-analyses as well. - (minor) "We will continue to promote such approaches, enriching metadata wherever possible." Is this going to be done via the CDCH? - (minor) "The ELIXIR Microbiome Community will also work to move the Marine Metagenomics domain in the RDMKit towards a more general Microbiome domain." What is the RDMKit? It is not explained, nor cited nor mentioned again. ### Tools - (minor) "will increase their use of BioContainers" should be "will increase the use of BioContainers" - (minor) "In order to make tools findable by the end users, the Community" - (minor) "workflow descriptions (e.g. Snakemake, CWL, Nextflow)" None of them have their references cited, is it an omission or space limitation? - (minor) "A current joint effort between the Microbiome and Galaxy Communities" ### Benchmarking - (major) "Benchmarking" it is the only item at this hierarchy level, meaning that this is useless for structuring the text. Please edit. - (minor) Review reference 2 published a recent review with guidelines to learn from that could have its place in this paragraph. - (minor) in the sentence: "As the Microbiome Community establishes, we will develop a broader understanding of the requirements of the Community , feed this to the Tools Platform, as well as seek opportunities to interact with the Tools Platform to capture the diversity of tools and their utility via such benchmarking activities.", is this the ELIXIR Microbiome Community, or the wide community of microbiome researchers? Is it to mean that the ELIXIR Microbiome Community is going to act as an interface between microbiome researchers and ELIXIR Tools/Infrastructure? ### Compute - (minor) The first sentence would fit better in the introduction. ### Interoperability - (minor) The reference 57 should be at the end of the sentence starting with "This effort was paralleled" not in the middle. - (minor) Reference 58 should be removed at it is a duplicate of reference 47. - (minor) Use the full text "Global Alliance for Genomics and Health" instead of GA4GH. - (minor) "is in the process of applying to be a n ELIXIR Recommended Interoperability Resource." - (minor) In the sentence: "This will require the development of new data Interoperability layers for data resources that are not normally focused in Microbiome data" I think the authors meant "used" instead of "focused", and "microbiome" instead of "Microbiome". ### Training - (minor) In "Platforms such as MGnify support large-scale services for most, if not all, steps of a microbiome study", I would suggest to remove the "if not all". - (minor) "with areas of expertise covering different environments, ‘omics approaches and data analysis pathways." I think the authors meant "learning paths", and I would refrain from using "pathways" as it also has a biological meaning. ### Context with other international initiatives - (major) Whilst I appreciated the emphasis that no initiative exists on its own, I feel the first paragraph on the Genom ic (please correct the typo) Standards Consortium feels lengthy for a manuscript whose topic is not the GSC. I would advise to summarize. In this respect, the second paragraph is particularly relevant to a tangible collaboration between ELIXIR Microbiome Community and GSC. - (major) The NMDC is discussed in this section but not part of the Table 3. - (minor) Is the mentioned M5 project still active as the website's last update is 2012? - (minor) "Combining the activities on standards concerning workflows [...]" does this means adding and providing workflows to the microbiome research community? - (minor) Given the emphasis on the fact that MicrobiomeSupport was a program, the authors could update the readers and indicate that it is now MicrobiomeSupport Association. ### Interaction with other key data resources beyond ELIXIR - (major) There is an order issue with the main text that a proofread could solve, as MG-RAST is explained and cited in the first paragraph but already mentioned upstream of the main text in the Interoperability section. ### Specific challenges and objectives of the ELIXIR Microbiome Community - (major) The strong claim "it is widely accepted that current short-read assembly-based methods do not generally work as well for soil microbiomes" would probably need at least one reference. - (major) "(iii) there is no centralised database collecting the millions of viral sequences". It seems to be the case indeed, and there are databases (~24) out there as recently compiled in Review reference 3. How ELIXIR Microbiome Community plans to integrate/aggregate these resources in a non-duplicating manner? - (minor) The word "through" is superfluous in the the sentence that starts with "This current limitation, [...]" and can be removed. - (minor) The CAMI was already explained and cited above, so the already defined acronym can be used. - (minor) Precise the area in :"Additionally, another key area of development of taxonomy [...]". - (minor) The sentence "Viruses, particularly those that infect bacteria, are found ubiquitously in all environments and play critical roles in community dynamics." belongs in an introduction, not so downstream of the manuscript. - (minor) The sentence starting the sixth paragraph could be precised as "The increase in metagenomic assemblies has resulted in a parallel increase in the number of predicted protein sequences, with sets of non-redundant proteins now in the billions." - (minor) Fix typo in "that are undetectable by current sequence based methods." - (minor) Would it make sense to also be able to access representative, of clusters for instance? "[...] develop new infrastructural frameworks for accessing slices of the data or adequate representatives based on the requirements." - (minor) What is the " expanded Microbiome Community"? It was never mentioned before. - (minor) A few comments on the sentence: "Within the Community, we will develop and promote standards around the analysis provenance (analytical metadata),". In my opinion and how it was already stated in the manuscript, it would make more sense to promote existing standards first and then develop if need be. There is no mention of other type of metadata, so the "analytical metadata" precision seems superfluous. I would suggest: "Within the Community, we will promote and develop standards regarding the analysis provenance," - (minor) Precise the term forms in "[...] ensuring that comput ing resources are accessible for performing the different forms of data analysis [...]", do the authors mean types of/steps in the data analysis? - (minor) Same argument as before regarding reinventing the wheel, I would swap the part of the sentence: "This may require the extensions to existing databases or development of new ones , but it requires an agreement from the research community to adopt them." - (minor) This part "Metaproteomics aims to elucidate the functional and taxonomic interplay of proteins in microbiomes," should have been in the Table 1, or to reuse the Table 1 here. - (minor) The tenth paragraph of this section starts with the mention of multiple major challenges, but detail "only" one of them. I would suggest to mention some of the others challenges. - (minor) The "MIA" method is just a hyperlink, without any reference. Either cite the website accordingly or add the reference. ### Table 4 - (major) I was surprised to see that the objective "Foster international collaborations between other resources providers and databases to ensure global harmonisation of e-infrastructures for microbiome research" was long-term, as I would have imagined that a gap analysis would be short-term to ensure we do not reinvent the wheel, especially given the others initiatives discussed in the manuscript. - (minor) If the Objective column starts with action verbs (which is a good idea), then it should be "Survey the needs" instead of "Survey of needs". - (minor) Specify the type of workflow with "Address knowledge gaps in generating and adopting data analysis workflows" - (minor) The verb is missing in " Teach a dvanced containerisation and cloud deployment" - (minor) The verb is missing in " Promote data analysis through the use of services". Is this ELIXIR services in general or specific ELIXIR Microbiome Community services? - (minor) Use "Share" instead of "Sharing" in the "Co-ordinate" entry. - (minor) The "Industry connection" entry does not fit the action verb pattern. A suggestion would be "Use ELIXIR and Node forums to understand pharmaceutical and biotechnological demands and current limitations impacting this sector." as the first sentence felt generic. - (minor) The verb is missing in " Design targeted training for different microbiome communities" - (minor) Use "findability" instead of "discoverability" for consistency and to fit with the FAIR principles. - (minor) Reorder the sentence to start with the action verb: " E stablish new standards for microbiome research, particularly with respect to data analysis reporting and contextual metadata reporting in conjunction with GSC " - (minor) Correct "Established" to "Establish" in the "Promoting new approaches" entry. Editorial comments: - (major) "Community" is used in upper-case and this is unclear in many instances whether the ELIXIR Marine Metagenomics Community is referred to, the ELIXIR Microbiome Community, or the broader microbiome research community. I suggest to use consistently the full term for the sake of transparency. Abbreviations like EMMC and EMC could be even more misleading in my opinion. - (major) The structure of the white paper is not evident as the hierarchy is indicated only by change in font size, and some sections are quite lengthy for a white paper that is supposed to be concise. I understand that this is a constraint from the Editor, but see Review reference 4 for a white paper with a more clearer structure. An alternative could be to use numbered sections. - References - (minor) The very first reference is oddly formatted in the text creating an artificial and confusing end of sentence. - (minor) Title of reference 9 is truncated and should be "Methods included: standardizing computational reuse and portability with the Common Workflow Language" - (minor) The superscript numbers of the numeric style of bibliography are in many instances after the final dot (see reference 2, 12-16, 17-19, 36, 37), after a comma (see reference 10, 20, 23), a bracket (see reference 60) or semi colon (see reference 65) when they should be before any of these symbols. - (minor) In the paragraph "Interactions with other ELIXIR Communities", we jump from reference 24 to 29 when the numeric style of bibliography (that was chosen by the authors) is expected to mirror the mentions in the manuscript. Please either have the Table 2 earlier in the paper, or change the order of the references. Is the topic of the opinion article discussed accurately in the context of the current literature? Yes Are all factual statements correct and adequately supported by citations? Partly Are arguments sufficiently supported by evidence from the published literature? Partly Are the conclusions drawn balanced and justified on the basis of the presented arguments? Partly References 1. Berg G, Rybakova D, Fischer D, Cernava T, et al.: Microbiome definition re-visited: old concepts and new challenges. Microbiome . 2020; 8 (1). Publisher Full Text 2. Ritsch M, Cassman NA, Saghaei S, Marz M: Navigating the Landscape: A Comprehensive Review of Current Virus Databases. Viruses . 2023; 15 (9). PubMed Abstract | Publisher Full Text 3. Brooks TG, Lahens NF, Mrčela A, Grant GR: Challenges and best practices in omics benchmarking. Nat Rev Genet . 2024. PubMed Abstract | Publisher Full Text 4. Vizcaíno JA, Walzer M, Jiménez RC, Bittremieux W, et al.: A community proposal to integrate proteomics activities in ELIXIR. F1000Res . 2017; 6 . PubMed Abstract | Publisher Full Text Competing Interests No competing interests were disclosed. Reviewer Expertise microbiome, bioinformatics, reproducible research, data and metadata standards I confirm that I have read this submission and believe that I have an appropriate level of expertise to confirm that it is of an acceptable scientific standard, however I have significant reservations, as outlined above. reply Respond to this report Responses (1) Author Response 10 Sep 2025 Bérénice Batut, Bioinformatics Group, Department of Computer Science, Albert-Ludwigs-University Freiburg, Freiburg, Germany The authors make the case for mutating the ELIXIR Marine Metagenomics Community initiative into an ELIXIR Microbiome Community. It is indeed timely to have an such a pan-European initiative especially as it draws on previous experience, albeit a more narrow biome. The emphasis on multi-omics integration is also a strong suit as it is a complex topic where the microbiome research community would need support and infrastructure. The authors map out existing (but not all) similar initiatives, even if it is unclear how they would work together. In my opinion, a couple of claims should be strengthened and the argumentative power of this white paper would benefit from minor reorganization, a text slim down and extra proofreading. To this end, I highlight the strengths and weaknesses of the manuscript below. ## Strengths - Table 3 is a great idea to map out the others initiatives but needs a bit rework to be straight to the point. - In the Data section, the part between "Throughout the lifetime of the ELIXIR Marine Metagenomics Community" and the end of the paragraph is really relevant and draws from the experience of the ELIXIR Community. These sentences should be more highlighted, maybe upstream. I do wonder how the authors considered how the others initiatives (mentioned in Table 3) could also contribute to the addition of metadata for a global curation effort, possibly via the mentioned Contextual Data Clearinghouse. - The data re-analysis and integration between PRIDE and MGnify (named MetaPUF) is a good use case of the efforts that the new ELIXIR Community can promote. - I did appreciate the clear objective at the end of the Compute section, to help with scaling up analyses as well as the plan for the interaction with the ELIXIR Training Platform which promise to give a boost in training the current and next generation of microbiome researchers. - I liked how the authors highlighted (using the example of viral sequences databases) that time and resources should not be wasted in duplicating works, and that initiatives such as the ELIXIR Microbiome Community can promote this. - The authors thought ahead to promote and develop methods to limit confounding factors when re-using datasets, especially in multi-omics settings, and added this to one of the ELIXIR Microbiome Community challenges. ## Weaknesses ### Introduction - (minor) The microbiome definition used in the paper was revisited in Review reference 1 to include the interactions as well as others molecules than genomes as part of the microbiome. I suggest to use this definition given the emphasis on additional omics made in the paper. We have added that the definition includes other molecules in addition to genomes. - (major) The first paragraph ends on challenges of our research community including how to make data FAIR. Given the experience gained by the ELIXIR Marine Metagenomics Community, I would have appreciated a sentence on why making our data FAIR is important (e.g., transparency to make the research more reproducible, accountability given the (public) source of funding). See major comment in next section. Thank you for this suggestion. We have added a few sentences to this effect in the manuscript. ### The scope of the ELIXIR Microbiome Community - (major) These following arguments for FAIR data should be in the Introduction section. "Moreover, when wishing to contextualise the results with similar experiments, the way a dataset has been produced and processed must be transparent to establish whether it is comparable (e.g. amplified sequence variants can only be compared when the same amplified regions are compared). Furthermore, when different methods are applied, best practices in data stewardship are required to ensure that the connectivity of the derived sequence data products, together with functional and taxonomic assertions are kept in context of the original sample/sequencing effort and associated contextual metadata." We have merged this part of the paragraph into the Introduction. - (minor) The very first sentence of this section blames researchers for misuse of a term "[...] regularly (incorrectly) used [...]". I suggest to rephrase the sentence to simply state the differences, or back up the misuse of the term with references or a survey. We have qualified this with an example of INSDC mislabelling. - (minor) "Fundamentally, the ELIXIR Microbiome Community is about providing the necessary infrastructures required to perform analysis of nucleotide sequence data derived from a microbiome" “Nucleotide” has been added to this sentence. - (minor) "Finally, and possible unique to this ELIXIR Community, is the variety of researchers". I am not sure that these features are unique to the ELIXIR Community but rather to the field of microbiome research. I suggest to tone it down by simply pointing out that the ELIXIR Community gathers a variety of researchers. Modified accordingly. ### Table 1 - (major) Metabolomics is listed in Table 1 but is omitted from the narrative in the paragraph. We have added a sentence in the paragraph. - (major) The Table 1 is not really used but could actually reduce some redundancies in the manuscript by providing a one-stop-shop to explain and detail these techniques. We have increased the cross linking to the table. ### Figure 1 - (major) The figure is early on in the manuscript and it is unclear which of the ELIXIR platforms and communities are already established, or foreseen. Especially since "current and future ELIXIR activities" is mentioned in the manuscript before referencing this figure. The Table 2 does not add more information in that regard. - (major) There is a need for different types of arrows as the same arrow represent interactions between communities or processes. Modified accordingly. - (minor) The "Data" node is quite generic, I was wondering whether it meant public repositories, institute repositories or both. Please be precise. "Data" refers here to the ELIXIR Data Platform - (minor) I guess the type of metadata illustrated in the figure is restricted to biological metadata, that is indeed collected during sampling, however, technical metadata such as the sequencing method used, the type of instrument or the library layout, are collected during the processing of the sample not only the sampling itself. This is everything from phenotypic, sample conditions, experimental methods - (minor) There should be an arrow from "Archive" back to "Data" when the data produced is deposited and then contributes back to public repositories? ### Interactions with other ELIXIR Communities * (major) The sentence "Similarly, many of the biodiversity approaches use marker gene amplification for studying environmental DNA (eDNA)" is redundant with the previous one and introduces a different nomenclature that was not used before (marker gene amplication vs metabarcoding) and the differences, if any, are not explained. I would suggest to remove the sentence. We have modified the sentence to refer to barcoding rather than removing the sentence. We feel it is important to mention eDNA. * (major) The term "Isolation of genomes" is misleading and I guess the authors used it as a shorthand for "The isolation of bacteria, its DNA extraction, genome sequencing and their annotation". Please rephrase to avoid misinterpretation. Plus I would argue that these steps are also done when deconstructing microbiomes via cultivation strategies and are therefore not so out-of-scope. This sentence has been removed as the reviewer is correct that cultivation of bacteria (and other organisms) from microbiomes is becoming an increasing trend. * (minor) "yet each one of these areas is far greater in scientific scope" feels exaggerated and vague. It should be rephrased. A suggestion is: "yet each one of these areas is too complex to be tacked individually" We have amended the sentence accordingly. * (minor) There is no link nor transition between the paragraph that starts with "In summary, " before the Table 3 and the paragraph after that starts with "Similarly, microbiome research". Please rephrase or edit to connect the two sections. Plus, whilst this is good to have a concrete example in the "Similarly" paragraph, the paragraph before was very broad and doing a summary. I would suggest to try reordering the two paragraphs and bring the example earlier for a smoother transition. The paragraphs have been reordered as suggested and slightly amended to improve readability. * (minor) Typo "Similarly, microbiome research has many translation al aspects" Fixed * (minor) The PET acronym is detailed but actually used only once, I am unsure if this is necessary. There is another instance of this abbreviation, so we have retained this abbreviation. ### Table 2 - (major) The text in the "Interaction" column needs rework as it does not use a consistent wording and could be more to the point, especially as it is a complement to the main text: - The Food and Nutrition entry is a question. - The Galaxy entry has an unnecessary return carriage, and a unspecific "ongoing evaluation study" that needs to be clarified. - The Plant Science entry starts with a generic sentence that could be in the introduction or removed for clarity. Plus I would add a clarification that "plants maintain or not their microbial communities across generations." - (minor) The status (e.g., currently active, inactive, planned, etc.) of each ELIXIR Communities would have been appreciated as it is missing also from the Figure 1. - (minor) The mention of "the field" for the Federated human data community is vague in a manuscript about gathering communities, which research field is implied? If this is the human microbiome research field as a whole, please indicate. In response to all of the above comments, we have substantially reworked table 2. The only comment that we have not addressed so specifically, is whether an activity is in progress or not, because the activities ebb and flow, and it is not always easy to say when there is a start or an end. However, being somewhat more focused in the content, we feel this revised table provides a stronger view of the direction that the community wants to take with the other communities. ### Table 3 - (major) Similarly to the Table 2 major comment, there is a lack of consistent wording in the "Aim" column that makes the Table 3 not as impactful as it should and could be. Maybe the authors could extract common/distinct features from each of these initiatives as an alternative way to the "Aim" column. A couple of suggestions for these features would be: is it a national initiative?, does it relates to data storage, data analysis, training? is it linked to ELIXIR? - (major) I was surprised not to see the NMDC listed in the table 3, especially when it is discussed in the main text. What about the NCCR Microbiomes initiative in Switzerland? I can understand that some initiatives are not included for space reason, but maybe state it in the legend of the table. - (minor) Is this table sorted? It seems not, but it could be by acronyms or names. - (minor) The aim for the NFDI4Microbiota is way too big a paragraph. The authors should reduce it for conciseness. - (minor) The Metaproteomics Initiative entry has a hyperlink and a reference when none of the others have. Please homogenise. - (minor) Some entries have country listed and some not. Please homogenise. - (minor) I am not questioning the existence of the European Reference Genome Atlas here, but how best to phrase its relevance to ELIXIR in the manuscript. There seems to have no prokaryotes genomes in their atlas, however, there seems to be a trove of fungi and protists genomes which are usually said to be understudied in microbiome. So I think there is a missed opportunity for the authors here to make the most out of this entry. - (minor) "With its headquarters in Bari (Apulia region)," seems irrelevant in the context of the table, please remove. As with Table 2, we have substantially reworked Table 3, including ordering by the reach of the effort, harmonising the language and reducing text. We have also included the relevance of the activity to the community. We have not included NMDC, as this is a US effort, and while important to the community, this table is focused on Europe efforts. ### Data - (minor) The end of the following sentence is redundant as the INSDC was introduced earlier already "INSDC, which in collaboration with the National Institute of Genetics DNA DataBank of Japan (DDBJ) and the United States National Center for Biotechnology’s (NCBI) GenBank and Sequence Read Archive (SRA), facilitate the deposition and global exchange of sequence data.". Please adjust accordingly. We have removed the redundancy here. - (minor) "A current challenge facing the field is connecting different multi ‘omics data that have been derived from the same sample." Is this going to be tackled by the ELIXIR Microbiome Community? If so, I would state that this is part of its objectives. We feel this is likely to be a joint effort across different communities. Thus, we have retained the sentence as is. - (minor) The end of sentence "[...]overarching context to the experiment, which can be important for meta-analyses." seems like an euphemism, I would suggest to replace with "[...]overarching context to the experiment, re-analyses or meta-analyses." to include re-analyses as well. We have added “re-analyses or” as suggested. - (minor) "We will continue to promote such approaches, enriching metadata wherever possible." Is this going to be done via the CDCH? While CDCH offers one approach, we believe that there are many different ways that metadata may be enriched, first by making scientists more aware of the need of submitting meta, promoting different, more specific checklists (e.g. STORMS or MicroB3) and mining metadata from published literature. This, specifically highlighting CDCH would not be appropriate. - (minor) "The ELIXIR Microbiome Community will also work to move the Marine Metagenomics domain in the RDMKit towards a more general Microbiome domain." What is the RDMKit? It is not explained, nor cited nor mentioned again. This has been expanded to the ELIXIR Research Data Management Kit and a link has been added to the text. ### Tools - (minor) "will increase their use of BioContainers" should be "will increase the use of BioContainers" Corrected. - (minor) "In order to make tools findable by the end users, the Community" Done. - (minor) "workflow descriptions (e.g. Snakemake, CWL, Nextflow)" None of them have their references cited, is it an omission or space limitation? Added references - (minor) "A current joint effort between the Microbiome and Galaxy Communities" Fixed ### Benchmarking - (major) "Benchmarking" it is the only item at this hierarchy level, meaning that this is useless for structuring the text. Please edit. Thank you for pointing this out. We have changed the sub-sub-heading into an introductory sentence to this paragraph. - (minor) Review reference 2 published a recent review with guidelines to learn from that could have its place in this paragraph. - (minor) in the sentence: "As the Microbiome Community establishes, we will develop a broader understanding of the requirements of the Community , feed this to the Tools Platform, as well as seek opportunities to interact with the Tools Platform to capture the diversity of tools and their utility via such benchmarking activities.", is this the ELIXIR Microbiome Community, or the wide community of microbiome researchers? Is it to mean that the ELIXIR Microbiome Community is going to act as an interface between microbiome researchers and ELIXIR Tools/Infrastructure? We clarified this to mean the wider community of microbiome researchers. ### Compute - (minor) The first sentence would fit better in the introduction. ### Interoperability - (minor) The reference 57 should be at the end of the sentence starting with "This effort was paralleled" not in the middle. This citation has been moved to the end of the sentence. - (minor) Reference 58 should be removed at it is a duplicate of reference 47. References have been fixed - (minor) Use the full text "Global Alliance for Genomics and Health" instead of GA4GH. This abbreviation has been expanded. - (minor) "is in the process of applying to be a n ELIXIR Recommended Interoperability Resource." ELIXIR has been added to this sentence. - (minor) In the sentence: "This will require the development of new data Interoperability layers for data resources that are not normally focused in Microbiome data" I think the authors meant "used" instead of "focused", and "microbiome" instead of "Microbiome". We have updated the sentence accordingly. ### Training - (minor) In "Platforms such as MGnify support large-scale services for most, if not all, steps of a microbiome study", I would suggest to remove the "if not all". Agreed. - (minor) "with areas of expertise covering different environments, ‘omics approaches and data analysis pathways." I think the authors meant "learning paths", and I would refrain from using "pathways" as it also has a biological meaning. Actually, we do mean data analysis, but have changed to data analysis strategies. This is in reference to the fact that there may be multiple different ways of analysing a data type, for example metagenomics can be analysed using tools such as Kraken or MetaPhlAn4, to full assembly, gene calling and functional analysis. ### Context with other international initiatives - (major) Whilst I appreciated the emphasis that no initiative exists on its own, I feel the first paragraph on the Genom ic (please correct the typo) Standards Consortium feels lengthy for a manuscript whose topic is not the GSC. I would advise to summarize. In this respect, the second paragraph is particularly relevant to a tangible collaboration between ELIXIR Microbiome Community and GSC. The typo has been corrected. - (major) The NMDC is discussed in this section but not part of the Table 3. Table 3 lists pan-European efforts, whereas the NMDC is a US specific initiative. Thus, we have not added this to the table. - (minor) Is the mentioned M5 project still active as the website's last update is 2012? While the website has not been updated, there are ongoing efforts to expand and revitalise this effort. RDF is a member of the GSC board and is promoting aspects of the M5 initiative. While the name may change, we feel this is useful to keep. However, as part of condensing the overall length of the manuscript, we have removed the reference to this initiative. - (minor) "Combining the activities on standards concerning workflows [...]" does this means adding and providing workflows to the microbiome research community? - (minor) Given the emphasis on the fact that MicrobiomeSupport was a program, the authors could update the readers and indicate that it is now MicrobiomeSupport Association. ### Interaction with other key data resources beyond ELIXIR - (major) There is an order issue with the main text that a proofread could solve, as MG-RAST is explained and cited in the first paragraph but already mentioned upstream of the main text in the Interoperability section. This has now been fixed. ### Specific challenges and objectives of the ELIXIR Microbiome Community - (major) The strong claim "it is widely accepted that current short-read assembly-based methods do not generally work as well for soil microbiomes" would probably need at least one reference. Added - (major) "(iii) there is no centralised database collecting the millions of viral sequences". It seems to be the case indeed, and there are databases (~24) out there as recently compiled in Review reference 3. How ELIXIR Microbiome Community plans to integrate/aggregate these resources in a non-duplicating manner? - (minor) The word "through" is superfluous in the the sentence that starts with "This current limitation, [...]" and can be removed. We have removed the sentence “through”. - (minor) The CAMI was already explained and cited above, so the already defined acronym can be used. This has been fixed. - (minor) Precise the area in :"Additionally, another key area of development of taxonomy [...]". Added - (minor) The sentence "Viruses, particularly those that infect bacteria, are found ubiquitously in all environments and play critical roles in community dynamics." belongs in an introduction, not so downstream of the manuscript. - (minor) The sentence starting the sixth paragraph could be precised as "The increase in metagenomic assemblies has resulted in a parallel increase in the number of predicted protein sequences, with sets of non-redundant proteins now in the billions." Added “predicted” to the sentence. - (minor) Fix typo in "that are undetectable by current sequence based methods." Fixed - sequenced -> sequence - (minor) Would it make sense to also be able to access representative, of clusters for instance? "[...] develop new infrastructural frameworks for accessing slices of the data or adequate representatives based on the requirements." Added - (minor) What is the " expanded Microbiome Community"? It was never mentioned before. Removed “expanded” from the sentence. - (minor) A few comments on the sentence: "Within the Community, we will develop and promote standards around the analysis provenance (analytical metadata),". In my opinion and how it was already stated in the manuscript, it would make more sense to promote existing standards first and then develop if need be. There is no mention of other type of metadata, so the "analytical metadata" precision seems superfluous. I would suggest: "Within the Community, we will promote and develop standards regarding the analysis provenance," We agree with the reviewer's comment and have re-ordered accordingly. - (minor) Precise the term forms in "[...] ensuring that comput ing resources are accessible for performing the different forms of data analysis [...]", do the authors mean types of/steps in the data analysis? “compute resources” is an accepted resource. We have removed the “different forms of” as we feel it is a little superfluous. We were meaning the different strategies, but it does not add to the sentence. - (minor) Same argument as before regarding reinventing the wheel, I would swap the part of the sentence: "This may require the extensions to existing databases or development of new ones , but it requires an agreement from the research community to adopt them." We agree, and have amended accordingly. - (minor) This part "Metaproteomics aims to elucidate the functional and taxonomic interplay of proteins in microbiomes," should have been in the Table 1, or to reuse the Table 1 here. We have highlighted this in table 1. We have kept the text here to provide the context for the rest of the sentence. - (minor) The tenth paragraph of this section starts with the mention of multiple major challenges, but detail "only" one of them. I would suggest to mention some of the others challenges. - (minor) The "MIA" method is just a hyperlink, without any reference. Either cite the website accordingly or add the reference. There is not a reference, and the hyperlink goes to the GitHub site, as requested by the authors. We have included “ https://github.com/microbiome/mia ” for completeness sake. ### Table 4 - (major) I was surprised to see that the objective "Foster international collaborations between other resources providers and databases to ensure global harmonisation of e-infrastructures for microbiome research" was long-term, as I would have imagined that a gap analysis would be short-term to ensure we do not reinvent the wheel, especially given the others initiatives discussed in the manuscript. - (minor) If the Objective column starts with action verbs (which is a good idea), then it should be "Survey the needs" instead of "Survey of needs". We have amended accordingly. - (minor) Specify the type of workflow with "Address knowledge gaps in generating and adopting data analysis workflows" Added - (minor) The verb is missing in " Teach a dvanced containerisation and cloud deployment" Added. - (minor) The verb is missing in " Promote data analysis through the use of services". Is this ELIXIR services in general or specific ELIXIR Microbiome Community services? Added Promote - generalised to ELIXIR services. - (minor) Use "Share" instead of "Sharing" in the "Co-ordinate" entry. Done - (minor) The "Industry connection" entry does not fit the action verb pattern. A suggestion would be "Use ELIXIR and Node forums to understand pharmaceutical and biotechnological demands and current limitations impacting this sector." as the first sentence felt generic. Done - (minor) The verb is missing in " Design targeted training for different microbiome communities" Done - (minor) Use "findability" instead of "discoverability" for consistency and to fit with the FAIR principles. Done . - (minor) Reorder the sentence to start with the action verb: " E stablish new standards for microbiome research, particularly with respect to data analysis reporting and contextual metadata reporting in conjunction with GSC " Done - (minor) Correct "Established" to "Establish" in the "Promoting new approaches" entry. Done Editorial comments: - (major) "Community" is used in upper-case and this is unclear in many instances whether the ELIXIR Marine Metagenomics Community is referred to, the ELIXIR Microbiome Community, or the broader microbiome research community. I suggest to use consistently the full term for the sake of transparency. Abbreviations like EMMC and EMC could be even more misleading in my opinion. We have prefixed all instances of “Community” with their explicit community name. - (major) The structure of the white paper is not evident as the hierarchy is indicated only by change in font size, and some sections are quite lengthy for a white paper that is supposed to be concise. I understand that this is a constraint from the Editor, but see Review reference 4 for a white paper with a more clearer structure. An alternative could be to use numbered sections. We have introduced section numbers to aid the structuring of the manuscript. - References - (minor) The very first reference is oddly formatted in the text creating an artificial and confusing end of sentence. - (minor) Title of reference 9 is truncated and should be "Methods included: standardizing computational reuse and portability with the Common Workflow Language" - (minor) The superscript numbers of the numeric style of bibliography are in many instances after the final dot (see reference 2, 12-16, 17-19, 36, 37), after a comma (see reference 10, 20, 23), a bracket (see reference 60) or semi colon (see reference 65) when they should be before any of these symbols. - (minor) In the paragraph "Interactions with other ELIXIR Communities", we jump from reference 24 to 29 when the numeric style of bibliography (that was chosen by the authors) is expected to mirror the mentions in the manuscript. Please either have the Table 2 earlier in the paper, or change the order of the references. References have been fixed View more View less Competing Interests No competing interests were disclosed. reply Respond Report a concern Pauvert C. Peer Review Report For: Establishing the ELIXIR Microbiome Community [version 1; peer review: 1 approved, 1 approved with reservations] . F1000Research 2024, 13 (ELIXIR):50 ( https://doi.org/10.5256/f1000research.158321.r244855) NOTE: it is important to ensure the information in square brackets after the title is included in this citation. The direct URL for this report is: https://f1000research.com/articles/13-50/v1#referee-response-244855 Alongside their report, reviewers assign a status to the article: Approved - the paper is scientifically sound in its current form and only minor, if any, improvements are suggested Approved with reservations - A number of small changes, sometimes more significant revisions are required to address specific details and improve the papers academic merit. Not approved - fundamental flaws in the paper seriously undermine the findings and conclusions Adjust parameters to alter display View on desktop for interactive features Includes Interactive Elements View on desktop for interactive features Competing Interests Policy Provide sufficient details of any financial or non-financial competing interests to enable users to assess whether your comments might lead a reasonable person to question your impartiality. Consider the following examples, but note that this is not an exhaustive list: Examples of 'Non-Financial Competing Interests' Within the past 4 years, you have held joint grants, published or collaborated with any of the authors of the selected paper. You have a close personal relationship (e.g. parent, spouse, sibling, or domestic partner) with any of the authors. You are a close professional associate of any of the authors (e.g. scientific mentor, recent student). You work at the same institute as any of the authors. You hope/expect to benefit (e.g. favour or employment) as a result of your submission. You are an Editor for the journal in which the article is published. Examples of 'Financial Competing Interests' You expect to receive, or in the past 4 years have received, any of the following from any commercial organisation that may gain financially from your submission: a salary, fees, funding, reimbursements. You expect to receive, or in the past 4 years have received, shared grant support or other funding with any of the authors. You hold, or are currently applying for, any patents or significant stocks/shares relating to the subject matter of the paper you are commenting on. Stay Updated Sign up for content alerts and receive a weekly or monthly email with all newly published articles Register with F1000Research Already registered? Sign in Not now, thanks close PLEASE NOTE If you are an AUTHOR of this article, please check that you signed in with the account associated with this article otherwise we cannot automatically identify your role as an author and your comment will be labelled as a “User Comment”. If you are a REVIEWER of this article, please check that you have signed in with the account associated with this article and then go to your account to submit your report, please do not post your review here. If you do not have access to your original account, please contact us . All commenters must hold a formal affiliation as per our Policies . The information that you give us will be displayed next to your comment. User comments must be in English, comprehensible and relevant to the article under discussion. We reserve the right to remove any comments that we consider to be inappropriate, offensive or otherwise in breach of the User Comment Terms and Conditions . Commenters must not use a comment for personal attacks. When criticisms of the article are based on unpublished data, the data should be made available. I accept the User Comment Terms and Conditions Please confirm that you accept the User Comment Terms and Conditions. Affiliation ✕ refresh Please enter your institution. Note: To add your institution or organisation, start typing the name and then select the correct name from the list. Where applicable, the name will appear in both the original language and in English. Do not paste in the name. If the name does not appear in the drop-down list, we will display the information you have entered. ✕ refresh Country/Region * USA UK Canada China France Germany Afghanistan Aland Islands Albania Algeria American Samoa Andorra Angola Anguilla Antarctica Antigua and Barbuda Argentina Armenia Aruba Australia Austria Azerbaijan Bahamas Bahrain Bangladesh Barbados Belarus Belgium Belize Benin Bermuda Bhutan Bolivia Bosnia and Herzegovina Botswana Bouvet Island Brazil British Indian Ocean Territory British Virgin Islands Brunei Bulgaria Burkina Faso Burundi Cambodia Cameroon Canada Cape Verde Cayman Islands Central African Republic Chad Chile China Christmas Island Cocos (Keeling) Islands Colombia Comoros Congo Cook Islands Costa Rica Cote d'Ivoire Croatia Cuba Cyprus Czech Republic Democratic Republic of the Congo Denmark Djibouti Dominica Dominican Republic Ecuador Egypt El Salvador Equatorial Guinea Eritrea Estonia Ethiopia Falkland Islands Faroe Islands Federated States of Micronesia Fiji Finland France French Guiana French Polynesia French Southern Territories Gabon Georgia Germany Ghana Gibraltar Greece Greenland Grenada Guadeloupe Guam Guatemala Guernsey Guinea Guinea-Bissau Guyana Haiti Heard Island and Mcdonald Islands Holy See (Vatican City State) Honduras Hong Kong Hungary Iceland India Indonesia Iran Iraq Ireland Israel Italy Jamaica Japan Jersey Jordan Kazakhstan Kenya Kiribati Kosovo (Serbia and Montenegro) Kuwait Kyrgyzstan Lao People's Democratic Republic Latvia Lebanon Lesotho Liberia Libya Liechtenstein Lithuania Luxembourg Macao Madagascar Malawi Malaysia Maldives Mali Malta Marshall Islands Martinique Mauritania Mauritius Mayotte Mexico Minor Outlying Islands of the United States Moldova Monaco Mongolia Montenegro Montserrat Morocco Mozambique Myanmar Namibia Nauru Nepal Netherlands Antilles New Caledonia New Zealand Nicaragua Niger Nigeria Niue Norfolk Island North Korea North Macedonia Northern Mariana Islands Norway Oman Pakistan Palau Palestinian Territory Panama Papua New Guinea Paraguay Peru Philippines Pitcairn Poland Portugal Puerto Rico Qatar Reunion Romania Russian Federation Rwanda Saint Helena Saint Kitts and Nevis Saint Lucia Saint Pierre and Miquelon Saint Vincent and the Grenadines Samoa San Marino Sao Tome and Principe Saudi Arabia Senegal Serbia Seychelles Sierra Leone Singapore Slovakia Slovenia Solomon Islands Somalia South Africa South Georgia and the South Sandwich Is South Korea South Sudan Spain Sri Lanka Sudan Suriname Svalbard and Jan Mayen Swaziland Sweden Switzerland Syria Taiwan Tajikistan Tanzania Thailand The Gambia The Netherlands Timor-Leste Togo Tokelau Tonga Trinidad and Tobago Tunisia Turkey Turkmenistan Turks and Caicos Islands Tuvalu UK USA Uganda Ukraine United Arab Emirates United States Virgin Islands Uruguay Uzbekistan Vanuatu Venezuela Vietnam Wallis and Futuna West Bank and Gaza Strip Western Sahara Yemen Zambia Zimbabwe Please select your country/region. You must enter a comment. Competing Interests Please disclose any competing interests that might be construed to influence your judgment of the article's or peer review report's validity or importance. Competing Interests Policy Provide sufficient details of any financial or non-financial competing interests to enable users to assess whether your comments might lead a reasonable person to question your impartiality. Consider the following examples, but note that this is not an exhaustive list: Examples of 'Non-Financial Competing Interests' Within the past 4 years, you have held joint grants, published or collaborated with any of the authors of the selected paper. You have a close personal relationship (e.g. parent, spouse, sibling, or domestic partner) with any of the authors. You are a close professional associate of any of the authors (e.g. scientific mentor, recent student). You work at the same institute as any of the authors. You hope/expect to benefit (e.g. favour or employment) as a result of your submission. You are an Editor for the journal in which the article is published. Examples of 'Financial Competing Interests' You expect to receive, or in the past 4 years have received, any of the following from any commercial organisation that may gain financially from your submission: a salary, fees, funding, reimbursements. You expect to receive, or in the past 4 years have received, shared grant support or other funding with any of the authors. You hold, or are currently applying for, any patents or significant stocks/shares relating to the subject matter of the paper you are commenting on. Please state your competing interests The comment has been saved. An error has occurred. Please try again. Cancel Post var lTitle = "Establishing the ELIXIR Microbiome Community".replace("'", ''); var linkedInUrl = "http://www.linkedin.com/shareArticle?url=https://f1000research.com/articles/13-50/v1" + "&title=" + encodeURIComponent(lTitle) + "&summary=" + encodeURIComponent('Read the article by '); var deliciousUrl = "https://del.icio.us/post?url=https://f1000research.com/articles/13-50/v1&title=" + encodeURIComponent(lTitle); var redditUrl = "http://reddit.com/submit?url=https://f1000research.com/articles/13-50/v1" + "&title=" + encodeURIComponent(lTitle); linkedInUrl += encodeURIComponent('Finn RD et al.'); var offsetTop = /chrome/i.test( navigator.userAgent ) ? 4 : -10; var addthis_config = { ui_offset_top: offsetTop, services_compact : "facebook,twitter,www.linkedin.com,www.mendeley.com,reddit.com", services_expanded : "facebook,twitter,www.linkedin.com,www.mendeley.com,reddit.com", services_custom : [ { name: "LinkedIn", url: linkedInUrl, icon:"/img/icon/at_linkedin.svg" }, { name: "Mendeley", url: "http://www.mendeley.com/import/?url=https://f1000research.com/articles/13-50/v1/mendeley", icon:"/img/icon/at_mendeley.svg" }, { name: "Reddit", url: redditUrl, icon:"/img/icon/at_reddit.svg" }, ] }; var addthis_share = { url: "https://f1000research.com/articles/13-50", templates : { twitter : "Establishing the ELIXIR Microbiome Community. Finn RD et al., published by " + "@F1000Research" + ", https://f1000research.com/articles/13-50/v1" } }; if (typeof(addthis) != "undefined"){ addthis.addEventListener('addthis.ready', checkCount); addthis.addEventListener('addthis.menu.share', checkCount); } $(".f1r-shares-twitter").attr("href", "https://twitter.com/intent/tweet?text=" + addthis_share.templates.twitter); $(".f1r-shares-facebook").attr("href", "https://www.facebook.com/sharer/sharer.php?u=" + addthis_share.url); $(".f1r-shares-linkedin").attr("href", addthis_config.services_custom[0].url); $(".f1r-shares-reddit").attr("href", addthis_config.services_custom[2].url); $(".f1r-shares-mendelay").attr("href", addthis_config.services_custom[1].url); function checkCount(){ setTimeout(function(){ $(".addthis_button_expanded").each(function(){ var count = $(this).text(); if (count !== "" && count != "0") $(this).removeClass("is-hidden"); else $(this).addClass("is-hidden"); }); }, 1000); } close How to cite this report {{reportCitation}} Cancel Copy Citation Details $(function(){R.ui.buttonDropdowns('.dropdown-for-downloads');}); $(function(){R.ui.toolbarDropdowns('.toolbar-dropdown-for-downloads');}); $.get("/articles/acj/144515/158321") new F1000.Clipboard(); new F1000.ThesaurusTermsDisplay("articles", "article", "158321"); $(document).ready(function() { $( "#frame1" ).on('load', function() { var mydiv = $(this).contents().find("div"); var h = mydiv.height(); console.log(h) }); var tooltipLivingFigure = jQuery(".interactive-living-figure-label .icon-more-info"), titleLivingFigure = tooltipLivingFigure.attr("title"); tooltipLivingFigure.simpletip({ fixed: true, position: ["-115", "30"], baseClass: 'small-tooltip', content:titleLivingFigure + " " }); tooltipLivingFigure.removeAttr("title"); $("body").on("click", ".cite-living-figure", function(e) { e.preventDefault(); var ref = $(this).attr("data-ref"); $(this).closest(".living-figure-list-container").find("#" + ref).fadeIn(200); }); $("body").on("click", ".close-cite-living-figure", function(e) { e.preventDefault(); $(this).closest(".popup-window-wrapper").fadeOut(200); }); $(document).on("mouseup", function(e) { var metricsContainer = $(".article-metrics-popover-wrapper"); if (!metricsContainer.is(e.target) && metricsContainer.has(e.target).length === 0) { $(".article-metrics-close-button").click(); } }); var articleId = $('#articleId').val(); if($("#main-article-count-box").attachArticleMetrics) { $("#main-article-count-box").attachArticleMetrics(articleId, { articleMetricsView: true }); } }); var figshareWidget = $(".new_figshare_widget"); if (figshareWidget.length > 0) { window.figshare.load("f1000", function(Widget) { // Select a tag/tags defined in your page. In this tag we will place the widget. _.map(figshareWidget, function(el){ var widget = new Widget({ articleId: $(el).attr("figshare_articleId") //height:300 // this is the height of the viewer part. [Default: 550] }); widget.initialize(); // initialize the widget widget.mount(el); // mount it in a tag that's on your page // this will save the widget on the global scope for later use from // your JS scripts. This line is optional. //window.widget = widget; }); }); } close Error Close Add Reset F1000.MICROSERVICES.AFFILIATION = ''; $(document).ready(function () { $('.js-affiliations-form').each((index, form) => { new AffiliationForm({ formId: form.id, institutionErrorSelector: '.comment-enter-institution', departmentErrorSelector: '.comment-enter-department', placeSelector: '.js-add-comment-place', stateSelector: '.js-add-comment-state', zipCodeSelector: '.js-add-comment-zipcode', countrySelector: '.js-add-comment-country', countryErrorSelector: '.comment-enter-country', }); }); }); $(document).ready(function () { var reportIds = { "412934": 11, "412935": 0, "412932": 0, "412933": 0, "412930": 0, "274944": 0, "274945": 0, "412931": 0, "412928": 0, "274946": 0, "412929": 0, "265756": 0, "265757": 0, "265758": 0, "265759": 0, "265755": 0, "265764": 0, "265760": 0, "265761": 0, "265762": 0, "265763": 0, "239543": 0, "239542": 0, "239541": 0, "239547": 0, "251067": 0, "239546": 0, "239545": 0, "239544": 0, "251071": 0, "239550": 0, "239549": 0, "239548": 0, "251075": 0, "251079": 0, "251083": 0, "251087": 0, "412360": 16, "412361": 0, "251091": 0, "251095": 0, "251099": 25, "251103": 0, "244839": 0, "244842": 0, "244840": 0, "244846": 0, "244845": 0, "244844": 0, "244849": 0, "244855": 49, "244853": 0, "244852": 0, "274940": 0, "274941": 0, "274942": 0, "274943": 0, "274937": 0, "274938": 0, "274939": 0, }; $(".referee-response-container,.js-referee-report").each(function(index, el) { var reportId = $(el).attr("data-reportid"), reportCount = reportIds[reportId] || 0; $(el).find(".comments-count-container,.js-referee-report-views").html(reportCount); }); var uuidInput = $("#article_uuid"), oldUUId = uuidInput.val(), newUUId = "db033e76-4bfd-4b54-be0f-12b4937eb36e"; uuidInput.val(newUUId); $("a[href*='article_uuid=']").each(function(index, el) { var newHref = $(el).attr("href").replace(oldUUId, newUUId); $(el).attr("href", newHref); }); }); An innovative open access publishing platform offering rapid publication and open peer review, whilst supporting data deposition and sharing. Browse Gateways Collections How it Works Contact For Developers Cookie Notice Privacy Notice RSS Submit Your Research Follow us © 2012-2026 F1000 Research Ltd. ISSN 2046-1402 | Legal | Partner of Research4Life • CrossRef • ORCID • FAIRSharing R.templateTests.simpleTemplate = R.template(' $text $text $text $text $text '); R.templateTests.runTests(); var F1000platform = new F1000.Platform({ name: "f1000research", displayName: "F1000Research", hostName: "f1000research.com", id: "1", editorialEmail: "[email protected]", infoEmail: "[email protected]", usePmcStats: true }); $(function(){R.ui.dropdowns('.dropdown-for-authors, .dropdown-for-about, .dropdown-for-myresearch');}); // $(function(){R.ui.dropdowns('.dropdown-for-referees');}); $(document).ready(function () { if ($(".cookie-warning").is(":visible")) { $(".sticky").css("margin-bottom", "35px"); $(".devices").addClass("devices-and-cookie-warning"); } $(".cookie-warning .close-button").click(function (e) { $(".devices").removeClass("devices-and-cookie-warning"); $(".sticky").css("margin-bottom", "0"); }); $("#tweeter-feed .tweet-message").each(function (i, message) { var self = $(message); self.html(linkify(self.html())); }); $(".partner").on("mouseenter mouseleave", function() { $(this).find(".gray-scale, .colour").toggleClass("is-hidden"); }); }); Sign In Remember me Forgotten your password? Sign In Cancel Email or password not correct. Please try again Please wait... $(function(){ // Note: All the setup needs to run against a name attribute and *not* the id due the clonish // nature of facebox... $("a[id=googleSignInButton]").click(function(event){ event.preventDefault(); $("input[id=oAuthSystem]").val("GOOGLE"); $("form[id=oAuthForm]").submit(); }); $("a[id=facebookSignInButton]").click(function(event){ event.preventDefault(); $("input[id=oAuthSystem]").val("FACEBOOK"); $("form[id=oAuthForm]").submit(); }); $("a[id=orcidSignInButton]").click(function(event){ event.preventDefault(); $("input[id=oAuthSystem]").val("ORCID"); $("form[id=oAuthForm]").submit(); }); }); If you've forgotten your password, please enter your email address below and we'll send you instructions on how to reset your password. The email address should be the one you originally registered with F1000. Email address not valid, please try again You registered with F1000 via Google, so we cannot reset your password. To sign in, please click here . If you still need help with your Google account password, please click here . You registered with F1000 via Facebook, so we cannot reset your password. To sign in, please click here . If you still need help with your Facebook account password, please click here . Code not correct, please try again Reset password Cancel Email us for further assistance. Server error, please try again. If your email address is registered with us, we will email you instructions to reset your password. If you think you should have received this email but it has not arrived, please check your spam filters and/or contact for further assistance. Please wait... Register $(document).ready(function () { signIn.createSignInAsRow($("#sign-in-form-gfb-popup")); $(".target-field").each(function () { var uris = $(this).val().split("/"); if (uris.pop() === "login") { $(this).val(uris.toString().replace(",","/")); } }); });

Text is read by the "Ask this paper" AI Q&A widget below. Extraction quality varies by source — PMC NXML preserves structure cleanly, OA-HTML may include some navigation residue, and OA-PDF can have broken hyphenation. The publisher copy (via DOI) is the canonical version.

My notes (saved in your browser only)

Ask this paper AI returns verbatim quotes from the full text · source: preprint-html

Answers must be backed by verbatim quotes from this paper's full text. Hallucinated quotes are dropped automatically; if no verbatim passage answers the question, we say so. How this works

Citation neighborhood (no data yet)

We don't have any in-corpus citations linked to this paper yet. This is a recent paper (2024) — citers typically take a year or two to land, and the OpenAlex reference graph may still be filling in.

Source provenance

europepmc
last seen: 2026-05-20T01:45:00.602351+00:00
unpaywall
last seen: 2026-05-20T11:00:21.680559+00:00
License: CC-BY-4.0