IsoSeq Clustering (Refine + Cluster)
Clustering: C57 I-K1 (individual)
Type
CWL
Status
succeeded
Engine
cwltool
Duration
0.2 h
Source Data
Pipeline
PacBio CCS (Subreads → HiFi)
IsoSeq Clustering (Refine + Cluster)
Run #79
(this run)
succeeded
1 sources
IsoSeq Annotation (Map + Collapse + SQANTI3)
Functional Annotation (TransDecoder + Pfam + SwissProt)
Combined From
- #74 — PacBio CCS (Subreads → HiFi) succeeded
Workflow
IsoSeq Clustering (Refine + Cluster)
#cwl
Software Tools
| Tool | Version | URL |
|---|---|---|
| cwltool | - | https://github.com/common-workflow-language/cwltool |
Results Summary
Input CCS Reads
115,377
FLNC Reads
112,575
Mean FLNC Length
0 nt
HQ Isoforms
0
LQ Isoforms
0
Clustering Ratio
0.0
FLNC / HQ isoforms
Mean FL Support
6.4
reads per isoform
Total FL Reads
74,169
Output Files
Provenance
| Execution | Expression quantification summary |
| Completed | 2026-03-04T12:46:04+00:00 |
RO-Crate 1.1
Workflow RO-Crate 1.0
FAIR
This analysis is packaged as a Research Object Crate
with machine-readable provenance and FAIR metadata.
RO-Crate Metadata (JSON-LD)
Show/hide raw JSON-LD
{
"@context": "https://w3id.org/ro/crate/1.1/context",
"@graph": [
{
"@id": "ro-crate-metadata.json",
"@type": "CreativeWork",
"about": {
"@id": "./"
},
"conformsTo": [
{
"@id": "https://w3id.org/ro/crate/1.1"
},
{
"@id": "https://w3id.org/workflowhub/workflow-ro-crate/1.0"
}
]
},
{
"@id": "./",
"@type": "Dataset",
"name": "IsoSeq Clustering (Refine + Cluster) \u2014 Run #79",
"description": "Generic IsoSeq3 clustering pipeline. Merges demultiplexed BAMs, runs primer removal + polyA filtering (refine), then clusters into HQ/LQ isoform consensus sequences. Compatible with Sequel I CCS (use_qvs=false) and Sequel II/IIe HiFi data.",
"datePublished": "2026-03-04",
"license": {
"@id": "https://creativecommons.org/licenses/by/4.0/"
},
"mainEntity": {
"@id": "isoseq_clustering.cwl"
},
"hasPart": [
{
"@id": "isoseq_clustering.cwl"
},
{
"@id": "job.yml"
},
{
"@id": "clustered.hq.bam"
},
{
"@id": "clustered.lq.bam"
},
{
"@id": "clustered.cluster_report.csv"
},
{
"@id": "clustered.cluster"
},
{
"@id": "flnc.filter_summary.report.json"
},
{
"@id": "flnc.bam"
},
{
"@id": "clustered.hq.bam.pbi"
},
{
"@id": "demux_primers.lima.summary"
},
{
"@id": "flnc.bam.pbi"
},
{
"@id": "clustered.lq.bam.pbi"
},
{
"@id": "results_summary.json"
},
{
"@id": "summary_extractor.py"
}
],
"mentions": [
{
"@id": "#execution"
},
{
"@id": "#summary-extraction"
}
]
},
{
"@id": "isoseq_clustering.cwl",
"@type": [
"File",
"SoftwareSourceCode",
"ComputationalWorkflow"
],
"name": "IsoSeq Clustering (Refine + Cluster)",
"description": "#cwl",
"programmingLanguage": {
"@id": "Generic IsoSeq3 clustering pipeline. Merges demultiplexed BAMs, runs primer removal + polyA filtering (refine), then clusters into HQ/LQ isoform consensus sequences. Compatible with Sequel I CCS (use_qvs=false) and Sequel II/IIe HiFi data."
},
"contentSize": "2.9 KB",
"sha256": "3cd8cfcc8caaf0fb4a964a16fc02edffbddb048e0c46b8c633c8fd3abf7efa08"
},
{
"@id": "#cwl",
"@type": "ComputerLanguage",
"name": "Common Workflow Language",
"url": {
"@id": "https://www.commonwl.org/"
},
"version": "1.2"
},
{
"@id": "#cwltool",
"@type": "SoftwareApplication",
"name": "cwltool",
"url": {
"@id": "https://github.com/common-workflow-language/cwltool"
}
},
{
"@id": "job.yml",
"@type": "File",
"name": "job.yml",
"description": "CWL job input parameters",
"encodingFormat": "text/yaml",
"contentSize": "274 B",
"sha256": "75851acb8daed3e8b1c23682cf900dfd5e98c26521b1b0e8a90065a7340ffe9a"
},
{
"@id": "clustered.hq.bam",
"@type": "File",
"name": "clustered.hq.bam",
"encodingFormat": "application/octet-stream",
"contentSize": "9.8 MB",
"sha256": "416def190f7944569271e7d760983a2014b196640c6314ca042a2d1429303a9e"
},
{
"@id": "clustered.lq.bam",
"@type": "File",
"name": "clustered.lq.bam",
"encodingFormat": "application/octet-stream",
"contentSize": "623 B",
"sha256": "a9d2d415fb4b0f8395e3aefadd7018ac358b907e84d51438fd4a396052d7806f"
},
{
"@id": "clustered.cluster_report.csv",
"@type": "File",
"name": "clustered.cluster_report.csv",
"encodingFormat": "text/csv",
"contentSize": "3.8 MB",
"sha256": "2ef3ffbcabf962a3dac2f893546d46b62acdcd6613d769088446597ecebd800f"
},
{
"@id": "clustered.cluster",
"@type": "File",
"name": "clustered.cluster",
"encodingFormat": "application/octet-stream",
"contentSize": "4.8 MB",
"sha256": "d26fa3f471b89d5a5e9659c12f8ac12f058fdebdc9a1f164f36b1905bfe71186"
},
{
"@id": "flnc.filter_summary.report.json",
"@type": "File",
"name": "flnc.filter_summary.report.json",
"encodingFormat": "application/json",
"contentSize": "819 B",
"sha256": "d1843d8a37672bf7616c79ea7e3d9b0647d5c3b6083f85e7ed88ea8356c350cd"
},
{
"@id": "flnc.bam",
"@type": "File",
"name": "flnc.bam",
"encodingFormat": "application/octet-stream",
"contentSize": "259.5 MB",
"sha256": "b559b9ef741c5eac8a500df8ed0b317a2531e06c0e0a594138a1434357097c77"
},
{
"@id": "clustered.hq.bam.pbi",
"@type": "File",
"name": "clustered.hq.bam.pbi",
"encodingFormat": "application/octet-stream",
"contentSize": "66.6 KB",
"sha256": "37f8468e5750999f11ec2ec58bfb34a28435097d28b30ce5f46c9b7dc224935d"
},
{
"@id": "demux_primers.lima.summary",
"@type": "File",
"name": "demux_primers.lima.summary",
"encodingFormat": "application/octet-stream",
"contentSize": "874 B",
"sha256": "5f2bbdc9c1e786d4b950d457939b8eef9e5ad8922cd052ddf60fd173d9c64d75"
},
{
"@id": "flnc.bam.pbi",
"@type": "File",
"name": "flnc.bam.pbi",
"encodingFormat": "application/octet-stream",
"contentSize": "1.2 MB",
"sha256": "02e5450dc4bc6cdb544f932136ef6c27f104e3b9e7f462028aac6176fcc48d2c"
},
{
"@id": "clustered.lq.bam.pbi",
"@type": "File",
"name": "clustered.lq.bam.pbi",
"encodingFormat": "application/octet-stream",
"contentSize": "65 B",
"sha256": "0898b440b5d691c4384b9cadaf7b007cfb06b947826234896416987b491df3ad"
},
{
"@id": "#execution",
"@type": "CreateAction",
"name": "IsoSeq Clustering (Refine + Cluster) execution",
"instrument": {
"@id": "isoseq_clustering.cwl"
},
"startTime": "2026-03-04T22:31:55+00:00",
"endTime": "2026-03-04T12:45:53+00:00",
"object": [
{
"@id": "job.yml"
}
],
"result": [
{
"@id": "clustered.hq.bam"
},
{
"@id": "clustered.lq.bam"
},
{
"@id": "clustered.cluster_report.csv"
},
{
"@id": "clustered.cluster"
},
{
"@id": "flnc.filter_summary.report.json"
},
{
"@id": "flnc.bam"
},
{
"@id": "clustered.hq.bam.pbi"
},
{
"@id": "demux_primers.lima.summary"
},
{
"@id": "flnc.bam.pbi"
},
{
"@id": "clustered.lq.bam.pbi"
}
]
},
{
"@id": "results_summary.json",
"@type": "File",
"name": "results_summary.json",
"description": "Derived summary statistics from pipeline outputs (CPM >= 1, uniquely mapped reads)",
"encodingFormat": "application/json",
"contentSize": "276 B",
"sha256": "703c8a51352c549a5660c41259bd19f9a547de9100545d5bc5ef171ddf630496"
},
{
"@id": "summary_extractor.py",
"@type": [
"File",
"SoftwareSourceCode"
],
"name": "Summary extraction script",
"description": "Python script that computed results_summary.json from pipeline outputs",
"programmingLanguage": {
"@id": "#python3"
}
},
{
"@id": "#python3",
"@type": "ComputerLanguage",
"name": "Python",
"url": {
"@id": "https://www.python.org/"
},
"version": "3"
},
{
"@id": "#summary-extraction",
"@type": "CreateAction",
"name": "Expression quantification summary",
"instrument": {
"@id": "summary_extractor.py"
},
"endTime": "2026-03-04T12:46:04+00:00",
"object": [
{
"@id": "OUT.read_assignments.tsv.gz"
},
{
"@id": "OUT.gene_counts.tsv"
},
{
"@id": "OUT.transcript_counts.tsv"
},
{
"@id": "OUT.extended_annotation.gtf"
},
{
"@id": "OUT.transcript_models.gtf"
}
],
"result": [
{
"@id": "results_summary.json"
}
]
}
]
}