This is an automated email from the ASF dual-hosted git repository.
morningman pushed a commit to branch master
in repository https://gitbox.apache.org/repos/asf/doris-website.git
The following commit(s) were added to refs/heads/master by this push:
new fe355c48ca8 weekly report 260921 (#4177)
fe355c48ca8 is described below
commit fe355c48ca842493fd76d809927eb345b42d0828
Author: Mingyu Chen (Rayner) <[email protected]>
AuthorDate: Mon Sep 28 14:58:22 2026 +0800
weekly report 260921 (#4177)
Co-authored-by: morningman <[email protected]>
---
src/pages/community-report/_reports/2026-09-21.mdx | 296 +++++++++++++++++++++
1 file changed, 296 insertions(+)
diff --git a/src/pages/community-report/_reports/2026-09-21.mdx
b/src/pages/community-report/_reports/2026-09-21.mdx
new file mode 100644
index 00000000000..c0577c1ad00
--- /dev/null
+++ b/src/pages/community-report/_reports/2026-09-21.mdx
@@ -0,0 +1,296 @@
+{/* Auto-generated by community-radar `report export-mdx` — do not edit by
hand.
+ `report` is English-only structured data laid out by doris-website. */}
+
+export const label = "Sep 21 – Sep 27, 2026";
+export const week = "Week 39, 2026";
+export const stats = [
+ {
+ "num": "97",
+ "label": "Merged PRs"
+ },
+ {
+ "num": "11",
+ "label": "New issues"
+ },
+ {
+ "num": "44",
+ "label": "Contributors"
+ }
+];
+
+export const report = {
+ "summary": {
+ "lead": "This week, Doris continued toward more complete lakehouse reads
and writes, more flexible semi-structured analytics, and broader query
capabilities for resource-constrained and highly concurrent workloads; 44
contributors helped merge 97 PRs, while the community opened 11 new issues.
Kernel and engineering remained the main focus, followed by real-time
warehousing, pointing to stronger execution foundations alongside broader
lakehouse interoperability and online query serving.",
+ "highlights": [
+ {
+ "title": "Toward a Complete Paimon Write Path",
+ "narrative": "The proposal forward-ports the complete Paimon write
path to master, including native bucket routing, JNI writing, and commit
payload transport, while upgrading to Paimon 1.4.2. It moves Doris toward
integrated lakehouse reads and writes.",
+ "prs": [
+ {
+ "num": 67395,
+ "title": "[feature](paimon) Forward-port Paimon table writes to
master",
+ "url": "https://github.com/apache/doris/pull/67395"
+ }
+ ]
+ },
+ {
+ "title": "Unified Lake and Changelog Reads with Fluss",
+ "narrative": "The proposed Fluss catalog supports log, primary-key,
and tiered tables. It merges snapshots with subsequent changes by key or scans
Paimon lake data alongside later logs, extending Doris toward unified analysis
of historical and recent data.",
+ "prs": [
+ {
+ "num": 66399,
+ "title": "[feat](fluss) Support reading Apache Fluss tables
through a fluss catalog",
+ "url": "https://github.com/apache/doris/pull/66399"
+ }
+ ]
+ },
+ {
+ "title": "Native Relational Operations for VARIANT",
+ "narrative": "The proposal connects VARIANT comparisons and hash joins
through a shared canonical comparator, enabling equality comparisons without
explicit casts and supporting multiple join and subquery forms. It makes
semi-structured values more directly usable in relational SQL analysis.",
+ "prs": [
+ {
+ "num": 67675,
+ "title": "[feature](variant) Support relational operations for
Variant",
+ "url": "https://github.com/apache/doris/pull/67675"
+ }
+ ]
+ },
+ {
+ "title": "Query Spill Expands to Windows and Object Storage",
+ "narrative": "The proposals add spill for eligible non-streaming
window functions and selectable S3 spill in cloud mode, with storage usage
visible through SHOW DATA. They aim to ease memory pressure from large
partitions and extend temporary storage options beyond local disks.",
+ "prs": [
+ {
+ "num": 68500,
+ "title": "[feature](be) Support spill for non-streaming window
functions",
+ "url": "https://github.com/apache/doris/pull/68500"
+ },
+ {
+ "num": 68032,
+ "title": "[feature](cloud) Spill to object storage in cloud mode
and report the traffic in SHOW DATA",
+ "url": "https://github.com/apache/doris/pull/68032"
+ }
+ ]
+ },
+ {
+ "title": "Batched Point Lookups for Concurrent Serving",
+ "narrative": "The proposals introduce optional RPC micro-batching and
batch row-store lookups, combining concurrent requests to the same BE and
grouping eligible literal IN lookup keys by tablet. They aim to reduce
scheduling and scanning overhead, extending Doris toward high-concurrency query
serving.",
+ "prs": [
+ {
+ "num": 67870,
+ "title": "[feature](perf)Support RPC micro-batching for
high-concurrency point queries",
+ "url": "https://github.com/apache/doris/pull/67870"
+ },
+ {
+ "num": 68446,
+ "title": "[improvement](point query) Add opt-in batch row-store
point queries",
+ "url": "https://github.com/apache/doris/pull/68446"
+ }
+ ]
+ }
+ ],
+ "numbers": {
+ "mergedPrs": 97,
+ "newIssues": 11,
+ "contributors": 44
+ }
+ },
+ "repos": [
+ {
+ "repo": "apache/doris",
+ "scenarios": [
+ {
+ "name": "Multi-modal Lakehouse",
+ "mergedNarrative": "For analytics combining AI and lake tables,
embedding models gain independent endpoint, credential and model settings, with
fallback to general AI configuration. Iceberg nested decimals now accept
precision widening at unchanged scale, improving service configuration
flexibility and schema-evolution compatibility.",
+ "merged": [
+ {
+ "num": 67673,
+ "title": "[Enhance](ai_func) Support dedicate embed properties
in AI RESOURCE",
+ "url": "https://github.com/apache/doris/pull/67673"
+ },
+ {
+ "num": 68378,
+ "title": "[fix](iceberg) Allow nested decimal precision
widening",
+ "url": "https://github.com/apache/doris/pull/68378"
+ }
+ ],
+ "inProgressNarrative": "Lakehouse work pushes selective partition
filtering into Hive Metastore to avoid full listings on heavily partitioned
tables. Newly opened Hive/Hudi partition-value pushdown would answer
partition-column queries without opening data files, extending metadata-driven
execution. Oracle floating-point mappings also broaden cross-source
compatibility.",
+ "inProgress": [
+ {
+ "num": 67725,
+ "title": "[improvement](hive) Push partition filters to HMS",
+ "url": "https://github.com/apache/doris/pull/67725"
+ },
+ {
+ "num": 68200,
+ "title": "[opt](jdbc) Support Oracle BINARY_FLOAT and
BINARY_DOUBLE types",
+ "url": "https://github.com/apache/doris/pull/68200"
+ },
+ {
+ "num": 68452,
+ "title": "[improvement](hive) Support
partition-column-value-only pushdown for Hive/Hudi scan",
+ "url": "https://github.com/apache/doris/pull/68452",
+ "isNew": true
+ }
+ ]
+ },
+ {
+ "name": "Real-time Data Warehouse",
+ "mergedNarrative": "Real-time warehouse work centers on cheaper
numeric DISTINCT merging, removing hash-set copies and intermediate-state
overhead. Lower VARIANT array ingestion CPU complements fixes to incremental
reads and cloud INSERT retries, improving processing reliability. Streaming
tests become more stable, while ST_Within extends spatial analysis.",
+ "merged": [
+ {
+ "num": 68350,
+ "title": "[improvement](be) Optimize numeric DISTINCT state
merging",
+ "url": "https://github.com/apache/doris/pull/68350"
+ }
+ ],
+ "inProgressNarrative": "Retention and ingestion work centers on
row/binlog TTL and one-time S3 imports. Configurable bucketing and
constraint-inferred colocated joins target data movement; gamma extends
analytics. Newly opened aggregate-state finalization and batch merging target
redundant aggregation, alongside new SQL plan management, load scheduling, and
cardinality-estimation work.",
+ "inProgress": [
+ {
+ "num": 65858,
+ "title": "[feature](storage) Support row-level TTL",
+ "url": "https://github.com/apache/doris/pull/65858"
+ },
+ {
+ "num": 66307,
+ "title": "[feature](fe) Add constraint-based colocate join
inference with distribution mappings",
+ "url": "https://github.com/apache/doris/pull/66307"
+ },
+ {
+ "num": 66477,
+ "title": "[feature](bucket) Support custom
distribution_hash_type for Hash Bucketing",
+ "url": "https://github.com/apache/doris/pull/66477"
+ },
+ {
+ "num": 68007,
+ "title": "[improvement](streaming) Support one-time S3 streaming
ingestion",
+ "url": "https://github.com/apache/doris/pull/68007"
+ },
+ {
+ "num": 68029,
+ "title": "[feature](function) Support gamma scalar function",
+ "url": "https://github.com/apache/doris/pull/68029"
+ },
+ {
+ "num": 68226,
+ "title": "[feature](storage) Support ROW binlog TTL",
+ "url": "https://github.com/apache/doris/pull/68226"
+ },
+ {
+ "num": 68312,
+ "title": "[feature](be) Add aggregate state finalizers and batch
merging",
+ "url": "https://github.com/apache/doris/pull/68312",
+ "isNew": true
+ },
+ {
+ "num": 68499,
+ "title": "[feature](plan)support sql plan management",
+ "url": "https://github.com/apache/doris/pull/68499",
+ "isNew": true
+ }
+ ]
+ },
+ {
+ "name": "Compute-Storage Separation & Cloud-Native",
+ "inProgressNarrative": "Cloud reads are moving toward access-aware
read-ahead, I/O coalescing, and cache-block hole filling to reduce excessive
point-query downloads and sequential-scan remote waits. BE-reported tablet
activity would also replace FE tracking windows, reducing frontend memory use
and providing more reliable activity inputs for cloud balancing.",
+ "inProgress": [
+ {
+ "num": 67292,
+ "title": "[feature](be) Add PageIO read-ahead, I/O coalescing,
and cache block hole filling",
+ "url": "https://github.com/apache/doris/pull/67292"
+ },
+ {
+ "num": 67621,
+ "title": "[improvement](tablet) Source active tablet stats from
BE reports",
+ "url": "https://github.com/apache/doris/pull/67621"
+ }
+ ]
+ },
+ {
+ "name": "Agent Observability",
+ "inProgressNarrative": "Search work targets correct join predicates
and field-specific analyzers, while smaller sparse VARIANT indexes and
per-field BM25 document counts aim to improve retrieval correctness, relevance,
and storage efficiency. Fuzzy matching broadens similarity analysis. Newly
opened parsing work makes JSON error behavior function-defined rather than
dependent on mutable BE settings.",
+ "inProgress": [
+ {
+ "num": 67932,
+ "title": "[feature](nereids) Support SEARCH in joins and
per-field analyzers",
+ "url": "https://github.com/apache/doris/pull/67932"
+ },
+ {
+ "num": 68202,
+ "title": "[improvement](inverted index) Shrink SNII indexes on
mostly NULL columns and use per-field BM25 document counts",
+ "url": "https://github.com/apache/doris/pull/68202"
+ },
+ {
+ "num": 67436,
+ "title": "[feature](function) Add jaro, jaro_winkler and
jaccard_similarity string functions",
+ "url": "https://github.com/apache/doris/pull/67436"
+ },
+ {
+ "num": 68447,
+ "title": "[improvement](variant) Remove BE configs that change
Variant JSON parse errors",
+ "url": "https://github.com/apache/doris/pull/68447",
+ "isNew": true
+ }
+ ]
+ },
+ {
+ "name": "Ecosystem Integration",
+ "inProgressNarrative": "Hive migration work targets missing
character-conversion semantics through compatible encode/decode functions,
reducing SQL rewrites. Restricting charset arguments to supported string
literals or NULL lets Doris select the conversion path before execution and
support vectorized processing.",
+ "inProgress": [
+ {
+ "num": 68131,
+ "title": "[feature](function) Support Hive-compatible encode and
decode",
+ "url": "https://github.com/apache/doris/pull/68131"
+ }
+ ]
+ },
+ {
+ "name": "Kernel & Engineering",
+ "mergedNarrative": "This week’s kernel work strengthens correctness
and stability under production workloads. Safer optimizer and conversion
semantics address wrong results, while fixes tackle crashes, stale data and
policy matching. Better materialized-view recovery, session lifecycles and
cache management reduce operational risk alongside scan overhead.",
+ "merged": [
+ {
+ "num": 66173,
+ "title": "[fix](routine-load) Persist and replay cancel reason,
and fix cleanup stall from corrupted final jobs",
+ "url": "https://github.com/apache/doris/pull/66173"
+ },
+ {
+ "num": 66530,
+ "title": "[opt](filescan) Reuse external scan tasks within a
statement",
+ "url": "https://github.com/apache/doris/pull/66530"
+ },
+ {
+ "num": 66963,
+ "title": "[improvement](ddl) Restrict state types to aggregate
tables",
+ "url": "https://github.com/apache/doris/pull/66963"
+ }
+ ],
+ "inProgressNarrative": "Shared asynchronous range-read
infrastructure and native 128-bit UUID support anchor work on remote I/O
efficiency and identifier storage. Broader hardening targets query correctness
and resource control; newly opened efforts address crashes, memory accounting,
cache and metadata lifecycles, and diagnostic gaps to make execution and
operations more dependable.",
+ "inProgress": [
+ {
+ "num": 67612,
+ "title": "[feature](be) Add file-range I/O coalescing and
asynchronous read infrastructure",
+ "url": "https://github.com/apache/doris/pull/67612"
+ },
+ {
+ "num": 67627,
+ "title": "[feature](datatype) Add UUID data type support",
+ "url": "https://github.com/apache/doris/pull/67627"
+ }
+ ]
+ }
+ ],
+ "demand": [
+ {
+ "name": "Multi-modal Lakehouse",
+ "narrative": "The community wants Doris’s Lance vector retrieval to
batch independent queries, each returning its own Top-K with a query index
identifying result ownership. Users could retrieve separate rankings in one
batch, distinct from combining vectors into one logical query or running
concurrent SQL statements.",
+ "refs": [
+ {
+ "num": 68394,
+ "title": "[Feature](lance) Support batch vector search with
per-query Top-K",
+ "url": "https://github.com/apache/doris/issues/68394"
+ }
+ ]
+ }
+ ]
+ }
+ ]
+};
---------------------------------------------------------------------
To unsubscribe, e-mail: [email protected]
For additional commands, e-mail: [email protected]