Azure
diff --git a/‎sdk/cosmos/azure-cosmos-spark_3-5/src/main/scala/com/azure/cosmos/spark/ItemsWriterBuilder.scala‎
Lines changed: 59 additions & 36 deletions b/‎sdk/cosmos/azure-cosmos-spark_3-5/src/main/scala/com/azure/cosmos/spark/ItemsWriterBuilder.scala‎
Lines changed: 59 additions & 36 deletions
diff --git a/‎sdk/cosmos/azure-cosmos-spark_3/docs/configuration-reference.md‎
Lines changed: 1 addition & 1 deletion b/‎sdk/cosmos/azure-cosmos-spark_3/docs/configuration-reference.md‎
Lines changed: 1 addition & 1 deletion
@@ -59,26 +59,29 @@ private class ItemsWriterBuilder
       )
     }
 
+    // Extract userConfig conversion to avoid repeated calls
+    private[this] val userConfigMap = userConfig.asCaseSensitiveMap().asScala.toMap
+
     private[this] val writeConfig = CosmosWriteConfig.parseWriteConfig(
-      userConfig.asCaseSensitiveMap().asScala.toMap,
+      userConfigMap,
       inputSchema
     )
 
     private[this] val containerConfig = CosmosContainerConfig.parseCosmosContainerConfig(
-      userConfig.asCaseSensitiveMap().asScala.toMap
+      userConfigMap
     )
 
     override def toBatch(): BatchWrite =
       new ItemsBatchWriter(
-        userConfig.asCaseSensitiveMap().asScala.toMap,
+        userConfigMap,
         inputSchema,
         cosmosClientStateHandles,
         diagnosticsConfig,
         sparkEnvironmentInfo)
 
     override def toStreaming: StreamingWrite =
       new ItemsBatchWriter(
-        userConfig.asCaseSensitiveMap().asScala.toMap,
+        userConfigMap,
         inputSchema,
         cosmosClientStateHandles,
         diagnosticsConfig,
@@ -88,6 +91,7 @@ private class ItemsWriterBuilder
 
     override def requiredDistribution(): Distribution = {
       if (writeConfig.bulkEnabled && writeConfig.bulkTransactional) {
+        log.logInfo("Transactional batch mode enabled - configuring data distribution by partition key columns")
         // For transactional writes, partition by all partition key columns
         val partitionKeyPaths = getPartitionKeyColumnNames()
         if (partitionKeyPaths.nonEmpty) {
@@ -125,44 +129,63 @@ private class ItemsWriterBuilder
 
     private def getPartitionKeyColumnNames(): Seq[String] = {
       try {
-        // Need to create a temporary container client to get partition key definition
-        val clientCacheItem = CosmosClientCache(
-          CosmosClientConfiguration(
-            userConfig.asCaseSensitiveMap().asScala.toMap,
-            ReadConsistencyStrategy.EVENTUAL,
-            sparkEnvironmentInfo
-          ),
-          Some(cosmosClientStateHandles.value.cosmosClientMetadataCaches),
-          "ItemsWriterBuilder-PKLookup"
-        )
-
-        val container = ThroughputControlHelper.getContainer(
-          userConfig.asCaseSensitiveMap().asScala.toMap,
-          containerConfig,
-          clientCacheItem,
-          None
-        )
-
-        val containerProperties = SparkBridgeInternal.getContainerPropertiesFromCollectionCache(container)
-        val partitionKeyDefinition = containerProperties.getPartitionKeyDefinition
-
-        // Release the client
-        clientCacheItem.close()
-
-        if (partitionKeyDefinition != null && partitionKeyDefinition.getPaths != null) {
-          val paths = partitionKeyDefinition.getPaths.asScala
-          paths.map(path => {
-            // Remove leading '/' from partition key path (e.g., "/pk" -> "pk")
-            if (path.startsWith("/")) path.substring(1) else path
-          }).toSeq
-        } else {
-          Seq.empty[String]
+        // Use loan pattern to ensure client is properly closed
+        using(createClientForPartitionKeyLookup()) { clientCacheItem =>
+          val container = ThroughputControlHelper.getContainer(
+            userConfigMap,
+            containerConfig,
+            clientCacheItem,
+            None
+          )
+
+          // Simplified retrieval using SparkBridgeInternal directly
+          val containerProperties = SparkBridgeInternal.getContainerPropertiesFromCollectionCache(container)
+          val partitionKeyDefinition = containerProperties.getPartitionKeyDefinition
+
+          extractPartitionKeyPaths(partitionKeyDefinition)
         }
       } catch {
         case ex: Exception =>
           log.logWarning(s"Failed to get partition key definition for transactional writes: ${ex.getMessage}")
           Seq.empty[String]
       }
     }
+
+    private def createClientForPartitionKeyLookup(): CosmosClientCacheItem = {
+      CosmosClientCache(
+        CosmosClientConfiguration(
+          userConfigMap,
+          ReadConsistencyStrategy.EVENTUAL,
+          sparkEnvironmentInfo
+        ),
+        Some(cosmosClientStateHandles.value.cosmosClientMetadataCaches),
+        "ItemsWriterBuilder-PKLookup"
+      )
+    }
+
+    private def extractPartitionKeyPaths(partitionKeyDefinition: com.azure.cosmos.models.PartitionKeyDefinition): Seq[String] = {
+      if (partitionKeyDefinition != null && partitionKeyDefinition.getPaths != null) {
+        val paths = partitionKeyDefinition.getPaths.asScala
+        if (paths.isEmpty) {
+          log.logError("Partition key definition has 0 columns - this should not happen for modern containers")
+        }
+        paths.map(path => {
+          // Remove leading '/' from partition key path (e.g., "/pk" -> "pk")
+          if (path.startsWith("/")) path.substring(1) else path
+        }).toSeq
+      } else {
+        log.logError("Partition key definition is null - this should not happen for modern containers")
+        Seq.empty[String]
+      }
+    }
+
+    // Scala loan pattern to ensure resources are properly cleaned up
+    private def using[A <: { def close(): Unit }, B](resource: A)(f: A => B): B = {
+      try {
+        f(resource)
+      } finally {
+        if (resource != null) resource.close()
+      }
+    }
   }
 }
@@ -67,7 +67,7 @@
 | `spark.cosmos.write.point.maxConcurrency`                       | None            | Cosmos DB Item Write Max concurrency. If not specified it will be determined based on the Spark executor VM Size                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                        |
 | `spark.cosmos.write.bulk.maxPendingOperations`                  | None            | Cosmos DB Item Write bulk mode maximum pending operations. Defines a limit of bulk operations being processed concurrently. If not specified it will be determined based on the Spark executor VM Size. If the volume of data is large for the provisioned throughput on the destination container, this setting can be adjusted by following the estimation of `1000 x Cores`                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                          |
 | `spark.cosmos.write.bulk.enabled`                               | `true`          | Cosmos DB Item Write bulk enabled                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                       |
-| `spark.cosmos.write.bulk.transactional`                         | `false`         | Enable transactional batch mode for bulk writes. When enabled, all operations for the same partition key are executed atomically (all succeed or all fail). Requires ordering and clustering by partition key columns. Only supports upsert operations. Cannot exceed 100 operations or 2MB per partition key. See [Transactional Batch documentation](https://learn.microsoft.com/azure/cosmos-db/nosql/transactional-batch) for details.                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                |\n| `spark.cosmos.write.bulk.targetedPayloadSizeInBytes`            | `220201`        | When the targeted payload size is reached for buffered documents, the request is sent to the backend. The default value is optimized for small documents <= 10 KB - when documents often exceed 110 KB, it can help to increase this value to up to about `1500000` (should still be smaller than 2 MB).                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                |
+| `spark.cosmos.write.bulk.transactional`                         | `false`         | Enable transactional batch mode for bulk writes. When enabled, all operations for the same partition key are executed atomically (all succeed or all fail). Requires ordering and clustering by partition key columns. Only supports upsert operations. Cannot exceed 100 operations or 2MB per partition key. **Note**: For containers using hierarchical partition keys (HPK), transactional scope applies only to **logical partitions** (complete partition key paths), not partial top-level keys. See [Transactional Batch documentation](https://learn.microsoft.com/azure/cosmos-db/nosql/transactional-batch) for details.                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                |\n| `spark.cosmos.write.bulk.targetedPayloadSizeInBytes`            | `220201`        | When the targeted payload size is reached for buffered documents, the request is sent to the backend. The default value is optimized for small documents <= 10 KB - when documents often exceed 110 KB, it can help to increase this value to up to about `1500000` (should still be smaller than 2 MB).                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                |
 | `spark.cosmos.write.bulk.initialBatchSize`                      | `100`           | Cosmos DB initial bulk micro batch size - a micro batch will be flushed to the backend when the number of documents enqueued exceeds this size - or the target payload size is met. The micro batch size is getting automatically tuned based on the throttling rate. By default the initial micro batch size is 100. Reduce this when you want to avoid that the first few requests consume too many RUs.                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                              |
 | `spark.cosmos.write.bulk.maxBatchSize`                          | `100`           | Cosmos DB max. bulk micro batch size - a micro batch will be flushed to the backend when the number of documents enqueued exceeds this size - or the target payload size is met. The micro batch size is getting automatically tuned based on the throttling rate. By default the max. micro batch size is 100. Use this setting only when migrating Spark 2.4 workloads - for other scenarios relying on the auto-tuning combined with throughput control will result in better experience.                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                            |
 | `spark.cosmos.write.flush.noProgress.maxIntervalInSeconds`      | `180`           | The time interval in seconds that write operations will wait when no progress can be made for bulk writes before forcing a retry. The retry will reinitialize the bulk write process - so, any delays on the retry can be sure to be actual service issues. The default value of 3 min should be sufficient to prevent false negatives when there is a short service-side write unavailability - like for partition splits or merges. Increase it only if you regularly see these transient errors to exceed a time period of 180 seconds.                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                                              |