@@ -798,10 +798,9 @@ func.func @depthwise_conv2d_dyn_w_h(%arg0: tensor<2x?x?x3xf32>, %arg1: tensor<3x
798
798
// CHECK: arith.subi
799
799
// CHECK: arith.muli
800
800
// CHECK: arith.divui
801
- // CHECK: [[CST0:%.+]] = arith.constant 0
802
801
// CHECK: %[[PADDED:.+]] = tensor.pad %arg0 low[0, 1, 3, 0] high[0, 2, 4, 0] {
803
802
// CHECK: ^bb0(%[[ARG3:[0-9a-zA-Z_]+]]: index, %[[ARG4:[0-9a-zA-Z_]+]]: index, %[[ARG5:[0-9a-zA-Z_]+]]: index, %[[ARG6:[0-9a-zA-Z_]+]]: index):
804
- // CHECK: tensor.yield [[CST0]] : f32
803
+ // CHECK: tensor.yield %cst : f32
805
804
// CHECK: } : tensor<2x?x?x3xf32> to tensor<2x?x?x3xf32>
806
805
// CHECK: %[[CONV:.+]] = linalg.depthwise_conv_2d_nhwc_hwcm {dilations = dense<[2, 1]> : tensor<2xi64>, strides = dense<[1, 2]> : tensor<2xi64>} ins(%[[PADDED]], %arg1 : tensor<2x?x?x3xf32>, tensor<3x6x3x5xf32>) outs(%{{.*}} : tensor<2x?x?x3x5xf32>) -> tensor<2x?x?x3x5xf32>
807
806
// CHECK: %[[COLLAPSED:.+]] = tensor.collapse_shape %[[CONV]] {{\[}}[0], [1], [2], [3, 4]]
@@ -813,30 +812,6 @@ func.func @depthwise_conv2d_dyn_w_h(%arg0: tensor<2x?x?x3xf32>, %arg1: tensor<3x
813
812
814
813
// -----
815
814
816
- // CHECK: #[[$MAP0:.*]] = affine_map<(d0, d1, d2, d3) -> (d3)>
817
- // CHECK: #[[$MAP1:.*]] = affine_map<(d0, d1, d2, d3) -> (d0, d1, d2, d3)>
818
-
819
- // CHECK-LABEL: @depthwise_int_conv_zero_zp
820
- func.func @depthwise_int_conv_zero_zp (%arg0 : tensor <1 x7 x5 x3 xi8 >, %arg1 : tensor <3 x1 x3 x11 xi8 >, %arg2 : tensor <33 xi32 >) -> () {
821
- // CHECK: [[INIT:%.+]] = tensor.empty()
822
- // CHECK: [[CST0:%.+]] = arith.constant 0
823
- // CHECK: [[FILL:%.+]] = linalg.fill ins([[CST0]]{{.*}}outs([[INIT]]
824
- // CHECK: [[OUT:%.+]] = tensor.empty()
825
- // CHECK: [[DEPTH:%.+]] = linalg.depthwise_conv_2d_nhwc_hwcm {dilations = dense<1> : tensor<2xi64>, strides = dense<1> : tensor<2xi64>} ins(%arg0, %arg1 : tensor<1x7x5x3xi8>, tensor<3x1x3x11xi8>) outs([[FILL]] : tensor<1x5x5x3x11xi32>)
826
- // CHECK: [[COLLAPSED:%.+]] = tensor.collapse_shape [[DEPTH]] {{\[}}[0], [1], [2], [3, 4]]
827
- // CHECK: [[BIAS:%.+]] = linalg.generic {indexing_maps = [#[[$MAP0]], #[[$MAP1]], #[[$MAP1]]], iterator_types = ["parallel", "parallel", "parallel", "parallel"]} ins(%arg2, [[COLLAPSED]] : tensor<33xi32>, tensor<1x5x5x33xi32>) outs([[OUT]] : tensor<1x5x5x33xi32>) {
828
- // CHECK: ^bb0(%[[ARG3:[0-9a-zA-Z_]+]]: i32, %[[ARG4:[0-9a-zA-Z_]+]]: i32, %[[ARG5:[0-9a-zA-Z_]+]]: i32):
829
- // CHECK: [[ADD:%.+]] = arith.addi %[[ARG3]], %[[ARG4]] : i32
830
- // CHECK: linalg.yield [[ADD]] : i32
831
- // CHECK: } -> tensor<1x5x5x33xi32>
832
- %input_zp = " tosa.const" () <{value = dense <0 > : tensor <1 xi8 >}> : () -> tensor <1 xi8 >
833
- %weight_zp = " tosa.const" () <{value = dense <0 > : tensor <1 xi8 >}> : () -> tensor <1 xi8 >
834
- %2 = tosa.depthwise_conv2d %arg0 , %arg1 , %arg2 , %input_zp , %weight_zp {acc_type = i32 , pad = array<i64 : 0 , 0 , 0 , 0 >, stride = array<i64 : 1 , 1 >, dilation = array<i64 : 1 , 1 > } : (tensor <1 x7 x5 x3 xi8 >, tensor <3 x1 x3 x11 xi8 >, tensor <33 xi32 >, tensor <1 xi8 >, tensor <1 xi8 >) -> tensor <1 x5 x5 x33 xi32 >
835
- return
836
- }
837
-
838
- // -----
839
-
840
815
// CHECK: #[[$MAP1:.+]] = affine_map<(d0, d1, d2, d3, d4) -> (d4)>
841
816
// CHECK: #[[$MAP2:.+]] = affine_map<(d0, d1, d2, d3, d4) -> (d0, d1, d2, d3, d4)>
842
817
0 commit comments