1#!/bin/sh 2# Copyright 2022 Google LLC 3# 4# This source code is licensed under the BSD-style license found in the 5# LICENSE file in the root directory of this source tree. 6 7################################### ARM NEON ################################## 8tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=8 -D DIV_ALGO=DIV -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-div-x8.c & 9tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=16 -D DIV_ALGO=DIV -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-div-x16.c & 10tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=24 -D DIV_ALGO=DIV -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-div-x24.c & 11tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=32 -D DIV_ALGO=DIV -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-div-x32.c & 12tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=40 -D DIV_ALGO=DIV -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-div-x40.c & 13tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=48 -D DIV_ALGO=DIV -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-div-x48.c & 14tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=56 -D DIV_ALGO=DIV -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-div-x56.c & 15tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=64 -D DIV_ALGO=DIV -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-div-x64.c & 16 17tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=8 -D DIV_ALGO=NR1FMA -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1fma-x8.c & 18tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=16 -D DIV_ALGO=NR1FMA -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1fma-x16.c & 19tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=24 -D DIV_ALGO=NR1FMA -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1fma-x24.c & 20tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=32 -D DIV_ALGO=NR1FMA -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1fma-x32.c & 21tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=40 -D DIV_ALGO=NR1FMA -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1fma-x40.c & 22tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=48 -D DIV_ALGO=NR1FMA -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1fma-x48.c & 23tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=56 -D DIV_ALGO=NR1FMA -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1fma-x56.c & 24tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=64 -D DIV_ALGO=NR1FMA -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1fma-x64.c & 25 26tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=8 -D DIV_ALGO=NR1RECPS -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1recps-x8.c & 27tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=16 -D DIV_ALGO=NR1RECPS -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1recps-x16.c & 28tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=24 -D DIV_ALGO=NR1RECPS -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1recps-x24.c & 29tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=32 -D DIV_ALGO=NR1RECPS -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1recps-x32.c & 30tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=40 -D DIV_ALGO=NR1RECPS -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1recps-x40.c & 31tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=48 -D DIV_ALGO=NR1RECPS -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1recps-x48.c & 32tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=56 -D DIV_ALGO=NR1RECPS -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1recps-x56.c & 33tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=64 -D DIV_ALGO=NR1RECPS -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1recps-x64.c & 34 35################################### x86 AVX2 ################################## 36tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=8 -D DIV_ALGO=div -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-div-x8.c & 37tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=16 -D DIV_ALGO=div -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-div-x16.c & 38tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=24 -D DIV_ALGO=div -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-div-x24.c & 39tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=32 -D DIV_ALGO=div -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-div-x32.c & 40tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=40 -D DIV_ALGO=div -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-div-x40.c & 41tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=48 -D DIV_ALGO=div -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-div-x48.c & 42tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=56 -D DIV_ALGO=div -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-div-x56.c & 43tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=64 -D DIV_ALGO=div -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-div-x64.c & 44 45tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=8 -D DIV_ALGO=rcp -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-rcp-x8.c & 46tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=16 -D DIV_ALGO=rcp -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-rcp-x16.c & 47tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=24 -D DIV_ALGO=rcp -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-rcp-x24.c & 48tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=32 -D DIV_ALGO=rcp -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-rcp-x32.c & 49tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=40 -D DIV_ALGO=rcp -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-rcp-x40.c & 50tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=48 -D DIV_ALGO=rcp -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-rcp-x48.c & 51tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=56 -D DIV_ALGO=rcp -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-rcp-x56.c & 52tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=64 -D DIV_ALGO=rcp -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-rcp-x64.c & 53 54################################## Unit tests ################################# 55tools/generate-vunary-test.py --spec test/f16-vsigmoid.yaml --output test/f16-vsigmoid.cc & 56 57wait 58