14a2e7687Sjeremylt // Copyright (c) 2017-2018, Lawrence Livermore National Security, LLC. 24a2e7687Sjeremylt // Produced at the Lawrence Livermore National Laboratory. LLNL-CODE-734707. 34a2e7687Sjeremylt // All Rights reserved. See files LICENSE and NOTICE for details. 44a2e7687Sjeremylt // 54a2e7687Sjeremylt // This file is part of CEED, a collection of benchmarks, miniapps, software 64a2e7687Sjeremylt // libraries and APIs for efficient high-order finite element and spectral 74a2e7687Sjeremylt // element discretizations for exascale applications. For more information and 84a2e7687Sjeremylt // source code availability see http://github.com/ceed. 94a2e7687Sjeremylt // 104a2e7687Sjeremylt // The CEED research is supported by the Exascale Computing Project 17-SC-20-SC, 114a2e7687Sjeremylt // a collaborative effort of two U.S. Department of Energy organizations (Office 124a2e7687Sjeremylt // of Science and the National Nuclear Security Administration) responsible for 134a2e7687Sjeremylt // the planning and preparation of a capable exascale ecosystem, including 144a2e7687Sjeremylt // software, applications, hardware, advanced system engineering and early 154a2e7687Sjeremylt // testbed platforms, in support of the nation's exascale computing imperative. 164a2e7687Sjeremylt 174a2e7687Sjeremylt #include <string.h> 184a2e7687Sjeremylt #include "ceed-blocked.h" 194a2e7687Sjeremylt #include "../ref/ceed-ref.h" 204a2e7687Sjeremylt 214a2e7687Sjeremylt static int CeedOperatorDestroy_Blocked(CeedOperator op) { 224a2e7687Sjeremylt int ierr; 234ce2993fSjeremylt CeedOperator_Blocked *impl; 244ce2993fSjeremylt ierr = CeedOperatorGetData(op, (void*)&impl); CeedChk(ierr); 254a2e7687Sjeremylt 264a2e7687Sjeremylt for (CeedInt i=0; i<impl->numein+impl->numeout; i++) { 274a2e7687Sjeremylt ierr = CeedElemRestrictionDestroy(&impl->blkrestr[i]); CeedChk(ierr); 284a2e7687Sjeremylt ierr = CeedVectorDestroy(&impl->evecs[i]); CeedChk(ierr); 294a2e7687Sjeremylt } 304a2e7687Sjeremylt ierr = CeedFree(&impl->blkrestr); CeedChk(ierr); 314a2e7687Sjeremylt ierr = CeedFree(&impl->evecs); CeedChk(ierr); 324a2e7687Sjeremylt ierr = CeedFree(&impl->edata); CeedChk(ierr); 334a2e7687Sjeremylt 34*aedaa0e5Sjeremylt for (CeedInt i=0; i<impl->numein; i++) { 35*aedaa0e5Sjeremylt ierr = CeedVectorDestroy(&impl->qvecsin[i]); CeedChk(ierr); 364a2e7687Sjeremylt } 37*aedaa0e5Sjeremylt ierr = CeedFree(&impl->qvecsin); CeedChk(ierr); 384a2e7687Sjeremylt 39*aedaa0e5Sjeremylt for (CeedInt i=0; i<impl->numeout; i++) { 40*aedaa0e5Sjeremylt ierr = CeedVectorDestroy(&impl->qvecsout[i]); CeedChk(ierr); 41*aedaa0e5Sjeremylt } 42*aedaa0e5Sjeremylt ierr = CeedFree(&impl->qvecsout); CeedChk(ierr); 434a2e7687Sjeremylt 44fe2413ffSjeremylt ierr = CeedFree(&impl); CeedChk(ierr); 454a2e7687Sjeremylt return 0; 464a2e7687Sjeremylt } 474a2e7687Sjeremylt 484a2e7687Sjeremylt /* 494a2e7687Sjeremylt Setup infields or outfields 504a2e7687Sjeremylt */ 51fe2413ffSjeremylt static int CeedOperatorSetupFields_Blocked(CeedQFunction qf, CeedOperator op, 52fe2413ffSjeremylt bool inOrOut, 534a2e7687Sjeremylt CeedElemRestriction *blkrestr, 54*aedaa0e5Sjeremylt CeedVector *evecs, 55*aedaa0e5Sjeremylt CeedVector *qvecs, CeedInt starte, 564a2e7687Sjeremylt CeedInt numfields, CeedInt Q) { 57*aedaa0e5Sjeremylt CeedInt dim, ierr, ncomp; 58*aedaa0e5Sjeremylt Ceed ceed; 59*aedaa0e5Sjeremylt ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr); 60d1bcdac9Sjeremylt CeedBasis basis; 61d1bcdac9Sjeremylt CeedElemRestriction r; 62*aedaa0e5Sjeremylt CeedOperatorField *opfields; 63*aedaa0e5Sjeremylt CeedQFunctionField *qffields; 64fe2413ffSjeremylt if (inOrOut) { 65*aedaa0e5Sjeremylt ierr = CeedOperatorGetFields(op, NULL, &opfields); 66fe2413ffSjeremylt CeedChk(ierr); 67*aedaa0e5Sjeremylt ierr = CeedQFunctionGetFields(qf, NULL, &qffields); 68fe2413ffSjeremylt CeedChk(ierr); 69fe2413ffSjeremylt } else { 70*aedaa0e5Sjeremylt ierr = CeedOperatorGetFields(op, &opfields, NULL); 71fe2413ffSjeremylt CeedChk(ierr); 72*aedaa0e5Sjeremylt ierr = CeedQFunctionGetFields(qf, &qffields, NULL); 73fe2413ffSjeremylt CeedChk(ierr); 74fe2413ffSjeremylt } 754a2e7687Sjeremylt const CeedInt blksize = 8; 764a2e7687Sjeremylt 774a2e7687Sjeremylt // Loop over fields 784a2e7687Sjeremylt for (CeedInt i=0; i<numfields; i++) { 79d1bcdac9Sjeremylt CeedEvalMode emode; 80*aedaa0e5Sjeremylt ierr = CeedQFunctionFieldGetEvalMode(qffields[i], &emode); CeedChk(ierr); 814a2e7687Sjeremylt 824a2e7687Sjeremylt if (emode != CEED_EVAL_WEIGHT) { 83*aedaa0e5Sjeremylt ierr = CeedOperatorFieldGetElemRestriction(opfields[i], &r); 84d1bcdac9Sjeremylt CeedChk(ierr); 85fe2413ffSjeremylt CeedElemRestriction_Ref *data; 86fe2413ffSjeremylt ierr = CeedElemRestrictionGetData(r, (void *)&data); 874ce2993fSjeremylt Ceed ceed; 884ce2993fSjeremylt ierr = CeedElemRestrictionGetCeed(r, &ceed); CeedChk(ierr); 894ce2993fSjeremylt CeedInt nelem, elemsize, ndof, ncomp; 904ce2993fSjeremylt ierr = CeedElemRestrictionGetNumElements(r, &nelem); CeedChk(ierr); 914ce2993fSjeremylt ierr = CeedElemRestrictionGetElementSize(r, &elemsize); CeedChk(ierr); 924ce2993fSjeremylt ierr = CeedElemRestrictionGetNumDoF(r, &ndof); CeedChk(ierr); 934ce2993fSjeremylt ierr = CeedElemRestrictionGetNumComponents(r, &ncomp); CeedChk(ierr); 944ce2993fSjeremylt ierr = CeedElemRestrictionCreateBlocked(ceed, nelem, elemsize, 954ce2993fSjeremylt blksize, ndof, ncomp, 964a2e7687Sjeremylt CEED_MEM_HOST, CEED_COPY_VALUES, 97*aedaa0e5Sjeremylt data->indices, &blkrestr[i+starte]); 984a2e7687Sjeremylt CeedChk(ierr); 99*aedaa0e5Sjeremylt ierr = CeedElemRestrictionCreateVector(blkrestr[i+starte], NULL, 100*aedaa0e5Sjeremylt &evecs[i+starte]); 1014a2e7687Sjeremylt CeedChk(ierr); 1024a2e7687Sjeremylt } 1034a2e7687Sjeremylt 1044a2e7687Sjeremylt switch(emode) { 1054a2e7687Sjeremylt case CEED_EVAL_NONE: 106*aedaa0e5Sjeremylt ierr = CeedQFunctionFieldGetNumComponents(qffields[i], &ncomp); 107d1bcdac9Sjeremylt CeedChk(ierr); 108*aedaa0e5Sjeremylt ierr = CeedVectorCreate(ceed, Q*ncomp, &qvecs[i]); CeedChk(ierr); 109*aedaa0e5Sjeremylt break; 110*aedaa0e5Sjeremylt case CEED_EVAL_INTERP: 111*aedaa0e5Sjeremylt ierr = CeedQFunctionFieldGetNumComponents(qffields[i], &ncomp); 112*aedaa0e5Sjeremylt CeedChk(ierr); 113*aedaa0e5Sjeremylt ierr = CeedVectorCreate(ceed, Q*ncomp*blksize, &qvecs[i]); CeedChk(ierr); 1144a2e7687Sjeremylt break; 1154a2e7687Sjeremylt case CEED_EVAL_GRAD: 116*aedaa0e5Sjeremylt ierr = CeedOperatorFieldGetBasis(opfields[i], &basis); CeedChk(ierr); 117*aedaa0e5Sjeremylt ierr = CeedQFunctionFieldGetNumComponents(qffields[i], &ncomp); 118d1bcdac9Sjeremylt ierr = CeedBasisGetDimension(basis, &dim); CeedChk(ierr); 119*aedaa0e5Sjeremylt ierr = CeedVectorCreate(ceed, Q*ncomp*dim*blksize, &qvecs[i]); CeedChk(ierr); 1204a2e7687Sjeremylt break; 1214a2e7687Sjeremylt case CEED_EVAL_WEIGHT: // Only on input fields 122*aedaa0e5Sjeremylt ierr = CeedOperatorFieldGetBasis(opfields[i], &basis); CeedChk(ierr); 123*aedaa0e5Sjeremylt ierr = CeedVectorCreate(ceed, Q*blksize, &qvecs[i]); CeedChk(ierr); 124d1bcdac9Sjeremylt ierr = CeedBasisApply(basis, blksize, CEED_NOTRANSPOSE, 125*aedaa0e5Sjeremylt CEED_EVAL_WEIGHT, NULL, qvecs[i]); CeedChk(ierr); 126*aedaa0e5Sjeremylt 1274a2e7687Sjeremylt break; 1284a2e7687Sjeremylt case CEED_EVAL_DIV: 1294a2e7687Sjeremylt break; // Not implimented 1304a2e7687Sjeremylt case CEED_EVAL_CURL: 1314a2e7687Sjeremylt break; // Not implimented 1324a2e7687Sjeremylt } 1334a2e7687Sjeremylt } 1344a2e7687Sjeremylt return 0; 1354a2e7687Sjeremylt } 1364a2e7687Sjeremylt 1374a2e7687Sjeremylt /* 1384a2e7687Sjeremylt CeedOperator needs to connect all the named fields (be they active or passive) 1394a2e7687Sjeremylt to the named inputs and outputs of its CeedQFunction. 1404a2e7687Sjeremylt */ 1414a2e7687Sjeremylt static int CeedOperatorSetup_Blocked(CeedOperator op) { 1424a2e7687Sjeremylt int ierr; 1434ce2993fSjeremylt bool setupdone; 1444ce2993fSjeremylt ierr = CeedOperatorGetSetupStatus(op, &setupdone); CeedChk(ierr); 1454ce2993fSjeremylt if (setupdone) return 0; 146*aedaa0e5Sjeremylt Ceed ceed; 147*aedaa0e5Sjeremylt ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr); 1484ce2993fSjeremylt CeedOperator_Blocked *impl; 1494ce2993fSjeremylt ierr = CeedOperatorGetData(op, (void*)&impl); CeedChk(ierr); 1504ce2993fSjeremylt CeedQFunction qf; 1514ce2993fSjeremylt ierr = CeedOperatorGetQFunction(op, &qf); CeedChk(ierr); 1524ce2993fSjeremylt CeedInt Q, numinputfields, numoutputfields; 1534ce2993fSjeremylt ierr = CeedOperatorGetNumQuadraturePoints(op, &Q); CeedChk(ierr); 1544a2e7687Sjeremylt ierr= CeedQFunctionGetNumArgs(qf, &numinputfields, &numoutputfields); 1554a2e7687Sjeremylt CeedChk(ierr); 156d1bcdac9Sjeremylt CeedOperatorField *opinputfields, *opoutputfields; 157d1bcdac9Sjeremylt ierr = CeedOperatorGetFields(op, &opinputfields, &opoutputfields); 158d1bcdac9Sjeremylt CeedChk(ierr); 159d1bcdac9Sjeremylt CeedQFunctionField *qfinputfields, *qfoutputfields; 160d1bcdac9Sjeremylt ierr = CeedQFunctionGetFields(qf, &qfinputfields, &qfoutputfields); 161d1bcdac9Sjeremylt CeedChk(ierr); 1624a2e7687Sjeremylt 1634a2e7687Sjeremylt // Allocate 164*aedaa0e5Sjeremylt ierr = CeedCalloc(numinputfields + numoutputfields, &impl->blkrestr); 1654a2e7687Sjeremylt CeedChk(ierr); 166*aedaa0e5Sjeremylt ierr = CeedCalloc(numinputfields + numoutputfields, &impl->evecs); 1674a2e7687Sjeremylt CeedChk(ierr); 168*aedaa0e5Sjeremylt ierr = CeedCalloc(numinputfields + numoutputfields, &impl->edata); 1694a2e7687Sjeremylt CeedChk(ierr); 1704a2e7687Sjeremylt 171*aedaa0e5Sjeremylt ierr = CeedCalloc(16, &impl->qvecsin); CeedChk(ierr); 172*aedaa0e5Sjeremylt ierr = CeedCalloc(16, &impl->qvecsout); CeedChk(ierr); 1734a2e7687Sjeremylt 174*aedaa0e5Sjeremylt impl->numein = numinputfields; impl->numeout = numoutputfields; 175*aedaa0e5Sjeremylt 1764a2e7687Sjeremylt // Set up infield and outfield pointer arrays 1774a2e7687Sjeremylt // Infields 178*aedaa0e5Sjeremylt ierr = CeedOperatorSetupFields_Blocked(qf, op, 0, impl->blkrestr, 179*aedaa0e5Sjeremylt impl->evecs, impl->qvecsin, 0, 180*aedaa0e5Sjeremylt numinputfields, Q); 1814a2e7687Sjeremylt CeedChk(ierr); 1824a2e7687Sjeremylt // Outfields 183*aedaa0e5Sjeremylt ierr = CeedOperatorSetupFields_Blocked(qf, op, 1, impl->blkrestr, 184*aedaa0e5Sjeremylt impl->evecs, impl->qvecsout, 185*aedaa0e5Sjeremylt numinputfields, numoutputfields, Q); 1864a2e7687Sjeremylt CeedChk(ierr); 187*aedaa0e5Sjeremylt 188*aedaa0e5Sjeremylt // Temporary Vector 189*aedaa0e5Sjeremylt ierr = CeedVectorCreate(ceed, 0, &impl->tempvec); CeedChk(ierr); 1904a2e7687Sjeremylt 1914ce2993fSjeremylt ierr = CeedOperatorSetSetupDone(op); CeedChk(ierr); 1924a2e7687Sjeremylt 1934a2e7687Sjeremylt return 0; 1944a2e7687Sjeremylt } 1954a2e7687Sjeremylt 1964a2e7687Sjeremylt static int CeedOperatorApply_Blocked(CeedOperator op, CeedVector invec, 1974a2e7687Sjeremylt CeedVector outvec, CeedRequest *request) { 1984a2e7687Sjeremylt int ierr; 1994ce2993fSjeremylt CeedOperator_Blocked *impl; 2004ce2993fSjeremylt ierr = CeedOperatorGetData(op, (void*)&impl); CeedChk(ierr); 2014ce2993fSjeremylt const CeedInt blksize = 8; 202d1bcdac9Sjeremylt CeedInt Q, elemsize, numinputfields, numoutputfields, numelements, ncomp; 2034ce2993fSjeremylt ierr = CeedOperatorGetNumElements(op, &numelements); CeedChk(ierr); 2044ce2993fSjeremylt ierr = CeedOperatorGetNumQuadraturePoints(op, &Q); CeedChk(ierr); 2054ce2993fSjeremylt CeedInt nblks = (numelements/blksize) + !!(numelements%blksize); 2064ce2993fSjeremylt CeedQFunction qf; 2074ce2993fSjeremylt ierr = CeedOperatorGetQFunction(op, &qf); CeedChk(ierr); 2084ce2993fSjeremylt ierr= CeedQFunctionGetNumArgs(qf, &numinputfields, &numoutputfields); 2094ce2993fSjeremylt CeedChk(ierr); 2104dccadb6Sjeremylt CeedTransposeMode lmode; 211d1bcdac9Sjeremylt CeedOperatorField *opinputfields, *opoutputfields; 212d1bcdac9Sjeremylt ierr = CeedOperatorGetFields(op, &opinputfields, &opoutputfields); 213d1bcdac9Sjeremylt CeedChk(ierr); 214d1bcdac9Sjeremylt CeedQFunctionField *qfinputfields, *qfoutputfields; 215d1bcdac9Sjeremylt ierr = CeedQFunctionGetFields(qf, &qfinputfields, &qfoutputfields); 216d1bcdac9Sjeremylt CeedChk(ierr); 217d1bcdac9Sjeremylt CeedEvalMode emode; 218d1bcdac9Sjeremylt CeedVector vec; 219d1bcdac9Sjeremylt CeedBasis basis; 220d1bcdac9Sjeremylt CeedElemRestriction Erestrict; 2214a2e7687Sjeremylt 2224a2e7687Sjeremylt // Setup 2234a2e7687Sjeremylt ierr = CeedOperatorSetup_Blocked(op); CeedChk(ierr); 2244a2e7687Sjeremylt 2254a2e7687Sjeremylt // Input Evecs and Restriction 2264a2e7687Sjeremylt for (CeedInt i=0; i<numinputfields; i++) { 227d1bcdac9Sjeremylt ierr = CeedQFunctionFieldGetEvalMode(qfinputfields[i], &emode); 228d1bcdac9Sjeremylt CeedChk(ierr); 2294a2e7687Sjeremylt if (emode == CEED_EVAL_WEIGHT) { // Skip 2304a2e7687Sjeremylt } else { 231d1bcdac9Sjeremylt // Get input vector 232d1bcdac9Sjeremylt ierr = CeedOperatorFieldGetVector(opinputfields[i], &vec); CeedChk(ierr); 233d1bcdac9Sjeremylt if (vec == CEED_VECTOR_ACTIVE) 234d1bcdac9Sjeremylt vec = invec; 2354a2e7687Sjeremylt // Restrict 2364dccadb6Sjeremylt ierr = CeedOperatorFieldGetLMode(opinputfields[i], &lmode); CeedChk(ierr); 2374a2e7687Sjeremylt ierr = CeedElemRestrictionApply(impl->blkrestr[i], CEED_NOTRANSPOSE, 238d1bcdac9Sjeremylt lmode, vec, impl->evecs[i], 2394a2e7687Sjeremylt request); CeedChk(ierr); CeedChk(ierr); 2404a2e7687Sjeremylt // Get evec 2414a2e7687Sjeremylt ierr = CeedVectorGetArrayRead(impl->evecs[i], CEED_MEM_HOST, 2424a2e7687Sjeremylt (const CeedScalar **) &impl->edata[i]); 2434a2e7687Sjeremylt CeedChk(ierr); 2444a2e7687Sjeremylt } 2454a2e7687Sjeremylt } 2464a2e7687Sjeremylt 2474a2e7687Sjeremylt // Output Evecs 2484a2e7687Sjeremylt for (CeedInt i=0; i<numoutputfields; i++) { 2494a2e7687Sjeremylt ierr = CeedVectorGetArray(impl->evecs[i+impl->numein], CEED_MEM_HOST, 2504a2e7687Sjeremylt &impl->edata[i + numinputfields]); CeedChk(ierr); 2514a2e7687Sjeremylt } 2524a2e7687Sjeremylt 2534a2e7687Sjeremylt // Loop through elements 2544a2e7687Sjeremylt for (CeedInt e=0; e<nblks*blksize; e+=blksize) { 2554a2e7687Sjeremylt // Input basis apply if needed 2564a2e7687Sjeremylt for (CeedInt i=0; i<numinputfields; i++) { 2574a2e7687Sjeremylt // Get elemsize, emode, ncomp 258d1bcdac9Sjeremylt ierr = CeedOperatorFieldGetElemRestriction(opinputfields[i], &Erestrict); 259d1bcdac9Sjeremylt CeedChk(ierr); 260d1bcdac9Sjeremylt ierr = CeedElemRestrictionGetElementSize(Erestrict, &elemsize); 261d1bcdac9Sjeremylt CeedChk(ierr); 262d1bcdac9Sjeremylt ierr = CeedQFunctionFieldGetEvalMode(qfinputfields[i], &emode); 263d1bcdac9Sjeremylt CeedChk(ierr); 264d1bcdac9Sjeremylt ierr = CeedQFunctionFieldGetNumComponents(qfinputfields[i], &ncomp); 265d1bcdac9Sjeremylt CeedChk(ierr); 2664a2e7687Sjeremylt // Basis action 2674a2e7687Sjeremylt switch(emode) { 2684a2e7687Sjeremylt case CEED_EVAL_NONE: 269*aedaa0e5Sjeremylt ierr = CeedVectorSetArray(impl->qvecsin[i], CEED_MEM_HOST, 270*aedaa0e5Sjeremylt CEED_USE_POINTER, 271*aedaa0e5Sjeremylt &impl->edata[i][e*Q*ncomp]); CeedChk(ierr); 2724a2e7687Sjeremylt break; 2734a2e7687Sjeremylt case CEED_EVAL_INTERP: 274*aedaa0e5Sjeremylt ierr = CeedOperatorFieldGetBasis(opinputfields[i], &basis); CeedChk(ierr); 275*aedaa0e5Sjeremylt ierr = CeedVectorSetArray(impl->tempvec, CEED_MEM_HOST, 276*aedaa0e5Sjeremylt CEED_USE_POINTER, 277*aedaa0e5Sjeremylt &impl->edata[i][e*elemsize*ncomp]); 2784a2e7687Sjeremylt CeedChk(ierr); 279d1bcdac9Sjeremylt ierr = CeedBasisApply(basis, blksize, CEED_NOTRANSPOSE, 280*aedaa0e5Sjeremylt CEED_EVAL_INTERP, impl->tempvec, 281*aedaa0e5Sjeremylt impl->qvecsin[i]); CeedChk(ierr); 2824a2e7687Sjeremylt break; 2834a2e7687Sjeremylt case CEED_EVAL_GRAD: 284*aedaa0e5Sjeremylt ierr = CeedOperatorFieldGetBasis(opinputfields[i], &basis); CeedChk(ierr); 285*aedaa0e5Sjeremylt ierr = CeedVectorSetArray(impl->tempvec, CEED_MEM_HOST, 286*aedaa0e5Sjeremylt CEED_USE_POINTER, 287*aedaa0e5Sjeremylt &impl->edata[i][e*elemsize*ncomp]); 2884a2e7687Sjeremylt CeedChk(ierr); 289d1bcdac9Sjeremylt ierr = CeedBasisApply(basis, blksize, CEED_NOTRANSPOSE, 290*aedaa0e5Sjeremylt CEED_EVAL_GRAD, impl->tempvec, 291*aedaa0e5Sjeremylt impl->qvecsin[i]); CeedChk(ierr); 2924a2e7687Sjeremylt break; 2934a2e7687Sjeremylt case CEED_EVAL_WEIGHT: 2944a2e7687Sjeremylt break; // No action 2954a2e7687Sjeremylt case CEED_EVAL_DIV: 2964a2e7687Sjeremylt break; // Not implimented 2974a2e7687Sjeremylt case CEED_EVAL_CURL: 2984a2e7687Sjeremylt break; // Not implimented 2994a2e7687Sjeremylt } 3004a2e7687Sjeremylt } 3014a2e7687Sjeremylt 3024a2e7687Sjeremylt // Output pointers 3034a2e7687Sjeremylt for (CeedInt i=0; i<numoutputfields; i++) { 304d1bcdac9Sjeremylt ierr = CeedQFunctionFieldGetEvalMode(qfoutputfields[i], &emode); 305d1bcdac9Sjeremylt CeedChk(ierr); 3064a2e7687Sjeremylt if (emode == CEED_EVAL_NONE) { 307d1bcdac9Sjeremylt ierr = CeedQFunctionFieldGetNumComponents(qfoutputfields[i], &ncomp); 308d1bcdac9Sjeremylt CeedChk(ierr); 309*aedaa0e5Sjeremylt ierr = CeedVectorSetArray(impl->qvecsout[i], CEED_MEM_HOST, 310*aedaa0e5Sjeremylt CEED_USE_POINTER, 311*aedaa0e5Sjeremylt &impl->edata[i + numinputfields][e*Q*ncomp]); 312*aedaa0e5Sjeremylt CeedChk(ierr); 3134a2e7687Sjeremylt } 3144a2e7687Sjeremylt } 3154a2e7687Sjeremylt // Q function 316*aedaa0e5Sjeremylt ierr = CeedQFunctionApply(qf, Q*blksize, impl->qvecsin, impl->qvecsout); 317*aedaa0e5Sjeremylt CeedChk(ierr); 3184a2e7687Sjeremylt 3194a2e7687Sjeremylt // Output basis apply if needed 3204a2e7687Sjeremylt for (CeedInt i=0; i<numoutputfields; i++) { 3214a2e7687Sjeremylt // Get elemsize, emode, ncomp 322d1bcdac9Sjeremylt ierr = CeedOperatorFieldGetElemRestriction(opoutputfields[i], &Erestrict); 323d1bcdac9Sjeremylt CeedChk(ierr); 324d1bcdac9Sjeremylt ierr = CeedElemRestrictionGetElementSize(Erestrict, &elemsize); 325d1bcdac9Sjeremylt CeedChk(ierr); 326d1bcdac9Sjeremylt ierr = CeedQFunctionFieldGetEvalMode(qfoutputfields[i], &emode); 327d1bcdac9Sjeremylt CeedChk(ierr); 328d1bcdac9Sjeremylt ierr = CeedQFunctionFieldGetNumComponents(qfoutputfields[i], &ncomp); 329d1bcdac9Sjeremylt CeedChk(ierr); 3304a2e7687Sjeremylt // Basis action 3314a2e7687Sjeremylt switch(emode) { 3324a2e7687Sjeremylt case CEED_EVAL_NONE: 3334a2e7687Sjeremylt break; // No action 3344a2e7687Sjeremylt case CEED_EVAL_INTERP: 335d1bcdac9Sjeremylt ierr = CeedOperatorFieldGetBasis(opoutputfields[i], &basis); 336d1bcdac9Sjeremylt CeedChk(ierr); 337*aedaa0e5Sjeremylt ierr = CeedVectorSetArray(impl->tempvec, CEED_MEM_HOST, 338*aedaa0e5Sjeremylt CEED_USE_POINTER, 3394a2e7687Sjeremylt &impl->edata[i + numinputfields][e*elemsize*ncomp]); 340*aedaa0e5Sjeremylt ierr = CeedBasisApply(basis, blksize, CEED_TRANSPOSE, 341*aedaa0e5Sjeremylt CEED_EVAL_INTERP, impl->qvecsout[i], 342*aedaa0e5Sjeremylt impl->tempvec); CeedChk(ierr); 3434a2e7687Sjeremylt break; 3444a2e7687Sjeremylt case CEED_EVAL_GRAD: 345d1bcdac9Sjeremylt ierr = CeedOperatorFieldGetBasis(opoutputfields[i], &basis); 346d1bcdac9Sjeremylt CeedChk(ierr); 347*aedaa0e5Sjeremylt ierr = CeedVectorSetArray(impl->tempvec, CEED_MEM_HOST, 348*aedaa0e5Sjeremylt CEED_USE_POINTER, 349*aedaa0e5Sjeremylt &impl->edata[i + numinputfields][e*elemsize*ncomp]); 350d1bcdac9Sjeremylt ierr = CeedBasisApply(basis, blksize, CEED_TRANSPOSE, 351*aedaa0e5Sjeremylt CEED_EVAL_GRAD, impl->qvecsout[i], 352*aedaa0e5Sjeremylt impl->tempvec); CeedChk(ierr); 3534a2e7687Sjeremylt break; 3544ce2993fSjeremylt case CEED_EVAL_WEIGHT: { 3554ce2993fSjeremylt Ceed ceed; 3564ce2993fSjeremylt ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr); 3574ce2993fSjeremylt return CeedError(ceed, 1, 3584a2e7687Sjeremylt "CEED_EVAL_WEIGHT cannot be an output evaluation mode"); 3594a2e7687Sjeremylt break; // Should not occur 3604ce2993fSjeremylt } 3614a2e7687Sjeremylt case CEED_EVAL_DIV: 3624a2e7687Sjeremylt break; // Not implimented 3634a2e7687Sjeremylt case CEED_EVAL_CURL: 3644a2e7687Sjeremylt break; // Not implimented 3654a2e7687Sjeremylt } 3664a2e7687Sjeremylt } 3674a2e7687Sjeremylt } 3684a2e7687Sjeremylt 3694a2e7687Sjeremylt // Zero lvecs 370d1bcdac9Sjeremylt for (CeedInt i=0; i<numoutputfields; i++) { 371d1bcdac9Sjeremylt ierr = CeedOperatorFieldGetVector(opoutputfields[i], &vec); CeedChk(ierr); 372d1bcdac9Sjeremylt if (vec == CEED_VECTOR_ACTIVE) 373d1bcdac9Sjeremylt vec = outvec; 374d1bcdac9Sjeremylt ierr = CeedVectorSetValue(vec, 0.0); CeedChk(ierr); 3754a2e7687Sjeremylt } 3764a2e7687Sjeremylt 3774a2e7687Sjeremylt // Output restriction 3784a2e7687Sjeremylt for (CeedInt i=0; i<numoutputfields; i++) { 3794a2e7687Sjeremylt // Restore evec 3804a2e7687Sjeremylt ierr = CeedVectorRestoreArray(impl->evecs[i+impl->numein], 3814a2e7687Sjeremylt &impl->edata[i + numinputfields]); CeedChk(ierr); 382d1bcdac9Sjeremylt // Get output vector 383d1bcdac9Sjeremylt ierr = CeedOperatorFieldGetVector(opoutputfields[i], &vec); CeedChk(ierr); 3844a2e7687Sjeremylt // Active 385d1bcdac9Sjeremylt if (vec == CEED_VECTOR_ACTIVE) 386d1bcdac9Sjeremylt vec = outvec; 3874a2e7687Sjeremylt // Restrict 3884dccadb6Sjeremylt ierr = CeedOperatorFieldGetLMode(opoutputfields[i], &lmode); CeedChk(ierr); 3894a2e7687Sjeremylt ierr = CeedElemRestrictionApply(impl->blkrestr[i+impl->numein], CEED_TRANSPOSE, 390d1bcdac9Sjeremylt lmode, impl->evecs[i+impl->numein], vec, 3914a2e7687Sjeremylt request); CeedChk(ierr); 392d1bcdac9Sjeremylt 3934a2e7687Sjeremylt } 3944a2e7687Sjeremylt 3954a2e7687Sjeremylt // Restore input arrays 3964a2e7687Sjeremylt for (CeedInt i=0; i<numinputfields; i++) { 397d1bcdac9Sjeremylt ierr = CeedQFunctionFieldGetEvalMode(qfinputfields[i], &emode); 398d1bcdac9Sjeremylt CeedChk(ierr); 3994a2e7687Sjeremylt if (emode == CEED_EVAL_WEIGHT) { // Skip 4004a2e7687Sjeremylt } else { 4014a2e7687Sjeremylt ierr = CeedVectorRestoreArrayRead(impl->evecs[i], 402d1bcdac9Sjeremylt (const CeedScalar **) &impl->edata[i]); 403d1bcdac9Sjeremylt CeedChk(ierr); 4044a2e7687Sjeremylt } 4054a2e7687Sjeremylt } 4064a2e7687Sjeremylt 4074a2e7687Sjeremylt return 0; 4084a2e7687Sjeremylt } 4094a2e7687Sjeremylt 4104a2e7687Sjeremylt int CeedOperatorCreate_Blocked(CeedOperator op) { 4114a2e7687Sjeremylt int ierr; 412fe2413ffSjeremylt Ceed ceed; 413fe2413ffSjeremylt ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr); 4144ce2993fSjeremylt CeedOperator_Blocked *impl; 4154a2e7687Sjeremylt 4164a2e7687Sjeremylt ierr = CeedCalloc(1, &impl); CeedChk(ierr); 417fe2413ffSjeremylt ierr = CeedOperatorSetData(op, (void *)&impl); 418fe2413ffSjeremylt 419fe2413ffSjeremylt ierr = CeedSetBackendFunction(ceed, "Operator", op, "Apply", 420fe2413ffSjeremylt CeedOperatorApply_Blocked); CeedChk(ierr); 421fe2413ffSjeremylt ierr = CeedSetBackendFunction(ceed, "Operator", op, "Destroy", 422fe2413ffSjeremylt CeedOperatorDestroy_Blocked); CeedChk(ierr); 4234a2e7687Sjeremylt return 0; 4244a2e7687Sjeremylt } 425